{"slug":"scale","name":"Scale","website":"https://scale.com","tagline":"Data-labeling incumbent extending into agent evals and RL environments","description":"Scale AI is the original data-labeling powerhouse for AI labs and enterprises, now building RL environments and agent evaluation products under its agents and RL environments group. After Meta's 2025 investment and the departure of CEO Alexandr Wang, it lost some lab customers but continues to push into environments.","domains":["multi-domain","coding","data-labeling","rlhf"],"locations":["San Francisco","New York","Washington DC","London"],"founders":[{"name":"Alexandr Wang","xHandle":"alexandr_wang"},{"name":"Lucy Guo","xHandle":null}],"teamSize":"250+","funding":"$1.6B+ raised; Meta invested $14.3B at ~$29B valuation (June 2025)","raising":null,"founded":2016,"notable":"Publishes SWE-bench Pro benchmark","entityType":"Incumbent","canonicalUrl":"https://envindex.com/startup/scale","research":{"startupSlug":"scale","entityType":"Incumbent","lastResearched":"2026-08-16","lastVerified":"2026-08-16","verificationStatus":"confirmed","confidence":"high","sourceIds":["scale-swe-pro","scale-swe-public"]},"sources":[{"id":"scale-swe-pro","url":"https://labs.scale.com/papers/swe_bench_pro","title":"SWE-Bench Pro: Can AI Agents Solve Long-Horizon Software Engineering Tasks?","publisher":"Scale Labs","sourceType":"paper","publishedAt":"2025-09-19","accessedAt":"2026-08-16","verificationStatus":"confirmed"},{"id":"scale-swe-public","url":"https://labs.scale.com/leaderboard/swe_bench_pro_public","title":"SWE-Bench Pro  -  Public Dataset","publisher":"Scale Labs","sourceType":"official","accessedAt":"2026-08-16","verificationStatus":"confirmed"}],"capabilityCatalog":{"startupSlug":"scale","reviewedAt":"2026-08-16","status":"cataloged","items":[{"name":"Agentic Data and Evaluations","kind":"Data service","summary":"Data and evaluation programs for agentic AI systems.","access":["Commercial"],"sourceUrl":"https://scale.com/enterprise/agentic-solutions","sourceLabel":"Agentic Solutions","verificationStatus":"company-reported","lastVerified":"2026-08-16"},{"name":"SWE-Bench Pro","kind":"Benchmark","summary":"A long-horizon software-engineering benchmark built from public and proprietary repositories.","access":["Public / Open","Open source"],"sourceUrl":"/environments/swe-bench-pro","sourceLabel":"EnvIndex catalog page","verificationStatus":"confirmed","lastVerified":"2026-08-16","artifactSlug":"swe-bench-pro"},{"name":"SWE-Bench Pro  -  Private Dataset","kind":"Benchmark","summary":"The proprietary-repository subset of SWE-Bench Pro, documented publicly but not distributed as an open dataset.","access":["Private catalog","Enterprise-only"],"sourceUrl":"/environments/swe-bench-pro-private","sourceLabel":"EnvIndex catalog page","verificationStatus":"company-reported","lastVerified":"2026-08-16","artifactSlug":"swe-bench-pro-private"}]},"artifacts":[{"slug":"swe-bench-pro","name":"SWE-Bench Pro","ownerSlug":"scale","artifactType":"Benchmark","availabilityKind":"public-sample","access":["Public / Open","Open source"],"summary":"A long-horizon software-engineering benchmark built from public and proprietary repositories.","description":"SWE-Bench Pro evaluates agents on realistic software-engineering problems. The published paper reports 1,865 problems from 41 actively maintained repositories; the public dataset and leaderboard are separately accessible.","domains":["coding","long-horizon"],"capabilities":{"long-horizon":"yes","code-execution":"yes","sandboxed":"yes","human-in-the-loop":"yes","expert-authored":"yes"},"taskCount":1865,"verifierType":"Repository test suites with human verification checkpoints","firstReleased":"2025-09-19","lastVerified":"2026-08-16","verificationStatus":"confirmed","confidence":"high","links":[{"label":"Research page","url":"https://labs.scale.com/papers/swe_bench_pro","kind":"paper"},{"label":"Public leaderboard","url":"https://labs.scale.com/leaderboard/swe_bench_pro_public","kind":"leaderboard"}],"sourceIds":["scale-swe-pro","scale-swe-public"],"relatedArtifactSlugs":["swe-bench-pro-private"]},{"slug":"swe-bench-pro-private","name":"SWE-Bench Pro  -  Private Dataset","ownerSlug":"scale","artifactType":"Benchmark","availabilityKind":"private-commercial","access":["Private catalog","Enterprise-only"],"summary":"The proprietary-repository subset of SWE-Bench Pro, documented publicly but not distributed as an open dataset.","description":"This private subset evaluates coding agents against commercial-grade proprietary repositories. Its existence and leaderboard are public; the underlying repository contents are not presented as publicly downloadable by EnvIndex.","domains":["coding","long-horizon","enterprise"],"capabilities":{"long-horizon":"yes","code-execution":"yes","sandboxed":"yes","human-in-the-loop":"yes","production-derived":"yes"},"verifierType":"Repository test suites","lastVerified":"2026-08-16","verificationStatus":"company-reported","confidence":"medium","links":[{"label":"Private dataset leaderboard","url":"https://labs.scale.com/leaderboard/swe_bench_pro_private","kind":"leaderboard"}],"sourceIds":["scale-swe-pro","scale-swe-private"],"relatedArtifactSlugs":["swe-bench-pro"],"editorialNote":"Private describes access, not verification. EnvIndex has not inspected the proprietary repository contents."}]}