{"slug":"surge-rlhf","name":"RLHF","ownerSlug":"surge","artifactType":"Dataset","availabilityKind":"custom-capability","access":["Custom-built","Commercial","Request demo"],"summary":"Human preference and reward data for reinforcement learning from human feedback.","description":"Surge reports generating preference and reward data intended to capture nuanced judgments for model alignment and post-training.","domains":["rlhf","alignment","multi-domain"],"capabilities":{"human-in-the-loop":"yes","expert-authored":"yes","multi-turn":"unknown"},"verifierType":"Human preference and reward signals","lastVerified":"2026-08-16","verificationStatus":"company-reported","confidence":"high","links":[{"label":"Product page","url":"https://surgehq.ai/products","kind":"official"}],"sourceIds":["surge-products"],"canonicalUrl":"https://envindex.com/environments/surge-rlhf","owner":{"slug":"surge","name":"Surge","website":"https://surgehq.ai","tagline":"Bootstrapped human-data leader with a dedicated RL environments org","description":"Surge AI is the largest human data / RLHF vendor by revenue ($1.2B+ in 2024) serving OpenAI, Google, Anthropic, and Meta, and has created an internal organization dedicated to RL environments. It builds expert-graded evals and environment products such as EnterpriseBench.","domains":["multi-domain","data-labeling","rlhf"],"locations":["San Francisco","New York","Seattle"],"founders":[{"name":"Edwin Chen","xHandle":null}],"teamSize":"101-250","funding":"Bootstrapped; reported in talks to raise ~$1B at $25B+ valuation (2025)","raising":true,"founded":2020,"notable":"~$9.9M revenue per employee; CoreCraft benchmark"},"sources":[{"id":"surge-products","url":"https://surgehq.ai/products","title":"Products","publisher":"Surge AI","sourceType":"official","accessedAt":"2026-08-16","verificationStatus":"company-reported"}]}