{"id":"61eb3030-04a2-404c-9c04-788f98a0f3cf","name":"Baseten","slug":"baseten","description":"Baseten is a San Francisco inference platform that sits on top of GPU capacity rather than owning it. Founded in 2019 by ex-Gumroad machine-learning engineers Tuhin Srivastava, Amir Haghighat, and Pankaj Gupta, the company packages open-source and customer-owned models behind a low-latency serving layer, its Truss packaging framework, and the Chains orchestration runtime for compound AI pipelines. In the neocloud landscape it is positioned as an inference-focused mid-market and long-tail platform, not a hyperscale training host: workloads run across the company's Inference Cloud on partner GPU capacity in AWS, GCP and dedicated colocation, with Baseten selling engineering support, autoscaling and cost optimization on top.\n\nThe company has become one of the more visible pure-play inference vendors thanks to a run of AI-native design-partner customers including Abridge, Cursor, Notion, Writer, Descript, Sourcegraph, Clay, Zed and Gamma. Baseten does not publish a fixed hourly GPU price sheet in the CoreWeave or Lambda style; pricing is a per-model, per-token or dedicated-deployment quote, which is part of how it differentiates against raw GPU-hour resellers.\n\nBaseten closed a $150 million Series D led by BOND in September 2025, with CapitalG, Conviction, IVP, Spark Capital, Greylock, 01A, Premji Invest, BoxGroup and Scribble Ventures participating, and Jay Simons joining the board. That round followed a $75 million Series C led by IVP and Spark Capital in February 2025. The company remains privately held.","status":"active","founded":2019,"hq":"San Francisco, California, United States","hqLat":null,"hqLng":null,"hqGeocodedAt":null,"hqAddress":null,"hqPrecision":null,"website":"https://www.baseten.co","fundingTotal":"300000000","type":null,"roles":[],"reviewStatus":"reviewed","heatIndex":0,"lifecycleState":"active","supersededByCompanyId":null,"sources":[{"url":"https://www.baseten.co/blog/announcing-baseten-150m-series-d/","date":"2025-09-05","title":"Announcing Baseten's $150M Series D","publisher":"Baseten"},{"url":"https://www.baseten.co/blog/announcing-baseten-75m-series-c/","date":"2025-02-25","title":"Announcing Baseten's $75M Series C","publisher":"Baseten"},{"url":"https://www.baseten.co/about/","title":"About Baseten","publisher":"Baseten"},{"url":"https://www.baseten.co/blog/","title":"Baseten engineering blog","publisher":"Baseten"}],"keyFacts":[{"label":"Headquarters","value":"San Francisco, California","sourceUrl":"https://www.baseten.co"},{"label":"Founded","value":"2019","sourceUrl":"https://www.baseten.co"},{"label":"Positioning","value":"Inference-focused platform serving mid-market AI-native customers; runs on partner GPU capacity rather than owning datacenters","sourceUrl":"https://www.baseten.co/blog/"},{"label":"Most recent round","value":"$150M Series D, September 5, 2025, led by BOND (CapitalG, Conviction, IVP, Spark, Greylock, 01A, Premji Invest, BoxGroup, Scribble Ventures participating); Jay Simons joined the board","sourceUrl":"https://www.baseten.co/blog/announcing-baseten-150m-series-d/"},{"label":"Prior round","value":"$75M Series C, February 25, 2025, co-led by IVP and Spark Capital with Greylock, Conviction, South Park Commons, Basecase, Lachy Groom and 01A (Adam Bain, Dick Costolo)","sourceUrl":"https://www.baseten.co/blog/announcing-baseten-75m-series-c/"},{"label":"Named customers","value":"Abridge, Cursor, Notion, Writer, Descript, Sourcegraph, Clay, Zed, Gamma, Bland, Lovable, OpenEvidence","sourceUrl":"https://www.baseten.co/about/"},{"label":"Datacenter model","value":"Multi-cloud Inference Cloud on partner GPU capacity (hyperscalers and colocation); no publicly disclosed owned MW or GPU inventory","sourceUrl":"https://www.baseten.co/blog/"},{"label":"Announced hourly pricing","value":"No standard public per-GPU-hour rate card; sold as per-model API pricing or dedicated deployments quoted per customer","sourceUrl":"https://www.baseten.co/blog/"},{"label":"Open-source stack","value":"Truss (model packaging) and Chains (compound AI orchestration) released under permissive license","sourceUrl":"https://www.baseten.co/blog/"},{"label":"Ownership status","value":"Private, venture-backed","sourceUrl":"https://www.baseten.co/blog/announcing-baseten-150m-series-d/"}],"aliases":[],"collisionRisk":"low","reviewNote":null,"atsProvider":null,"atsSlug":null,"factoryLocations":null,"manufacturingCapacity":null,"orgType":null,"burnSignal":null,"litigationEvents":null,"insurancePostureDisclosed":null,"unitEconomicsDisclosed":null,"tickerSymbol":null,"stockExchange":null,"secCik":null,"revenueUsd":null,"revenueBasis":null,"revenueAsOf":null,"employeeCount":null,"employeeCountBasis":null,"employeeCountAsOf":null,"marketCapUsd":null,"marketCapAsOf":null,"marketCapSource":null,"createdAt":"2026-08-28T01:48:22.713Z","updatedAt":"2026-08-28T03:45:17.421Z","jsonLd":{"@context":"https://schema.org","@type":"Organization","@id":"https://registry.deploy.report/companies/baseten","url":"https://registry.deploy.report/companies/baseten","name":"Baseten","description":"Baseten is a San Francisco inference platform that sits on top of GPU capacity rather than owning it. Founded in 2019 by ex-Gumroad machine-learning engineers Tuhin Srivastava, Amir Haghighat, and Pankaj Gupta, the company packages open-source and customer-owned models behind a low-latency serving layer, its Truss packaging framework, and the Chains orchestration runtime for compound AI pipelines. In the neocloud landscape it is positioned as an inference-focused mid-market and long-tail platform, not a hyperscale training host: workloads run across the company's Inference Cloud on partner GPU capacity in AWS, GCP and dedicated colocation, with Baseten selling engineering support, autoscaling and cost optimization on top.\n\nThe company has become one of the more visible pure-play inference vendors thanks to a run of AI-native design-partner customers including Abridge, Cursor, Notion, Writer, Descript, Sourcegraph, Clay, Zed and Gamma. Baseten does not publish a fixed hourly GPU price sheet in the CoreWeave or Lambda style; pricing is a per-model, per-token or dedicated-deployment quote, which is part of how it differentiates against raw GPU-hour resellers.\n\nBaseten closed a $150 million Series D led by BOND in September 2025, with CapitalG, Conviction, IVP, Spark Capital, Greylock, 01A, Premji Invest, BoxGroup and Scribble Ventures participating, and Jay Simons joining the board. That round followed a $75 million Series C led by IVP and Spark Capital in February 2025. The company remains privately held.","identifier":"61eb3030-04a2-404c-9c04-788f98a0f3cf","foundingDate":"2019","address":"San Francisco, California, United States","publisher":{"@id":"https://deploy.report/#organization"}},"framework_metadata":{"framework_schema_version":"0.1.0","verification_status":"verified","maturity_stage":null,"lifecycle_state":null,"architectural_position":{"cohort":null,"sub_cohorts":[]},"within_cohort_verified_vs_claimed_pair":null,"cap_flags":[],"verification_depth":{"sources_count":4,"primary_source_types":[]}}}