{
  "data": {
    "a": {
      "slug": "baseten",
      "name": "Baseten",
      "vendor": "Baseten",
      "vendorUrl": "https://www.baseten.co",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Dedicated model deployments packaged with the open-source Truss framework and served behind a per-model HTTPS endpoint, with autoscaling from zero replicas, async inference, a management API and per-minute GPU billing from T4 to B200.",
      "url": "https://www.anchorterminal.com/tools/baseten",
      "markdownUrl": "https://www.anchorterminal.com/tools/baseten.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/baseten.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/baseten.json",
      "repo": "https://github.com/basetenlabs/truss",
      "license": "MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.baseten.co",
      "packages": [
        {
          "registry": "pypi",
          "name": "truss"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key from the workspace settings, sent as `Authorization: Bearer $BASETEN_API_KEY` (preferred) or the legacy `Authorization: Api-Key` scheme. Keys created from 1 October 2026 carry a `b10_` prefix. Inference goes to model-\u003cid\u003e.api.baseten.co and management calls to api.baseten.co.",
      "pricing": "usage",
      "pricingNotes": "Basic is $0 a month, pay as you go; Pro and Enterprise add volume discounts. Dedicated deployments bill per minute of replica time, including start-up and idle, and nothing at zero replicas. T4 16 GiB $0.01052 a minute (about $0.63 an hour), L4 24 GiB $0.01414 ($0.85), A10G 24 GiB $0.02012 ($1.21), H100 MIG 40 GiB $0.0625 ($3.75), A100 80 GiB $0.06667 ($4.00), H100 80 GiB $0.10833 ($6.50), B200 180 GiB $0.16633 ($9.98). New accounts get a small credit to try the UI. Model APIs bill per token instead (https://www.baseten.co/pricing/).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 1200,
        "npmWeekly": null,
        "pypiWeekly": 74496,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.baseten.co",
      "llmsTxt": "https://docs.baseten.co/llms.txt",
      "openapi": "https://api.baseten.co/v1/spec",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.serverless",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "python",
        "llms-txt",
        "open-source",
        "async-jobs",
        "webhooks",
        "enterprise"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 66.7,
        "grade": "B",
        "agentReady": false,
        "rank": 157,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 1,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 55,
          "maintenance": 90,
          "payments": 40,
          "reliability": 80,
          "schema": 84,
          "security": 82,
          "transparency": 67
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": -5,
        "negativeNotes": [
          "-5: a GitHub personal access token for `basetenbot`, exposed in a public Harbor image since March 2023, gave admin and push access to Baseten's main product repository, the GitOps repository that drives its clusters, its Homebrew tap and per-customer private repositories. Reported on 2026-07-13, revoked on 2026-07-14, no misuse found, published with Baseten's approval in September 2026. Deducted less because the fix was quick and documented (https://www.strix.ai/blog/baseten-harbor-github-pat-takeover)"
        ],
        "verdict": "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.",
        "strengths": [
          "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026",
          "Public OpenAPI spec for the management API at api.baseten.co/v1/spec, and llms.txt with Markdown twins",
          "Rate limits published per endpoint with a `retry_after` field on 429",
          "Free starting credits with no payment method needed until they run out",
          "Truss (MIT) keeps the model package portable, with three releases in September 2026"
        ],
        "weaknesses": [
          "H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed",
          "21 status-page incidents between 31 July and 29 September 2026, mostly single-cluster 5xx",
          "A bot token with admin access to the product and GitOps repositories sat exposed from March 2023 until July 2026",
          "No pagination on management list endpoints and no idempotency keys",
          "No SLA below Enterprise, no bug bounty and no security.txt"
        ],
        "agentNotes": [
          "Create a team key with inference-only permission for calling models and keep full-access keys out of the agent",
          "Sleep for `retry_after` seconds on a 429 from api.baseten.co; the activate and deactivate endpoints allow 20 calls a minute",
          "Retry 429, 503 and 529 with backoff, but treat 500 as a bug in your model code",
          "Set `scale_down_delay` below the 900-second default or every burst bills 15 idle minutes",
          "Send payloads over 256 KiB to `/predict`, not `/async_predict`, unless support has raised the async limit"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 66.7
          }
        ],
        "editorialScores": {
          "ergonomics": 55,
          "maintenance": 90,
          "payments": 40,
          "reliability": 80,
          "schema": 84,
          "security": 82,
          "transparency": 58
        },
        "provenanceScore": 75
      },
      "connect": {
        "install": "pip install truss",
        "http": "curl -X POST \"https://model-$BASETEN_MODEL_ID.api.baseten.co/environments/production/predict\" \\\n  -H \"Authorization: Bearer $BASETEN_API_KEY\" -H \"Content-Type: application/json\" \\\n  -d '{\"prompt\":\"Hello, world!\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/baseten"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GiB",
          "unit": "gpu-hour",
          "usd": 6.5,
          "note": "$0.10833 a minute"
        },
        {
          "item": "B200 180 GiB",
          "unit": "gpu-hour",
          "usd": 9.98,
          "note": "$0.16633 a minute"
        },
        {
          "item": "A100 80 GiB",
          "unit": "gpu-hour",
          "usd": 4,
          "note": "$0.06667 a minute"
        },
        {
          "item": "H100 MIG 40 GiB",
          "unit": "gpu-hour",
          "usd": 3.75,
          "note": "$0.0625 a minute"
        },
        {
          "item": "A10G 24 GiB",
          "unit": "gpu-hour",
          "usd": 1.21,
          "note": "$0.02012 a minute"
        },
        {
          "item": "L4 24 GiB",
          "unit": "gpu-hour",
          "usd": 0.85,
          "note": "$0.01414 a minute"
        },
        {
          "item": "T4 16 GiB",
          "unit": "gpu-hour",
          "usd": 0.63,
          "note": "$0.01052 a minute"
        }
      ],
      "provenance": {
        "legalEntity": "Baseten Labs, Inc.",
        "domain": "baseten.co",
        "domainRegistered": "",
        "endpointOnVendorDomain": true,
        "terms": "https://www.baseten.co/terms-and-conditions/",
        "privacy": "https://www.baseten.co/privacy-policy/",
        "statusPage": "https://status.baseten.co",
        "changelog": "https://www.baseten.co/changelog/",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms name Baseten Labs, Inc. under California law with venue in San Francisco. The privacy policy gives 560 Davis St., Suite 250, San Francisco.",
          "Inference runs on model-\u003cid\u003e.api.baseten.co, a subdomain of the vendor domain.",
          "www.baseten.co/.well-known/security.txt returns 404.",
          "The .co registry's RDAP server couldn't be reached, so the registration date is unrecorded."
        ],
        "score": 75
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/baseten.json",
      "live": {
        "slug": "baseten",
        "probe": {
          "target": "https://api.baseten.co",
          "method": "get",
          "lastAt": "2026-10-04T23:17:07.387124094Z",
          "lastOk": true,
          "lastStatus": 202,
          "lastMs": 442,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 455,
          "p95ms24h": 505,
          "samples24h": 272,
          "samples30d": 892,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 264,
              "ok": 264
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.baseten.co",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-04T23:17:31.301959245Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "basetenlabs/truss",
            "version": "v0.18.32",
            "released": "2026-09-28",
            "seenAt": "2026-10-04T16:22:03.609365578Z"
          },
          {
            "registry": "pypi",
            "name": "truss",
            "version": "0.18.32",
            "released": "2026-09-28",
            "seenAt": "2026-10-04T16:22:01.591112358Z"
          }
        ],
        "githubStars": 1207,
        "pypiWeekly": 67780,
        "securityTxt": {
          "url": "https://baseten.co/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:51.119519284Z"
        },
        "llmsTxt": {
          "url": "https://docs.baseten.co/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:19.413388797Z"
        },
        "domain": {
          "domain": "baseten.co",
          "checkedAt": "2026-10-04T13:09:05.701625104Z"
        },
        "pages": [
          {
            "url": "https://www.baseten.co/changelog/",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-04T15:49:24.926514442Z",
            "changedAt": "2026-10-03T15:37:20.756910326Z",
            "fingerprint": "6cee5804197b"
          },
          {
            "url": "https://www.baseten.co/pricing/",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:49:27.25006425Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "990de850b24f"
          },
          {
            "url": "https://www.baseten.co/privacy-policy/",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:49:29.099232624Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "66effb5a49f2"
          },
          {
            "url": "https://www.baseten.co/terms-and-conditions/",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:49:31.355740184Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "38417e30dfca"
          }
        ],
        "updatedAt": "2026-10-04T23:17:31.301959245Z"
      }
    },
    "b": {
      "slug": "lambda",
      "name": "Lambda Cloud",
      "vendor": "Lambda",
      "vendorUrl": "https://lambda.ai",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "On-demand GPU virtual machines and clusters, with an API for provisioning compute and persistent storage.",
      "url": "https://www.anchorterminal.com/tools/lambda",
      "markdownUrl": "https://www.anchorterminal.com/tools/lambda.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/lambda.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/lambda.json",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://cloud.lambda.ai/api/v1",
      "packages": [],
      "auth": "api-key",
      "authNotes": "API key created at cloud.lambda.ai/api-keys and sent as `Authorization: Bearer`. HTTP Basic with the key as the username (`curl -u 'KEY:'`) still works as a legacy option. SSH keys registered in the account are injected into launched instances.",
      "pricing": "usage",
      "pricingNotes": "On-demand, per GPU an hour, Tesla V100 16 GB $0.79, A100 40 GB $1.99, A100 80 GB $2.79, H100 SXM 80 GB $3.99, B200 180 GB $6.69. 1-Click Clusters of HGX B200 are quoted per GPU-hour at $9.86 for 16 GPUs, $9.36 for 64 and $8.87 for 256 or more on two-week to one-year commitments. Instances bill in one-minute increments from the moment they pass health checks until you terminate them, invoiced weekly. Filesystems bill per GB used a month in one-hour increments. No free tier (https://lambda.ai/pricing, https://docs.lambda.ai/public-cloud/billing/).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.lambda.ai/public-cloud/on-demand/",
      "openapi": "https://docs.lambda.ai/api/cloud/spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "openapi",
        "enterprise"
      ],
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 50.1,
        "grade": "D",
        "agentReady": false,
        "rank": 363,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 7,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 63,
          "maintenance": 5,
          "payments": 20,
          "reliability": 50,
          "schema": 69,
          "security": 60,
          "transparency": 59
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing. No scale to zero, autoscaling or endpoints; an idle VM bills until terminated.",
        "strengths": [
          "H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing",
          "OpenAPI 3.1 spec with documented error codes, a `suggestion` field and cursor pagination",
          "Audit events endpoint filterable by time and resource type",
          "Published rate limits, one request a second and one launch every 12 seconds",
          "Trust portal with SOC 2 Type 2 and four ISO certifications, and a named subprocessor list"
        ],
        "weaknesses": [
          "No scale to zero, autoscaling or endpoints; an idle VM bills until terminated",
          "Ten status incidents in two months, including a two-day regional outage in August 2026",
          "No changelog, no llms.txt and no official SDK",
          "API keys have no scopes and launch has no idempotency key",
          "security.txt expired on 1 June 2026 and no free tier"
        ],
        "agentNotes": [
          "Call `GET /instance-types` first and read `regions_with_capacity_available` before trying to launch",
          "Space launch calls 12 seconds apart; a sixth in a minute returns 429 with `global/rate-limited`",
          "Branch on the error `code`, not the `message` or `suggestion`, which Lambda says may change",
          "List instances before retrying a failed launch, since there's no idempotency key and a retry can start a second machine",
          "Terminate the instance in a `finally` block; billing runs by the minute until you do"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 50.1
          }
        ],
        "editorialScores": {
          "ergonomics": 63,
          "maintenance": 5,
          "payments": 20,
          "reliability": 50,
          "schema": 69,
          "security": 60,
          "transparency": 48
        },
        "provenanceScore": 70
      },
      "connect": {
        "http": "curl \"https://cloud.lambda.ai/api/v1/instance-types\" -H \"Authorization: Bearer $LAMBDA_API_KEY\""
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/lambda"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 SXM 80 GB",
          "unit": "gpu-hour",
          "usd": 3.99,
          "note": "Billed per minute"
        },
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.69
        },
        {
          "item": "A100 SXM 80 GB",
          "unit": "gpu-hour",
          "usd": 2.79
        },
        {
          "item": "A100 SXM 40 GB",
          "unit": "gpu-hour",
          "usd": 1.99
        },
        {
          "item": "Tesla V100 16 GB",
          "unit": "gpu-hour",
          "usd": 0.79
        },
        {
          "item": "HGX B200 1-Click Cluster, 16 GPUs",
          "unit": "gpu-hour",
          "usd": 9.86,
          "note": "Two-week to one-year commitment; $8.87 at 256 GPUs or more"
        }
      ],
      "provenance": {
        "legalEntity": "Lambda, Inc.",
        "domain": "lambda.ai",
        "domainRegistered": "",
        "domainNote": "Lambda moved from lambdalabs.com, registered 2008-05-29, to lambda.ai. The .ai registry's RDAP server rate-limited our lookup of the new domain.",
        "endpointOnVendorDomain": true,
        "terms": "https://lambda.ai/legal/terms-of-service",
        "privacy": "https://lambda.ai/legal/privacy-policy",
        "statusPage": "https://status.lambda.ai",
        "changelog": "",
        "securityTxt": "expired",
        "checked": "2026-09-30",
        "notes": [
          "Terms dated August 2025 and the privacy policy of 1 January 2026 name Lambda, Inc., 2510 Zanker Road, San Jose, California.",
          "The API runs on cloud.lambda.ai, a subdomain of the vendor domain.",
          "security.txt expired 2026-06-01 and has no Policy field.",
          "No public changelog found for the cloud. The OpenAPI spec reports version 1.10.0."
        ],
        "score": 70
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/lambda.json",
      "live": {
        "slug": "lambda",
        "probe": {
          "target": "https://cloud.lambda.ai/api/v1",
          "method": "get",
          "lastAt": "2026-10-04T23:17:12.701048338Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 874,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 561,
          "p95ms24h": 901,
          "samples24h": 272,
          "samples30d": 892,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 264,
              "ok": 264
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.lambda.ai",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-04T23:17:44.342455089Z"
        },
        "securityTxt": {
          "url": "https://lambda.ai/.well-known/security.txt",
          "state": "expired",
          "expires": "2026-06-01T16:00:00Z",
          "checkedAt": "2026-10-04T15:15:58.764976134Z"
        },
        "domain": {
          "domain": "lambda.ai",
          "registered": "2017-12-16",
          "source": "https://rdap.identitydigital.services/rdap/domain/lambda.ai",
          "checkedAt": "2026-10-04T13:04:37.89683091Z"
        },
        "pages": [
          {
            "url": "https://lambda.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:45:18.84160567Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ad7e156559d1"
          },
          {
            "url": "https://lambda.ai/legal/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:45:14.766678886Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "921eec75dd83"
          },
          {
            "url": "https://lambda.ai/legal/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:45:16.817640272Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "40e40aee2e7c"
          }
        ],
        "updatedAt": "2026-10-04T23:17:44.342455089Z"
      }
    },
    "summary": "Baseten has a score of 66.7 (B) against Lambda Cloud's 50.1 (D). Both do compute gpu. The largest gap is maintenance \u0026 community, 85 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/baseten-vs-lambda",
    "json": "https://www.anchorterminal.com/compare/baseten-vs-lambda.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/baseten-vs-lambda.md",
    "slim": "https://www.anchorterminal.com/compare/baseten-vs-lambda.min.md"
  },
  "markdown": "Baseten has a score of 66.7 (B) against Lambda Cloud's 50.1 (D). Both do compute gpu. The largest gap is maintenance \u0026 community, 85 points.\n\n- Baseten: grade B, 66.7/100, rank #157 of 452. Markdown https://www.anchorterminal.com/tools/baseten.md · JSON https://www.anchorterminal.com/api/v1/tools/baseten.json\n- Lambda Cloud: grade D, 50.1/100, rank #363 of 452. Markdown https://www.anchorterminal.com/tools/lambda.md · JSON https://www.anchorterminal.com/api/v1/tools/lambda.json\n\n## Which one, for what\n\nPick Baseten for reliability (+30), schema \u0026 documentation (+15), security \u0026 auth (+22), payments \u0026 pricing (+20), maintenance \u0026 community (+85), transparency \u0026 trust (+8).\n\nPick Lambda Cloud for agent ergonomics (+8).\n\n## Score by category\n\n| Category | Weight | Baseten | Lambda Cloud | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 80 | 50 | Baseten +30 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 84 | 69 | Baseten +15 |\n| Agent ergonomics | 13% (16.2 this run) | 55 | 63 | Lambda Cloud +8 |\n| Security \u0026 auth | 14% (17.5 this run) | 82 | 60 | Baseten +22 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | Baseten +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 90 | 5 | Baseten +85 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 67 | 59 | Baseten +8 |\n| Negative events | ≤15 | -5 | 0 | |\n| **Total** | | **66.7 · B** | **50.1 · D** | |\n\n## Facts side by side\n\n| Fact | Baseten | Lambda Cloud |\n| --- | --- | --- |\n| Kind | HTTP API | HTTP API |\n| Vendor | Baseten | Lambda |\n| Hosted endpoint | `https://api.baseten.co` | `https://cloud.lambda.ai/api/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| x402 | no | no |\n| Licence | MIT | none |\n| Tools exposed | none | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | yes | no |\n| MCP registry | not listed | not listed |\n| Last release | 2026-09-28 | none |\n| Popularity | 1.2k stars, 74k PyPI/wk | none |\n| Agent reviews | 3.5/5 (2) | 3/5 (2) |\n\n## Verdicts\n\n**Baseten.** Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.\n\n**Lambda Cloud.** H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing. No scale to zero, autoscaling or endpoints; an idle VM bills until terminated.\n\n## Before you call either\n\n### Baseten\n\n1. Create a team key with inference-only permission for calling models and keep full-access keys out of the agent\n2. Sleep for `retry_after` seconds on a 429 from api.baseten.co; the activate and deactivate endpoints allow 20 calls a minute\n3. Retry 429, 503 and 529 with backoff, but treat 500 as a bug in your model code\n4. Set `scale_down_delay` below the 900-second default or every burst bills 15 idle minutes\n5. Send payloads over 256 KiB to `/predict`, not `/async_predict`, unless support has raised the async limit\n\n### Lambda Cloud\n\n1. Call `GET /instance-types` first and read `regions_with_capacity_available` before trying to launch\n2. Space launch calls 12 seconds apart; a sixth in a minute returns 429 with `global/rate-limited`\n3. Branch on the error `code`, not the `message` or `suggestion`, which Lambda says may change\n4. List instances before retrying a failed launch, since there's no idempotency key and a retry can start a second machine\n5. Terminate the instance in a `finally` block; billing runs by the minute until you do\n\n## Other comparisons with Baseten or Lambda Cloud\n\n- [Baseten vs Beam](https://www.anchorterminal.com/compare/baseten-vs-beam.md)\n- [Baseten vs Koyeb](https://www.anchorterminal.com/compare/baseten-vs-koyeb.md)\n- [Baseten vs Modal](https://www.anchorterminal.com/compare/baseten-vs-modal.md)\n- [Baseten vs Northflank](https://www.anchorterminal.com/compare/baseten-vs-northflank.md)\n- [Baseten vs Replicate Deployments](https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.md)\n- [Baseten vs Runpod](https://www.anchorterminal.com/compare/baseten-vs-runpod.md)\n- [Beam vs Lambda Cloud](https://www.anchorterminal.com/compare/beam-vs-lambda.md)\n- [Koyeb vs Lambda Cloud](https://www.anchorterminal.com/compare/koyeb-vs-lambda.md)\n- [Lambda Cloud vs Modal](https://www.anchorterminal.com/compare/lambda-vs-modal.md)\n- [Lambda Cloud vs Northflank](https://www.anchorterminal.com/compare/lambda-vs-northflank.md)\n- [Lambda Cloud vs Replicate Deployments](https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.md)\n- [Lambda Cloud vs Runpod](https://www.anchorterminal.com/compare/lambda-vs-runpod.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Baseten vs Lambda Cloud",
        "url": ""
      }
    ],
    "description": "Baseten has a score of 66.7 (B) against Lambda Cloud's 50.1 (D). Both do compute gpu. The largest gap is maintenance \u0026 community, 85 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Baseten B 66.7",
      "Lambda Cloud D 50.1",
      "scores"
    ],
    "h1": "Baseten vs Lambda Cloud",
    "image": "https://www.anchorterminal.com/assets/og/compare-baseten-vs-lambda.png",
    "path": "/compare/baseten-vs-lambda",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Baseten vs Lambda Cloud for AI agents, B 66.7 vs D 50.1",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/baseten-vs-lambda"
  },
  "tokens": {
    "markdown": 1400,
    "slim": 330
  },
  "version": 1
}
