{
  "data": {
    "a": {
      "slug": "baseten",
      "name": "Baseten",
      "vendor": "Baseten",
      "vendorUrl": "https://www.baseten.co",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Dedicated model deployments packaged with the open-source Truss framework and served behind a per-model HTTPS endpoint, with autoscaling from zero replicas, async inference, a management API and per-minute GPU billing from T4 to B200.",
      "url": "https://www.anchorterminal.com/tools/baseten",
      "markdownUrl": "https://www.anchorterminal.com/tools/baseten.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/baseten.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/baseten.json",
      "repo": "https://github.com/basetenlabs/truss",
      "license": "MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.baseten.co",
      "packages": [
        {
          "registry": "pypi",
          "name": "truss"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key from the workspace settings, sent as `Authorization: Bearer $BASETEN_API_KEY` (preferred) or the legacy `Authorization: Api-Key` scheme. Keys created from 1 October 2026 carry a `b10_` prefix. Inference goes to model-\u003cid\u003e.api.baseten.co and management calls to api.baseten.co.",
      "pricing": "usage",
      "pricingNotes": "Basic is $0 a month, pay as you go; Pro and Enterprise add volume discounts. Dedicated deployments bill per minute of replica time, including start-up and idle, and nothing at zero replicas. T4 16 GiB $0.01052 a minute (about $0.63 an hour), L4 24 GiB $0.01414 ($0.85), A10G 24 GiB $0.02012 ($1.21), H100 MIG 40 GiB $0.0625 ($3.75), A100 80 GiB $0.06667 ($4.00), H100 80 GiB $0.10833 ($6.50), B200 180 GiB $0.16633 ($9.98). New accounts get a small credit to try the UI. Model APIs bill per token instead (https://www.baseten.co/pricing/).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 1200,
        "npmWeekly": null,
        "pypiWeekly": 74496,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.baseten.co",
      "llmsTxt": "https://docs.baseten.co/llms.txt",
      "openapi": "https://api.baseten.co/v1/spec",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.serverless",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "python",
        "llms-txt",
        "open-source",
        "async-jobs",
        "webhooks",
        "enterprise"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 66.5,
        "grade": "B",
        "agentReady": false,
        "rank": 264,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 2,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 55,
          "maintenance": 90,
          "payments": 40,
          "reliability": 80,
          "schema": 84,
          "security": 82,
          "transparency": 65
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": -5,
        "negativeNotes": [
          "-5: a GitHub personal access token for `basetenbot`, exposed in a public Harbor image since March 2023, gave admin and push access to Baseten's main product repository, the GitOps repository that drives its clusters, its Homebrew tap and per-customer private repositories. Reported on 2026-07-13, revoked on 2026-07-14, no misuse found, published with Baseten's approval in September 2026. Deducted less because the fix was quick and documented (https://www.strix.ai/blog/baseten-harbor-github-pat-takeover)"
        ],
        "verdict": "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.",
        "bestFor": "Teams that want one model behind a production endpoint with real autoscaling knobs, environments and scoped keys.",
        "strengths": [
          "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026",
          "Public OpenAPI spec for the management API at api.baseten.co/v1/spec, and llms.txt with Markdown twins",
          "Rate limits published per endpoint with a `retry_after` field on 429",
          "Free starting credits with no payment method needed until they run out",
          "Truss (MIT) keeps the model package portable, with three releases in September 2026"
        ],
        "weaknesses": [
          "H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed",
          "21 status-page incidents between 31 July and 29 September 2026, mostly single-cluster 5xx",
          "A bot token with admin access to the product and GitOps repositories sat exposed from March 2023 until July 2026",
          "No pagination on management list endpoints and no idempotency keys",
          "No SLA below Enterprise, no bug bounty and no security.txt"
        ],
        "agentNotes": [
          "Create a team key with inference-only permission for calling models and keep full-access keys out of the agent",
          "Sleep for `retry_after` seconds on a 429 from api.baseten.co; the activate and deactivate endpoints allow 20 calls a minute",
          "Retry 429, 503 and 529 with backoff, but treat 500 as a bug in your model code",
          "Set `scale_down_delay` below the 900-second default or every burst bills 15 idle minutes",
          "Send payloads over 256 KiB to `/predict`, not `/async_predict`, unless support has raised the async limit"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 66.5
          }
        ],
        "editorialScores": {
          "ergonomics": 55,
          "maintenance": 90,
          "payments": 40,
          "reliability": 80,
          "schema": 84,
          "security": 82,
          "transparency": 58
        },
        "provenanceScore": 71
      },
      "connect": {
        "install": "pip install truss",
        "http": "curl -X POST \"https://model-$BASETEN_MODEL_ID.api.baseten.co/environments/production/predict\" \\\n  -H \"Authorization: Bearer $BASETEN_API_KEY\" -H \"Content-Type: application/json\" \\\n  -d '{\"prompt\":\"Hello, world!\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/baseten"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GiB",
          "unit": "gpu-hour",
          "usd": 6.5,
          "note": "$0.10833 a minute"
        },
        {
          "item": "B200 180 GiB",
          "unit": "gpu-hour",
          "usd": 9.98,
          "note": "$0.16633 a minute"
        },
        {
          "item": "A100 80 GiB",
          "unit": "gpu-hour",
          "usd": 4,
          "note": "$0.06667 a minute"
        },
        {
          "item": "H100 MIG 40 GiB",
          "unit": "gpu-hour",
          "usd": 3.75,
          "note": "$0.0625 a minute"
        },
        {
          "item": "A10G 24 GiB",
          "unit": "gpu-hour",
          "usd": 1.21,
          "note": "$0.02012 a minute"
        },
        {
          "item": "L4 24 GiB",
          "unit": "gpu-hour",
          "usd": 0.85,
          "note": "$0.01414 a minute"
        },
        {
          "item": "T4 16 GiB",
          "unit": "gpu-hour",
          "usd": 0.63,
          "note": "$0.01052 a minute"
        }
      ],
      "provenance": {
        "legalEntity": "Baseten Labs, Inc.",
        "domain": "baseten.co",
        "domainRegistered": "",
        "endpointOnVendorDomain": true,
        "terms": "https://www.baseten.co/terms-and-conditions/",
        "privacy": "https://www.baseten.co/privacy-policy/",
        "statusPage": "https://status.baseten.co",
        "changelog": "https://www.baseten.co/changelog/",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms name Baseten Labs, Inc. under California law with venue in San Francisco. The privacy policy gives 560 Davis St., Suite 250, San Francisco.",
          "Inference runs on model-\u003cid\u003e.api.baseten.co, a subdomain of the vendor domain.",
          "www.baseten.co/.well-known/security.txt returns 404.",
          "The .co registry's RDAP server couldn't be reached, so the registration date is unrecorded."
        ],
        "score": 71
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/baseten.json",
      "live": {
        "slug": "baseten",
        "probe": {
          "target": "https://api.baseten.co",
          "method": "get",
          "lastAt": "2026-10-09T10:42:37.475251192Z",
          "lastOk": true,
          "lastStatus": 202,
          "lastMs": 443,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 459,
          "p95ms24h": 490,
          "samples24h": 260,
          "samples30d": 2098,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 114,
              "ok": 114
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.baseten.co",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T10:41:26.112736642Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "basetenlabs/truss",
            "version": "v0.18.33",
            "released": "2026-10-06",
            "seenAt": "2026-10-08T16:02:10.626982738Z"
          },
          {
            "registry": "pypi",
            "name": "truss",
            "version": "0.18.33",
            "released": "2026-10-06",
            "seenAt": "2026-10-08T16:02:06.626291704Z"
          }
        ],
        "githubStars": 1214,
        "pypiWeekly": 62737,
        "securityTxt": {
          "url": "https://baseten.co/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:38:37.892521999Z"
        },
        "llmsTxt": {
          "url": "https://docs.baseten.co/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:06.558742297Z"
        },
        "domain": {
          "domain": "baseten.co",
          "checkedAt": "2026-10-04T13:09:05.701625104Z"
        },
        "pages": [
          {
            "url": "https://www.baseten.co/changelog/",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:31.551352111Z",
            "changedAt": "2026-10-06T16:14:05.469135663Z",
            "fingerprint": "fcc7d4bab97d"
          },
          {
            "url": "https://www.baseten.co/pricing/",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:33.809519828Z",
            "changedAt": "2026-10-06T16:14:07.823033015Z",
            "fingerprint": "e41ad7119a6b"
          },
          {
            "url": "https://www.baseten.co/privacy-policy/",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:35.678713031Z",
            "changedAt": "2026-10-06T16:14:09.752979771Z",
            "fingerprint": "09df4e188f65"
          },
          {
            "url": "https://www.baseten.co/terms-and-conditions/",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:38.014111749Z",
            "changedAt": "2026-10-06T16:14:11.853258349Z",
            "fingerprint": "a8b301619021"
          }
        ],
        "updatedAt": "2026-10-09T10:42:37.475251192Z"
      }
    },
    "answer": "Nebius AI Cloud and Baseten score within a point of each other on agent readiness, 67.2 (B) and 66.5 (B). Baseten leads on reliability, schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community.",
    "b": {
      "slug": "nebius-ai-cloud",
      "name": "Nebius AI Cloud",
      "vendor": "Nebius",
      "vendorUrl": "https://nebius.com",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Nebius AI Cloud rents NVIDIA GPU virtual machines and InfiniBand clusters, with managed Kubernetes, Slurm and Serverless AI jobs and endpoints for containers. Resources are managed through REST and gRPC APIs, a CLI, a Terraform provider and SDKs.",
      "url": "https://www.anchorterminal.com/tools/nebius-ai-cloud",
      "markdownUrl": "https://www.anchorterminal.com/tools/nebius-ai-cloud.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/nebius-ai-cloud.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/nebius-ai-cloud.json",
      "repo": "https://github.com/nebius/api",
      "license": "Proprietary service under the Nebius Services Agreement. The API definitions, the Go, Python and JavaScript SDKs and the MCP server on GitHub are MIT",
      "transports": [
        "http",
        "stdio"
      ],
      "remoteUrl": "https://api.nebius.cloud",
      "packages": [
        {
          "registry": "pypi",
          "name": "nebius"
        },
        {
          "registry": "npm",
          "name": "@nebius/js-sdk"
        },
        {
          "registry": "go",
          "name": "github.com/nebius/gosdk"
        }
      ],
      "auth": "mixed",
      "authNotes": "Self-serve. A person signs up in the web console with a Google, GitHub or Microsoft account. Every API call takes `Authorization: Bearer` with an access token valid for 12 hours. A user gets one from `nebius iam get-access-token`. A service account uploads an RSA public key (an authorised key, with optional expiry), signs a five-minute RS256 JWT and exchanges it at `https://auth.eu.nebius.com/oauth2/token/exchange`. Permissions come from group roles (`auditor`, `viewer`, `editor`, `admin` and service roles) granted on a tenant, project or resource. Serverless AI endpoints take their own token set in `spec.authToken`.",
      "pricing": "usage",
      "pricingNotes": "Pay as you go, billed by the second, with no free tier or trial found. On-demand per GPU-hour from 1 October 2026, B300 $9.50, B200 $8.50, H200 $5.40, H100 $4.50, RTX PRO 6000 $1.80, and L40S $1.35 plus vCPU and RAM. Preemptible GPUs are spot priced from $0.79. Serverless AI bills at Compute prices and a stopped endpoint bills nothing. Adding a card charges $25 to the balance. Commitment discounts go through sales (https://docs.nebius.com/compute/resources/pricing, https://nebius.com/prices).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs index, the OpenAPI document or the price list (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 4137,
        "pypiWeekly": 468852,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.nebius.com/",
      "llmsTxt": "https://docs.nebius.com/llms.txt",
      "openapi": "https://api.nebius.cloud/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "openapi",
        "llms-txt",
        "grpc",
        "terraform",
        "mcp",
        "python",
        "typescript",
        "go",
        "status-page",
        "soc2",
        "enterprise"
      ],
      "lastRelease": "2026-10-07",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 67.2,
        "grade": "B",
        "agentReady": false,
        "rank": 240,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 1,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 82,
          "payments": 20,
          "reliability": 57,
          "schema": 78,
          "security": 81,
          "transparency": 77
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "One API definition generates the REST and gRPC interfaces, the CLI, Terraform provider and three SDKs, with a 602-operation OpenAPI document, `X-Idempotency-Key` and role-scoped service accounts. The status page lists 14 major incidents between 14 July and 8 October 2026, no request rate limits were found, and signup needs a browser and a card.",
        "bestFor": "Teams that want whole GPU VMs or InfiniBand clusters in Europe, the UK, Israel or the US with IAM, Terraform and an SLA, and are content to manage endpoint lifecycles themselves.",
        "strengths": [
          "OpenAPI 3.0.3 document at `https://api.nebius.cloud/openapi.json` with 602 operations, generated from the same protobuf definitions as the gRPC API, CLI, Terraform provider and SDKs",
          "`X-Idempotency-Key` header for modifying calls, and a `retry_type` field on errors that says whether to retry the call",
          "Service accounts sign in with an uploaded RSA key and receive 12-hour tokens, with roles granted per tenant, project or resource",
          "Per-second billing with public prices, and a 99.5 per cent monthly uptime commitment per virtual machine",
          "llms.txt, every docs page as Markdown, and a keyless docs MCP server at `https://docs.nebius.com/mcp`"
        ],
        "weaknesses": [
          "Status page lists 14 incidents marked major between 14 July and 8 October 2026, including about 21 hours of partial degradation in us-central1 on 19 August",
          "No request rate limits with numbers and no Retry-After guidance found in the reviewed documentation",
          "Serverless AI endpoints run on one container VM that is started and stopped by hand; no autoscaling or scale to zero found",
          "No free tier or trial found. Adding a card at signup charges $25 to the balance, and signup is a browser flow through Google, GitHub or Microsoft",
          "The OpenAPI document has no examples, documents only 200 responses and reports its version as `version not set`"
        ],
        "agentNotes": [
          "Use a service account with an authorised key, then exchange a five-minute RS256 JWT at `https://auth.eu.nebius.com/oauth2/token/exchange` for a 12-hour Bearer token",
          "Send `X-Idempotency-Key` with a random UUID on every create, update and delete, since a 504 can follow a call that succeeded",
          "Poll the returned operation (`/ai/v1/endpoints/operations/{id}`) until `status` is set; concurrent operations on one resource are not supported",
          "Stop or delete endpoints when idle. A stopped endpoint bills nothing, a stopped Devlab or VM still bills for its disk",
          "Check region support first. Serverless AI is absent from `eu-south1` and `us-north1`, and each GPU platform exists in one to four regions",
          "Run the beta `nebius/mcp-server` with safe mode on (the default); `nebius_cli_execute` can run any CLI command when `SAFE_MODE=false`"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 67.2
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 82,
          "payments": 20,
          "reliability": 57,
          "schema": 78,
          "security": 81,
          "transparency": 70
        },
        "provenanceScore": 83
      },
      "connect": {
        "install": "curl -sSL https://artifacts.nebius.cloud/cli/install.sh | bash",
        "http": "curl --request GET --url 'https://api.nebius.cloud/iam/v1/profiles' --header 'Authorization: Bearer \u003caccess_token\u003e'",
        "config": {
          "mcpServers": {
            "Nebius MCP Server": {
              "args": [
                "--refresh-package",
                "nebius-mcp-server",
                "nebius-mcp-server@git+https://github.com/nebius/mcp-server@main"
              ],
              "command": "uvx",
              "env": {}
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/nebius-ai-cloud"
      },
      "sameCompany": [
        "nebius-token-factory-fine-tuning"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "NVIDIA H200 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 5.4,
          "note": "Billed per second. $4.50 before 1 October 2026"
        },
        {
          "item": "NVIDIA H100 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 4.5,
          "note": "$3.85 before 1 October 2026"
        },
        {
          "item": "NVIDIA B200 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 8.5
        },
        {
          "item": "NVIDIA B300 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 9.5
        },
        {
          "item": "NVIDIA RTX PRO 6000, on demand",
          "unit": "gpu-hour",
          "usd": 1.8
        },
        {
          "item": "NVIDIA L40S, GPU only",
          "unit": "gpu-hour",
          "usd": 1.35,
          "note": "vCPU ($0.01 to $0.012 an hour) and RAM ($0.0032 a GiB-hour) are billed separately"
        },
        {
          "item": "NVIDIA H200 NVLink, preemptible",
          "unit": "gpu-hour",
          "usd": 0.79,
          "note": "Minimum spot price from 8 October 2026. The spot price can change every 15 minutes"
        },
        {
          "item": "Network SSD disk",
          "unit": "gb-month",
          "usd": 0.071,
          "note": "Per GiB for 730 hours"
        }
      ],
      "provenance": {
        "legalEntity": "Nebius B.V.",
        "domain": "nebius.com",
        "domainRegistered": "2004-06-26",
        "endpointOnVendorDomain": false,
        "terms": "https://docs.nebius.com/legal/agreement",
        "privacy": "https://docs.nebius.com/legal/privacy",
        "statusPage": "https://status.nebius.com",
        "changelog": "https://docs.nebius.com/cli/release-notes",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "notes": [
          "The Services Agreement (published 15 September 2026, effective 28 September 2026) names Nebius B.V. under Dutch law for customers outside the United States and Israel, with other Nebius entities for those two countries. The separate Terms of Use page covers the website.",
          "The privacy policy (23 September 2026) gives Nebius B.V., Burgerweeshuispad 101, 1076ER Amsterdam, and says data processed for customers as a processor falls under the DPA at https://docs.nebius.com/legal/dpa.",
          "The API answers at api.nebius.cloud and tokens are exchanged at auth.eu.nebius.com. nebius.cloud is a second domain that Nebius's docs name for the API and the CLI installer.",
          "security.txt at nebius.com gives security@nebius.com and expires 2027-12-31. It has no Policy field.",
          "The changelog link is the CLI release notes, which are generated from the API. No separate API changelog was found.",
          "RDAP for nebius.com gives a registration date of 2004-06-26."
        ],
        "score": 83
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/nebius-ai-cloud.json",
      "live": {
        "slug": "nebius-ai-cloud",
        "probe": {
          "target": "https://api.nebius.cloud",
          "method": "get",
          "lastAt": "2026-10-09T10:42:50.420638285Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 228,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 164,
          "p95ms24h": 237,
          "samples24h": 33,
          "samples30d": 33,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 33,
              "ok": 33
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.nebius.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T10:41:49.609986027Z"
        },
        "updatedAt": "2026-10-09T10:42:50.420638285Z"
      }
    },
    "facts": [
      {
        "a": "HTTP API",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Baseten",
        "b": "Nebius",
        "name": "Vendor"
      },
      {
        "a": "https://api.baseten.co",
        "b": "https://api.nebius.cloud",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP, stdio",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "OAuth or key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "not published",
        "b": "$1.35 per GPU-hour",
        "name": "Price for compute gpu"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT",
        "b": "Proprietary service under the Nebius Services Agreement. The API definitions, the Go, Python and JavaScript SDKs and the MCP server on GitHub are MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-28",
        "b": "2026-10-07",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "2026-09-28",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2026-09-23",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "1.2k stars, 74k PyPI/wk",
        "b": "4.1k npm/wk, 469k PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "3.5/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Nebius AI Cloud and Baseten score within a point of each other on agent readiness, 67.2 (B) and 66.5 (B). Baseten leads on reliability, schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community.",
        "question": "Which is better for AI agents, Baseten or Nebius AI Cloud?"
      },
      {
        "answer": "Baseten needs an API key. Nebius AI Cloud takes an API key or an OAuth sign-in.",
        "question": "Do Baseten and Nebius AI Cloud need an API key?"
      },
      {
        "answer": "Yes. Baseten has a hosted endpoint at https://api.baseten.co and Nebius AI Cloud at https://api.nebius.cloud.",
        "question": "Can an agent call Baseten and Nebius AI Cloud without installing anything?"
      },
      {
        "answer": "Baseten is open source (MIT). No open-source release is listed for Nebius AI Cloud.",
        "question": "Are Baseten and Nebius AI Cloud open source?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 80 against 57",
          "Schema \u0026 documentation, 84 against 78",
          "Payments \u0026 pricing, 40 against 20",
          "Maintenance \u0026 community, 90 against 82"
        ],
        "also": [
          "Open source"
        ],
        "goodFor": "Teams that want one model behind a production endpoint with real autoscaling knobs, environments and scoped keys.",
        "slug": "baseten",
        "watchFor": "H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed"
      },
      {
        "aheadOn": [
          "Agent ergonomics, 77 against 55",
          "Transparency \u0026 trust, 77 against 65"
        ],
        "also": [
          "Runs on your own machine",
          "No incidents deducted, where Baseten loses 5 points for them"
        ],
        "goodFor": "Teams that want whole GPU VMs or InfiniBand clusters in Europe, the UK, Israel or the US with IAM, Terraform and an SLA, and are content to manage endpoint lifecycles themselves.",
        "slug": "nebius-ai-cloud",
        "watchFor": "Status page lists 14 incidents marked major between 14 July and 8 October 2026, including about 21 hours of partial degradation in us-central1 on 19 August"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-beam.json",
        "title": "Baseten vs Beam",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-beam"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
        "title": "Baseten vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-coreweave.json",
        "title": "Baseten vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-hugging-face-inference-endpoints.json",
        "title": "Baseten vs Hugging Face Inference Endpoints",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-hugging-face-inference-endpoints"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-hyperbolic.json",
        "title": "Baseten vs Hyperbolic",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-hyperbolic"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-koyeb.json",
        "title": "Baseten vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-lambda.json",
        "title": "Baseten vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-modal.json",
        "title": "Baseten vs Modal",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-northflank.json",
        "title": "Baseten vs Northflank",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.json",
        "title": "Baseten vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-runpod.json",
        "title": "Baseten vs Runpod",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-thunder-compute.json",
        "title": "Baseten vs Thunder Compute",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-thunder-compute"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-vast-ai.json",
        "title": "Baseten vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-verda.json",
        "title": "Baseten vs Verda",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-verda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-nebius-ai-cloud.json",
        "title": "Beam vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/beam-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-nebius-ai-cloud.json",
        "title": "Cerebrium vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-nebius-ai-cloud.json",
        "title": "CoreWeave vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-nebius-ai-cloud.json",
        "title": "Hugging Face Inference Endpoints vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/hyperbolic-vs-nebius-ai-cloud.json",
        "title": "Hyperbolic vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/hyperbolic-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-nebius-ai-cloud.json",
        "title": "Koyeb vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-nebius-ai-cloud.json",
        "title": "Lambda Cloud vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-nebius-ai-cloud.json",
        "title": "Modal vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/modal-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-northflank.json",
        "title": "Nebius AI Cloud vs Northflank",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy.json",
        "title": "Nebius AI Cloud vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-runpod.json",
        "title": "Nebius AI Cloud vs Runpod",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-thunder-compute.json",
        "title": "Nebius AI Cloud vs Thunder Compute",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-thunder-compute"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-vast-ai.json",
        "title": "Nebius AI Cloud vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-verda.json",
        "title": "Nebius AI Cloud vs Verda",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-verda"
      }
    ],
    "scores": [
      {
        "baseten": 80,
        "by": 23,
        "edge": "baseten",
        "key": "reliability",
        "name": "Reliability",
        "nebius-ai-cloud": 57,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "baseten": 84,
        "by": 6,
        "edge": "baseten",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "nebius-ai-cloud": 78,
        "weight": 13
      },
      {
        "baseten": 55,
        "by": 22,
        "edge": "nebius-ai-cloud",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "nebius-ai-cloud": 77,
        "weight": 13
      },
      {
        "baseten": 82,
        "by": 1,
        "edge": "baseten",
        "key": "security",
        "name": "Security \u0026 auth",
        "nebius-ai-cloud": 81,
        "weight": 14
      },
      {
        "baseten": 40,
        "by": 20,
        "edge": "baseten",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "nebius-ai-cloud": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "baseten": 90,
        "by": 8,
        "edge": "baseten",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "nebius-ai-cloud": 82,
        "weight": 7
      },
      {
        "baseten": 65,
        "by": 12,
        "edge": "nebius-ai-cloud",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "nebius-ai-cloud": 77,
        "weight": 7
      }
    ],
    "summary": "Nebius AI Cloud and Baseten score within a point of each other on agent readiness, 67.2 (B) and 66.5 (B). Baseten leads on reliability, schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community. Both do compute gpu.",
    "verdicts": {
      "baseten": "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.",
      "nebius-ai-cloud": "One API definition generates the REST and gRPC interfaces, the CLI, Terraform provider and three SDKs, with a 602-operation OpenAPI document, `X-Idempotency-Key` and role-scoped service accounts. The status page lists 14 major incidents between 14 July and 8 October 2026, no request rate limits were found, and signup needs a browser and a card."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud",
    "json": "https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud.md",
    "slim": "https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud.min.md"
  },
  "markdown": "Nebius AI Cloud and Baseten score within a point of each other on agent readiness, 67.2 (B) and 66.5 (B). Baseten leads on reliability, schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community. Both do compute gpu.\n\n- Baseten: grade B, 66.5/100, rank #264 of 842. Markdown https://www.anchorterminal.com/tools/baseten.md · JSON https://www.anchorterminal.com/api/v1/tools/baseten.json\n- Nebius AI Cloud: grade B, 67.2/100, rank #240 of 842. Markdown https://www.anchorterminal.com/tools/nebius-ai-cloud.md · JSON https://www.anchorterminal.com/api/v1/tools/nebius-ai-cloud.json\n\n## Which one, for what\n\n### Baseten (B)\n\nGood for: Teams that want one model behind a production endpoint with real autoscaling knobs, environments and scoped keys.\n\nAhead on:\n- Reliability, 80 against 57\n- Schema \u0026 documentation, 84 against 78\n- Payments \u0026 pricing, 40 against 20\n- Maintenance \u0026 community, 90 against 82\n\nAlso in its favour:\n- Open source\n\nWatch for: H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed\n\n### Nebius AI Cloud (B)\n\nGood for: Teams that want whole GPU VMs or InfiniBand clusters in Europe, the UK, Israel or the US with IAM, Terraform and an SLA, and are content to manage endpoint lifecycles themselves.\n\nAhead on:\n- Agent ergonomics, 77 against 55\n- Transparency \u0026 trust, 77 against 65\n\nAlso in its favour:\n- Runs on your own machine\n- No incidents deducted, where Baseten loses 5 points for them\n\nWatch for: Status page lists 14 incidents marked major between 14 July and 8 October 2026, including about 21 hours of partial degradation in us-central1 on 19 August\n\n\n## Score by category\n\n| Category | Weight | Baseten | Nebius AI Cloud | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 80 | 57 | Baseten +23 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 84 | 78 | Baseten +6 |\n| Agent ergonomics | 13% (16.2 this run) | 55 | 77 | Nebius AI Cloud +22 |\n| Security \u0026 auth | 14% (17.5 this run) | 82 | 81 | Baseten +1 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | Baseten +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 90 | 82 | Baseten +8 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 65 | 77 | Nebius AI Cloud +12 |\n| Negative events | ≤15 | -5 | 0 | |\n| **Total** | | **66.5 · B** | **67.2 · B** | |\n\n## Facts side by side\n\n| Fact | Baseten | Nebius AI Cloud |\n| --- | --- | --- |\n| Kind | HTTP API | HTTP API |\n| Vendor | Baseten | Nebius |\n| Hosted endpoint | `https://api.baseten.co` | `https://api.nebius.cloud` |\n| Transports | HTTP | HTTP, stdio |\n| Auth | API key | OAuth or key |\n| Pricing | Pay per use | Pay per use |\n| Price for compute gpu | not published | $1.35 per GPU-hour |\n| x402 | no | no |\n| Licence | MIT | Proprietary service under the Nebius Services Agreement. The API definitions, the Go, Python and JavaScript SDKs and the MCP server on GitHub are MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-28 | 2026-10-07 |\n| Terms last updated | no date given | 2026-09-28 |\n| Privacy policy last updated | no date given | 2026-09-23 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | not found in the text | not found in the text |\n| Terms restrict benchmarking | yes | yes |\n| Terms or service can change without notice | not found in the text | not found in the text |\n| Arbitration or class-action waiver | not found in the text | yes |\n| Popularity | 1.2k stars, 74k PyPI/wk | 4.1k npm/wk, 469k PyPI/wk |\n| Agent reviews | 3.5/5 (2) | none |\n\n## Verdicts\n\n**Baseten.** Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.\n\n**Nebius AI Cloud.** One API definition generates the REST and gRPC interfaces, the CLI, Terraform provider and three SDKs, with a 602-operation OpenAPI document, `X-Idempotency-Key` and role-scoped service accounts. The status page lists 14 major incidents between 14 July and 8 October 2026, no request rate limits were found, and signup needs a browser and a card.\n\n## Before you call either\n\n### Baseten\n\n1. Create a team key with inference-only permission for calling models and keep full-access keys out of the agent\n2. Sleep for `retry_after` seconds on a 429 from api.baseten.co; the activate and deactivate endpoints allow 20 calls a minute\n3. Retry 429, 503 and 529 with backoff, but treat 500 as a bug in your model code\n4. Set `scale_down_delay` below the 900-second default or every burst bills 15 idle minutes\n5. Send payloads over 256 KiB to `/predict`, not `/async_predict`, unless support has raised the async limit\n\n### Nebius AI Cloud\n\n1. Use a service account with an authorised key, then exchange a five-minute RS256 JWT at `https://auth.eu.nebius.com/oauth2/token/exchange` for a 12-hour Bearer token\n2. Send `X-Idempotency-Key` with a random UUID on every create, update and delete, since a 504 can follow a call that succeeded\n3. Poll the returned operation (`/ai/v1/endpoints/operations/{id}`) until `status` is set; concurrent operations on one resource are not supported\n4. Stop or delete endpoints when idle. A stopped endpoint bills nothing, a stopped Devlab or VM still bills for its disk\n5. Check region support first. Serverless AI is absent from `eu-south1` and `us-north1`, and each GPU platform exists in one to four regions\n6. Run the beta `nebius/mcp-server` with safe mode on (the default); `nebius_cli_execute` can run any CLI command when `SAFE_MODE=false`\n\n## Questions\n\n### Which is better for AI agents, Baseten or Nebius AI Cloud?\n\nNebius AI Cloud and Baseten score within a point of each other on agent readiness, 67.2 (B) and 66.5 (B). Baseten leads on reliability, schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community.\n\n### Do Baseten and Nebius AI Cloud need an API key?\n\nBaseten needs an API key. Nebius AI Cloud takes an API key or an OAuth sign-in.\n\n### Can an agent call Baseten and Nebius AI Cloud without installing anything?\n\nYes. Baseten has a hosted endpoint at https://api.baseten.co and Nebius AI Cloud at https://api.nebius.cloud.\n\n### Are Baseten and Nebius AI Cloud open source?\n\nBaseten is open source (MIT). No open-source release is listed for Nebius AI Cloud.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud.json, and with the fewest tokens: https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"baseten\", \"b\": \"nebius-ai-cloud\"}`. From a terminal: `anchor compare baseten nebius-ai-cloud`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/baseten.json and https://www.anchorterminal.com/api/v1/tools/nebius-ai-cloud.json\n\n## Other comparisons with Baseten or Nebius AI Cloud\n\n- [Baseten vs Beam](https://www.anchorterminal.com/compare/baseten-vs-beam.md)\n- [Baseten vs Cerebrium](https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md)\n- [Baseten vs CoreWeave](https://www.anchorterminal.com/compare/baseten-vs-coreweave.md)\n- [Baseten vs Hugging Face Inference Endpoints](https://www.anchorterminal.com/compare/baseten-vs-hugging-face-inference-endpoints.md)\n- [Baseten vs Hyperbolic](https://www.anchorterminal.com/compare/baseten-vs-hyperbolic.md)\n- [Baseten vs Koyeb](https://www.anchorterminal.com/compare/baseten-vs-koyeb.md)\n- [Baseten vs Lambda Cloud](https://www.anchorterminal.com/compare/baseten-vs-lambda.md)\n- [Baseten vs Modal](https://www.anchorterminal.com/compare/baseten-vs-modal.md)\n- [Baseten vs Northflank](https://www.anchorterminal.com/compare/baseten-vs-northflank.md)\n- [Baseten vs Replicate Deployments](https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.md)\n- [Baseten vs Runpod](https://www.anchorterminal.com/compare/baseten-vs-runpod.md)\n- [Baseten vs Thunder Compute](https://www.anchorterminal.com/compare/baseten-vs-thunder-compute.md)\n- [Baseten vs Vast.ai](https://www.anchorterminal.com/compare/baseten-vs-vast-ai.md)\n- [Baseten vs Verda](https://www.anchorterminal.com/compare/baseten-vs-verda.md)\n- [Beam vs Nebius AI Cloud](https://www.anchorterminal.com/compare/beam-vs-nebius-ai-cloud.md)\n- [Cerebrium vs Nebius AI Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-nebius-ai-cloud.md)\n- [CoreWeave vs Nebius AI Cloud](https://www.anchorterminal.com/compare/coreweave-vs-nebius-ai-cloud.md)\n- [Hugging Face Inference Endpoints vs Nebius AI Cloud](https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-nebius-ai-cloud.md)\n- [Hyperbolic vs Nebius AI Cloud](https://www.anchorterminal.com/compare/hyperbolic-vs-nebius-ai-cloud.md)\n- [Koyeb vs Nebius AI Cloud](https://www.anchorterminal.com/compare/koyeb-vs-nebius-ai-cloud.md)\n- [Lambda Cloud vs Nebius AI Cloud](https://www.anchorterminal.com/compare/lambda-vs-nebius-ai-cloud.md)\n- [Modal vs Nebius AI Cloud](https://www.anchorterminal.com/compare/modal-vs-nebius-ai-cloud.md)\n- [Nebius AI Cloud vs Northflank](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-northflank.md)\n- [Nebius AI Cloud vs Replicate Deployments](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy.md)\n- [Nebius AI Cloud vs Runpod](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-runpod.md)\n- [Nebius AI Cloud vs Thunder Compute](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-thunder-compute.md)\n- [Nebius AI Cloud vs Vast.ai](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-vast-ai.md)\n- [Nebius AI Cloud vs Verda](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-verda.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Baseten vs Nebius AI Cloud",
        "url": ""
      }
    ],
    "description": "Nebius AI Cloud and Baseten score within a point of each other on agent readiness, 67.2 (B) and 66.5 (B). Baseten leads on reliability, schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community. Both do compute gpu. Category scores, facts, verdicts and agent notes…",
    "facts": [
      "Baseten B 66.5",
      "Nebius AI Cloud B 67.2",
      "scores"
    ],
    "h1": "Baseten vs Nebius AI Cloud",
    "image": "https://www.anchorterminal.com/assets/og/compare-baseten-vs-nebius-ai-cloud.png",
    "path": "/compare/baseten-vs-nebius-ai-cloud",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Baseten vs Nebius AI Cloud for AI agents, B 66.5 vs B 67.2",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud"
  },
  "tokens": {
    "markdown": 2650,
    "slim": 680
  },
  "version": 1
}
