{
  "data": {
    "a": {
      "slug": "cerebrium",
      "name": "Cerebrium",
      "vendor": "Cerebrium Inc.",
      "vendorUrl": "https://www.cerebrium.ai",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Cerebrium is a serverless platform for running your own models and code on GPUs and CPUs. A CLI packages code into containers served as REST, streaming and WebSocket endpoints, managed through a REST API.",
      "url": "https://www.anchorterminal.com/tools/cerebrium",
      "markdownUrl": "https://www.anchorterminal.com/tools/cerebrium.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cerebrium.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cerebrium.json",
      "repo": "https://github.com/CerebriumAI/cerebrium",
      "license": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://rest.cerebrium.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cerebrium"
        }
      ],
      "auth": "api-key",
      "authNotes": "Two credentials. A service account token, created in the dashboard or over the API with an expiry of up to one year and a list of granted projects, authenticates the CLI (`CEREBRIUM_SERVICE_ACCOUNT_TOKEN`) and the management API at rest.cerebrium.ai as `Authorization: Bearer`. A project API key (a JWT) authenticates calls to deployed endpoints, and only when `cerebrium.toml` sets `disable_auth = false`. Signup and `cerebrium login` are browser flows.",
      "pricing": "freemium",
      "pricingNotes": "Hobby plan is $0 a month plus compute, Standard $100 a month plus compute, Enterprise on request. GPU, CPU and memory bill per second, from T4 at $0.000164 a second ($0.59 an hour) to H100 at $0.000944 ($3.40) and B200 at $0.00167 ($6.01). CPU $0.00000655 a vCPU-second, memory $0.00000222 a GB-second, storage $0.05 a GB-month after 100 GB free. Listed rates are for the default interruptible tier, and `protected` compute costs twice as much. Cold-start time is free, builds and model initialisation are billed. An account can start on Hobby without a contract. The pricing page does not say whether a card is needed or state a free compute allowance (https://www.cerebrium.ai/pricing, https://cerebrium.ai/docs/calculating-cost).",
      "priceSummary": "$0.0236 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI spec or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": 920,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://cerebrium.ai/docs",
      "llmsTxt": "https://cerebrium.ai/docs/llms.txt",
      "openapi": "https://s3.eu-west-1.amazonaws.com/www.cerebrium.ai/openapi_spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "serverless",
        "gpu",
        "cli",
        "openapi",
        "llms-txt",
        "python",
        "async-jobs",
        "multi-region",
        "status-page",
        "soc2",
        "hipaa"
      ],
      "lastRelease": "2026-09-16",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.3,
        "grade": "C",
        "agentReady": false,
        "rank": 512,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
        "bestFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "strengths": [
          "Per-second prices for ten GPU types published without a login, from T4 at $0.59 an hour to B200 at $6.01",
          "Public OpenAPI 3.0 spec for the management API at rest.cerebrium.ai, with 94 operations, plus llms.txt and Markdown docs",
          "Service account tokens carry an expiry of up to one year and a list of 1 to 50 granted projects",
          "Audit log of 22 actions with actor, IP address and outcome, readable over the API on Standard and Enterprise",
          "Status page with 11 components and 90 days of incident history, and five CLI releases between 7 August and 16 September 2026"
        ],
        "weaknesses": [
          "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it",
          "No request rate limits, 429 guidance, idempotency keys or SLA found in the reviewed documentation",
          "Several multi-hour degradations of the Inference API between 15 July and 2 September 2026, and a 10-minute outage on 21 July",
          "No deprecation policy, platform changelog or public subprocessor list found",
          "The CLI stores tokens in plaintext in `~/.cerebrium/config.yaml` with mode 0644, per its own SECURITY.md"
        ],
        "agentNotes": [
          "Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL",
          "Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser",
          "Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours",
          "Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller",
          "Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.3
          }
        ],
        "editorialScores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 45
        },
        "provenanceScore": 80
      },
      "connect": {
        "install": "pip install cerebrium \u0026\u0026 cerebrium login",
        "http": "curl --location --request POST 'https://api.cerebrium.ai/v4/p-xxxxxxxx/{app-name}/{function}' \\\n  --header 'Authorization: Bearer \u003cJWT_TOKEN\u003e' \\\n  --header 'Content-Type: application/json' \\\n  --data '{\"function_param\": \"data\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/cerebrium"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.01,
          "note": "$0.00167 a second, interruptible tier"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.2,
          "note": "$0.001166 a second, interruptible tier"
        },
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.4,
          "note": "$0.000944 a second, interruptible tier"
        },
        {
          "item": "RTX PRO 6000 96 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second, interruptible tier"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.1,
          "note": "$0.000583 a second, interruptible tier"
        },
        {
          "item": "A100 40 GB",
          "unit": "gpu-hour",
          "usd": 2,
          "note": "$0.000555 a second, interruptible tier"
        },
        {
          "item": "L40s 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second, interruptible tier"
        },
        {
          "item": "A10 24 GB",
          "unit": "gpu-hour",
          "usd": 1.1,
          "note": "$0.000306 a second, interruptible tier"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second, interruptible tier"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second, interruptible tier"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.0236,
          "note": "$0.00000655 a vCPU-second, memory extra at $0.00000222 a GB-second"
        },
        {
          "item": "Persistent storage",
          "unit": "gb-month",
          "usd": 0.05,
          "note": "First 100 GB free"
        },
        {
          "item": "Standard plan",
          "unit": "month",
          "usd": 100,
          "note": "Plus compute"
        }
      ],
      "provenance": {
        "legalEntity": "Cerebrium Inc.",
        "domain": "cerebrium.ai",
        "domainRegistered": "2021-06-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cerebrium.ai/terms-of-service",
        "privacy": "https://www.cerebrium.ai/privacy",
        "statusPage": "https://status.cerebrium.ai",
        "changelog": "https://github.com/CerebriumAI/cerebrium/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The terms of service name Cerebrium Inc and say they are governed by the laws of the United Kingdom. The privacy policy names Cerebrium Inc. as data controller at 251 Little Falls Drive, Wilmington, Delaware.",
          "The terms of service are the only terms Cerebrium publishes. They cover accounts, subscriptions and the Service, and the OpenAPI spec names them as the API's licence. Neither document states a date.",
          "Deployed endpoints answer at api.cerebrium.ai and the management API at rest.cerebrium.ai. The OpenAPI file is served from an AWS S3 bucket.",
          "cerebrium.ai/.well-known/security.txt and cerebrium.ai/security.txt returned 404 on 8 October 2026. The CLI repository's SECURITY.md and the docs give security@cerebrium.ai.",
          "The changelog link is the CLI's GitHub releases. No platform changelog was found. Domain registration date from RDAP."
        ],
        "score": 80
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cerebrium.json",
      "live": {
        "slug": "cerebrium",
        "probe": {
          "target": "https://rest.cerebrium.ai",
          "method": "get",
          "lastAt": "2026-10-08T19:52:47.839655439Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 248,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 252,
          "p95ms24h": 326,
          "samples24h": 27,
          "samples30d": 27,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 27,
              "ok": 27
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cerebrium.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T19:38:18.718586754Z"
        },
        "pages": [
          {
            "url": "https://www.cerebrium.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:57.20274404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "be1ada48db42"
          },
          {
            "url": "https://www.cerebrium.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:59.281048683Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "017a8d228768"
          },
          {
            "url": "https://www.cerebrium.ai/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:27:01.554163635Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a99804810a11"
          }
        ],
        "updatedAt": "2026-10-08T19:52:47.839655439Z"
      }
    },
    "answer": "Vast.ai scores 62.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Cerebrium leads on payments \u0026 pricing.",
    "b": {
      "slug": "vast-ai",
      "name": "Vast.ai",
      "vendor": "Vast.ai Inc.",
      "vendorUrl": "https://vast.ai",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Vast.ai is a marketplace for renting GPUs by the second from independent hosts and data centres, as Docker instances, virtual machines or autoscaling serverless endpoints. Agents use the `vastai` CLI, a Python SDK or a REST API.",
      "url": "https://www.anchorterminal.com/tools/vast-ai",
      "markdownUrl": "https://www.anchorterminal.com/tools/vast-ai.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/vast-ai.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/vast-ai.json",
      "repo": "https://github.com/vast-ai/vast-cli",
      "license": "Proprietary service under Vast.ai's Terms of Use Agreement. The `vastai` CLI and Python SDK on GitHub are MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://console.vast.ai/api/v0",
      "packages": [
        {
          "registry": "pypi",
          "name": "vastai"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key from the console's Keys page, sent as `Authorization: Bearer` to https://console.vast.ai/api/v0, or stored once with `vastai set api-key`. Access is self-serve after a browser signup and email verification. A key has full account access by default. Scoped keys take any of 11 permission categories (`instance_read`, `instance_write`, `billing_write` and others) and constraints that narrow an endpoint to resource IDs. Keys are shown once, and can be reset or deleted with immediate effect. The rate-limit page names an `api_key` query parameter as part of a caller's identity, while the authentication page documents only the header.",
      "pricing": "usage",
      "pricingNotes": "Prepaid credit with a $5 minimum deposit and no free tier or trial found. Each host sets its own rate per listing, so there is no fixed price list. Live rates are public at vast.ai/pricing and through `vastai search offers`. GPU time bills per second while an instance runs, storage per second while it exists (stopped included) and bandwidth per byte. Interruptible rentals are bid-priced and reserved rentals are prepaid for 1, 3 or 6 months. Serverless adds no fee on top of instance rates. Spent credit is not refunded (https://docs.vast.ai/guides/reference/billing.md, https://docs.vast.ai/guides/instances/pricing.md, checked 2026-10-08).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs index, llms.txt, the agent pricing page or the OpenAPI file (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 223,
        "npmWeekly": null,
        "pypiWeekly": 32699,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.vast.ai",
      "llmsTxt": "https://docs.vast.ai/llms.txt",
      "openapi": "https://docs.vast.ai/api-reference/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.containers",
        "compute.serverless",
        "compute.endpoints"
      ],
      "tags": [
        "hosted",
        "marketplace",
        "usage-priced",
        "prepaid",
        "api-key",
        "scoped-keys",
        "openapi",
        "llms-txt",
        "python",
        "cli",
        "agent-skill",
        "serverless",
        "spot",
        "status-page",
        "soc2"
      ],
      "lastRelease": "2026-10-02",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 62.6,
        "grade": "B",
        "agentReady": false,
        "rank": 331,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 4,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 57,
          "maintenance": 81,
          "payments": 20,
          "reliability": 65,
          "schema": 78,
          "security": 72,
          "transparency": 62
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "API keys can be limited by permission category and by resource ID, the REST API has a public OpenAPI 3.1 file, and the CLI ships weekly with a skill file for coding agents. Machines belong to independent hosts, prices move with the market, rate-limit thresholds are unpublished, and there is no SLA or free tier.",
        "bestFor": "Cost-sensitive training, batch work and self-managed inference where the agent can search listings, set a price cap and tolerate host variance or interruption.",
        "strengths": [
          "API keys take 11 permission categories and per-endpoint constraints on resource IDs, and can be reset or deleted at once",
          "Public OpenAPI 3.1 file with 89 operations, llms.txt and Markdown docs pages",
          "MIT CLI and Python SDK, v1.8.3 on 2 October 2026, with 12 tagged releases since 27 July 2026",
          "Per-second billing, with serverless workers charged at the same rates as rented instances",
          "Vulnerability disclosure policy with safe harbour and triage within 5 business days, and an account audit log from the CLI"
        ],
        "weaknesses": [
          "No SLA. The terms say availability is not guaranteed and the service can change without notice",
          "Rate-limit thresholds are unpublished and 429 responses carry no `Retry-After` header",
          "No free tier. Credit is prepaid with a $5 minimum deposit after a browser signup and email verification",
          "Hosts are independent. The security FAQ says individual hosts may have less formal security than Secure Cloud data centres",
          "The terms of 1 September 2026 prohibit scripts and automated access to the services without a separate written agreement, which conflicts with the public API"
        ],
        "agentNotes": [
          "Create a scoped key with `vastai create api-key --permissions` for the agent. A default key has full account access, billing and key management included",
          "Pass `--raw` on every CLI command for JSON output, and `-y` on `vastai destroy instance`, which otherwise waits for a confirmation prompt",
          "Register an SSH key with `vastai create ssh-key` before creating an instance, or the host is unreachable",
          "Destroy instances when finished. A stopped instance still bills storage, and a zero balance without a saved card leads to deletion",
          "Back off on HTTP 429 yourself when calling REST directly. The CLI retries 429 three times from 0.15 seconds",
          "Filter searches with `verified=true` or the Secure Cloud option for sensitive data, and set a `dph_total` price cap"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 62.6
          }
        ],
        "editorialScores": {
          "ergonomics": 57,
          "maintenance": 81,
          "payments": 20,
          "reliability": 65,
          "schema": 78,
          "security": 72,
          "transparency": 45
        },
        "provenanceScore": 79
      },
      "connect": {
        "install": "pip install vastai",
        "http": "curl -s -H \"Authorization: Bearer $VAST_API_KEY\" \\\n  \"https://console.vast.ai/api/v0/users/current/\"",
        "claudeCode": "/plugin marketplace add vast-ai/vast-claude-plugin\n/plugin install vastai"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/vast-ai"
      },
      "area": "models",
      "provenance": {
        "legalEntity": "Vast.ai Inc.",
        "domain": "vast.ai",
        "domainRegistered": "2017-12-16",
        "endpointOnVendorDomain": true,
        "terms": "https://vast.ai/terms",
        "privacy": "https://vast.ai/privacy",
        "statusPage": "https://status.vast.ai",
        "changelog": "https://github.com/vast-ai/vast-cli/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The Terms of Use Agreement, version date 1 September 2026, names Vast.ai Inc. and its affiliates, covers the website and the rental service, and is governed by California law with JAMS arbitration in the Los Angeles area.",
          "The privacy policy, version date 18 June 2025, names Vast.ai Inc. and a Privacy Officer at contact@vast.ai. It is written for users of the website and contains an unfilled `[INSERT HYPERLINK]` placeholder.",
          "A Data Processing Agreement at vast.ai/data-processing-agreement forms part of the terms, with standard contractual clauses and five named sub-processors (Stripe, Google, Meta, Twitter, Microsoft).",
          "The REST API answers at console.vast.ai and serverless routing at run.vast.ai, both on the vendor's domain. Rented machines belong to independent hosts.",
          "vast.ai/.well-known/security.txt returns 404. A vulnerability disclosure policy dated 23 July 2025 at vast.ai/vulnerability-disclosure-policy sends reports to security@vast.ai.",
          "No API changelog was found in the docs index. The changelog link is the CLI and SDK repository's release list, which has no CHANGELOG file.",
          "RDAP for vast.ai gives a registration date of 2017-12-16."
        ],
        "score": 79
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/vast-ai.json",
      "live": {
        "slug": "vast-ai",
        "probe": {
          "target": "https://console.vast.ai/api/v0",
          "method": "get",
          "lastAt": "2026-10-08T19:53:05.315409291Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 341,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 340,
          "p95ms24h": 407,
          "samples24h": 27,
          "samples30d": 27,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 27,
              "ok": 27
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.vast.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T19:39:18.949416936Z"
        },
        "pages": [
          {
            "url": "https://docs.vast.ai/guides/instances/pricing.md",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:19:43.22968168Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "7cd305f9206a"
          },
          {
            "url": "https://vast.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:25:30.293684318Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "afe890f5aa20"
          },
          {
            "url": "https://vast.ai/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:25:32.570996649Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "652ca94c3dd7"
          }
        ],
        "updatedAt": "2026-10-08T19:53:05.315409291Z"
      }
    },
    "facts": [
      {
        "a": "Model platform",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Cerebrium Inc.",
        "b": "Vast.ai Inc.",
        "name": "Vendor"
      },
      {
        "a": "https://rest.cerebrium.ai",
        "b": "https://console.vast.ai/api/v0",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
        "b": "Proprietary service under Vast.ai's Terms of Use Agreement. The `vastai` CLI and Python SDK on GitHub are MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-16",
        "b": "2026-10-02",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "2026-09-01",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2025-06-18",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "920 PyPI/wk",
        "b": "223 stars, 33k PyPI/wk",
        "name": "Popularity"
      }
    ],
    "faq": [
      {
        "answer": "Vast.ai scores 62.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Cerebrium leads on payments \u0026 pricing.",
        "question": "Which is better for AI agents, Cerebrium or Vast.ai?"
      },
      {
        "answer": "Yes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Vast.ai at https://console.vast.ai/api/v0.",
        "question": "Can an agent call Cerebrium and Vast.ai without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Payments \u0026 pricing, 30 against 20"
        ],
        "also": null,
        "goodFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "slug": "cerebrium",
        "watchFor": "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it"
      },
      {
        "aheadOn": [
          "Reliability, 65 against 48",
          "Schema \u0026 documentation, 78 against 70",
          "Agent ergonomics, 57 against 49",
          "Security \u0026 auth, 72 against 60",
          "Maintenance \u0026 community, 81 against 75"
        ],
        "also": null,
        "goodFor": "Cost-sensitive training, batch work and self-managed inference where the agent can search listings, set a price cap and tolerate host variance or interruption.",
        "slug": "vast-ai",
        "watchFor": "No SLA. The terms say availability is not guaranteed and the service can change without notice"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
        "title": "Baseten vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-vast-ai.json",
        "title": "Baseten vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-cerebrium.json",
        "title": "Beam vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/beam-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-vast-ai.json",
        "title": "Beam vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/beam-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.json",
        "title": "Cerebrium vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json",
        "title": "Cerebrium vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json",
        "title": "Cerebrium vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.json",
        "title": "Cerebrium vs Modal",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank.json",
        "title": "Cerebrium vs Northflank",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
        "title": "Cerebrium vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json",
        "title": "Cerebrium vs Runpod",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-vast-ai.json",
        "title": "CoreWeave vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-vast-ai.json",
        "title": "Koyeb vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-vast-ai.json",
        "title": "Lambda Cloud vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-vast-ai.json",
        "title": "Modal vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/modal-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/northflank-vs-vast-ai.json",
        "title": "Northflank vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/northflank-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai.json",
        "title": "Replicate Deployments vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/runpod-vs-vast-ai.json",
        "title": "Runpod vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/runpod-vs-vast-ai"
      }
    ],
    "scores": [
      {
        "by": 17,
        "cerebrium": 48,
        "edge": "vast-ai",
        "key": "reliability",
        "name": "Reliability",
        "vast-ai": 65,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 8,
        "cerebrium": 70,
        "edge": "vast-ai",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "vast-ai": 78,
        "weight": 13
      },
      {
        "by": 8,
        "cerebrium": 49,
        "edge": "vast-ai",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "vast-ai": 57,
        "weight": 13
      },
      {
        "by": 12,
        "cerebrium": 60,
        "edge": "vast-ai",
        "key": "security",
        "name": "Security \u0026 auth",
        "vast-ai": 72,
        "weight": 14
      },
      {
        "by": 10,
        "cerebrium": 30,
        "edge": "cerebrium",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "vast-ai": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 6,
        "cerebrium": 75,
        "edge": "vast-ai",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "vast-ai": 81,
        "weight": 7
      },
      {
        "by": 1,
        "cerebrium": 63,
        "edge": "cerebrium",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "vast-ai": 62,
        "weight": 7
      }
    ],
    "summary": "Vast.ai scores 62.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Cerebrium leads on payments \u0026 pricing. Both do compute gpu.",
    "verdicts": {
      "cerebrium": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
      "vast-ai": "API keys can be limited by permission category and by resource ID, the REST API has a public OpenAPI 3.1 file, and the CLI ships weekly with a skill file for coding agents. Machines belong to independent hosts, prices move with the market, rate-limit thresholds are unpublished, and there is no SLA or free tier."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai",
    "json": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.md",
    "slim": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.min.md"
  },
  "markdown": "Vast.ai scores 62.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Cerebrium leads on payments \u0026 pricing. Both do compute gpu.\n\n- Cerebrium: grade C, 55.3/100, rank #512 of 722. Markdown https://www.anchorterminal.com/tools/cerebrium.md · JSON https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n- Vast.ai: grade B, 62.6/100, rank #331 of 722. Markdown https://www.anchorterminal.com/tools/vast-ai.md · JSON https://www.anchorterminal.com/api/v1/tools/vast-ai.json\n\n## Which one, for what\n\n### Cerebrium (C)\n\nGood for: Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.\n\nAhead on:\n- Payments \u0026 pricing, 30 against 20\n\nWatch for: `disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it\n\n### Vast.ai (B)\n\nGood for: Cost-sensitive training, batch work and self-managed inference where the agent can search listings, set a price cap and tolerate host variance or interruption.\n\nAhead on:\n- Reliability, 65 against 48\n- Schema \u0026 documentation, 78 against 70\n- Agent ergonomics, 57 against 49\n- Security \u0026 auth, 72 against 60\n- Maintenance \u0026 community, 81 against 75\n\nWatch for: No SLA. The terms say availability is not guaranteed and the service can change without notice\n\n\n## Score by category\n\n| Category | Weight | Cerebrium | Vast.ai | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 48 | 65 | Vast.ai +17 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 70 | 78 | Vast.ai +8 |\n| Agent ergonomics | 13% (16.2 this run) | 49 | 57 | Vast.ai +8 |\n| Security \u0026 auth | 14% (17.5 this run) | 60 | 72 | Vast.ai +12 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 20 | Cerebrium +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 81 | Vast.ai +6 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 63 | 62 | Cerebrium +1 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **55.3 · C** | **62.6 · B** | |\n\n## Facts side by side\n\n| Fact | Cerebrium | Vast.ai |\n| --- | --- | --- |\n| Kind | Model platform | HTTP API |\n| Vendor | Cerebrium Inc. | Vast.ai Inc. |\n| Hosted endpoint | `https://rest.cerebrium.ai` | `https://console.vast.ai/api/v0` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | Proprietary service under Cerebrium's terms of service. The CLI is MIT | Proprietary service under Vast.ai's Terms of Use Agreement. The `vastai` CLI and Python SDK on GitHub are MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-16 | 2026-10-02 |\n| Terms last updated | no date given | 2026-09-01 |\n| Privacy policy last updated | no date given | 2025-06-18 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | yes |\n| Terms restrict benchmarking | not found in the text | yes |\n| Terms or service can change without notice | yes | yes |\n| Arbitration or class-action waiver | not found in the text | yes |\n| Popularity | 920 PyPI/wk | 223 stars, 33k PyPI/wk |\n\n## Verdicts\n\n**Cerebrium.** Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.\n\n**Vast.ai.** API keys can be limited by permission category and by resource ID, the REST API has a public OpenAPI 3.1 file, and the CLI ships weekly with a skill file for coding agents. Machines belong to independent hosts, prices move with the market, rate-limit thresholds are unpublished, and there is no SLA or free tier.\n\n## Before you call either\n\n### Cerebrium\n\n1. Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL\n2. Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser\n3. Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours\n4. Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller\n5. Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate\n\n### Vast.ai\n\n1. Create a scoped key with `vastai create api-key --permissions` for the agent. A default key has full account access, billing and key management included\n2. Pass `--raw` on every CLI command for JSON output, and `-y` on `vastai destroy instance`, which otherwise waits for a confirmation prompt\n3. Register an SSH key with `vastai create ssh-key` before creating an instance, or the host is unreachable\n4. Destroy instances when finished. A stopped instance still bills storage, and a zero balance without a saved card leads to deletion\n5. Back off on HTTP 429 yourself when calling REST directly. The CLI retries 429 three times from 0.15 seconds\n6. Filter searches with `verified=true` or the Secure Cloud option for sensitive data, and set a `dph_total` price cap\n\n## Questions\n\n### Which is better for AI agents, Cerebrium or Vast.ai?\n\nVast.ai scores 62.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Cerebrium leads on payments \u0026 pricing.\n\n### Can an agent call Cerebrium and Vast.ai without installing anything?\n\nYes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Vast.ai at https://console.vast.ai/api/v0.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cerebrium\", \"b\": \"vast-ai\"}`. From a terminal: `anchor compare cerebrium vast-ai`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cerebrium.json and https://www.anchorterminal.com/api/v1/tools/vast-ai.json\n\n## Other comparisons with Cerebrium or Vast.ai\n\n- [Baseten vs Cerebrium](https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md)\n- [Baseten vs Vast.ai](https://www.anchorterminal.com/compare/baseten-vs-vast-ai.md)\n- [Beam vs Cerebrium](https://www.anchorterminal.com/compare/beam-vs-cerebrium.md)\n- [Beam vs Vast.ai](https://www.anchorterminal.com/compare/beam-vs-vast-ai.md)\n- [Cerebrium vs CoreWeave](https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.md)\n- [Cerebrium vs Koyeb](https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.md)\n- [Cerebrium vs Lambda Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-lambda.md)\n- [Cerebrium vs Modal](https://www.anchorterminal.com/compare/cerebrium-vs-modal.md)\n- [Cerebrium vs Northflank](https://www.anchorterminal.com/compare/cerebrium-vs-northflank.md)\n- [Cerebrium vs Replicate Deployments](https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md)\n- [Cerebrium vs Runpod](https://www.anchorterminal.com/compare/cerebrium-vs-runpod.md)\n- [CoreWeave vs Vast.ai](https://www.anchorterminal.com/compare/coreweave-vs-vast-ai.md)\n- [Koyeb vs Vast.ai](https://www.anchorterminal.com/compare/koyeb-vs-vast-ai.md)\n- [Lambda Cloud vs Vast.ai](https://www.anchorterminal.com/compare/lambda-vs-vast-ai.md)\n- [Modal vs Vast.ai](https://www.anchorterminal.com/compare/modal-vs-vast-ai.md)\n- [Northflank vs Vast.ai](https://www.anchorterminal.com/compare/northflank-vs-vast-ai.md)\n- [Replicate Deployments vs Vast.ai](https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai.md)\n- [Runpod vs Vast.ai](https://www.anchorterminal.com/compare/runpod-vs-vast-ai.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-08",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cerebrium vs Vast.ai",
        "url": ""
      }
    ],
    "description": "Vast.ai scores 62.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Cerebrium leads on payments \u0026 pricing. Both do compute gpu. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Cerebrium C 55.3",
      "Vast.ai B 62.6",
      "scores"
    ],
    "h1": "Cerebrium vs Vast.ai",
    "image": "https://www.anchorterminal.com/assets/og/compare-cerebrium-vs-vast-ai.png",
    "path": "/compare/cerebrium-vs-vast-ai",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cerebrium vs Vast.ai for AI agents, C 55.3 vs B 62.6 | Anchor Terminal",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai"
  },
  "tokens": {
    "markdown": 2200,
    "slim": 630
  },
  "version": 1
}
