{
  "data": {
    "a": {
      "slug": "deepinfra",
      "name": "DeepInfra",
      "vendor": "Deep Infra Inc.",
      "vendorUrl": "https://deepinfra.com",
      "kind": "model",
      "category": "inference",
      "summary": "DeepInfra is a hosted inference API for open-weight and some third-party models, covering chat, embeddings, reranking, image, video and speech. It answers OpenAI-style and Anthropic-style calls at api.deepinfra.com with a Bearer key.",
      "url": "https://www.anchorterminal.com/tools/deepinfra",
      "markdownUrl": "https://www.anchorterminal.com/tools/deepinfra.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/deepinfra.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/deepinfra.json",
      "repo": "https://github.com/deepinfra/deepinfra-python",
      "license": "Proprietary service under the DeepInfra Terms of Service. The Python and Node SDKs and the docs repository are MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.deepinfra.com/v1/openai",
      "packages": [
        {
          "registry": "pypi",
          "name": "deepinfra"
        },
        {
          "registry": "npm",
          "name": "deepinfra"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve Bearer key from the dashboard at https://deepinfra.com/dash/api_keys after a browser sign-up with Google, GitHub, email or Okta SSO. The Anthropic-style routes also take the key in `x-api-key`. Keys can carry an IP allowlist and a monthly spending limit, and a key can mint scoped JWTs limited by model, expiry and spend.",
      "pricing": "usage",
      "pricingNotes": "Pay per token, image, audio minute or GPU-hour, with no free tier. An account must add a card or prepay. DeepSeek-V4-Flash-0731 costs $0.06 in and $0.18 out per 1M tokens, the priority tier is 1.5x, flex 0.8x and batch 20 per cent off (https://deepinfra.com/pricing, checked 2026-10-08).",
      "priceSummary": "from $0.06 / 1M in",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs index, the docs repository, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 21,
        "npmWeekly": 1250,
        "pypiWeekly": 50,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.deepinfra.com/",
      "rateLimitsUrl": "https://docs.deepinfra.com/account/rate-limits",
      "llmsTxt": "https://docs.deepinfra.com/llms.txt",
      "openapi": "https://api.deepinfra.com/openapi.json",
      "capabilities": [
        "inference.llm",
        "inference.open-weights",
        "embed.text",
        "rerank",
        "image.generate",
        "speech.stt",
        "speech.tts",
        "compute.batch"
      ],
      "tags": [
        "hosted",
        "model",
        "open-weights",
        "usage-based",
        "card-required",
        "openapi",
        "llms-txt",
        "openai-compatible",
        "anthropic-compatible",
        "batch",
        "prompt-caching",
        "python",
        "typescript",
        "status-page",
        "soc2"
      ],
      "lastRelease": "2026-10-07",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 63,
        "grade": "B",
        "agentReady": false,
        "rank": 371,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 9,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 65,
          "payments": 20,
          "reliability": 70,
          "schema": 69,
          "security": 64,
          "transparency": 67
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "The model list, context sizes and per-token prices are readable without a key, and keys can carry an IP allowlist, a monthly spending cap and model-limited JWTs. Deprecated models get one week's notice and are then redirected to another model, there is no changelog or SLA, and an account needs a card or prepayment before any call.",
        "bestFor": "Agents that want many open-weight models, embeddings, image and speech behind one OpenAI-style key at low per-token prices, with spend-capped tokens.",
        "strengths": [
          "`GET /v1/openai/models` answered without a key on 8 October 2026 with 181 models, each with context length, output cap and per-token prices",
          "Scoped JWTs limit a token to named models, an expiry and a USD spending limit, and each API key can carry an IP allowlist and a monthly cap",
          "The terms commit to zero data retention and no training on Customer Data, and say that clause controls over the privacy policy and docs",
          "Batch API at 20 per cent off, a flex tier at 0.8x, prompt caching with cached-input prices and server-side fallback across up to four models",
          "Status page with 90-day history for the API, the website and 154 models, plus a dated list of scheduled deprecations"
        ],
        "weaknesses": [
          "A deprecated model gets at least one week's notice, and requests are then forwarded to a replacement model under the old id",
          "The status feed lists 50 automated major-outage observations on single models since 3 August 2026, several longer than 24 hours, with no written incident notes",
          "No free tier. The pricing page says an account must add a card or prepay before using the service",
          "No changelog, SLA, error reference or `Retry-After` header was found in the reviewed documentation",
          "The sub-processor list names three companies and is dated 6 September 2024, while the docs say Google and Anthropic receive data for their models"
        ],
        "agentNotes": [
          "Call `GET https://api.deepinfra.com/v1/openai/models` at start-up for ids, context sizes and prices. No key is needed",
          "Check the `model` field of each response. After a deprecation date, requests to the old id are served by a replacement model",
          "Ask the account owner for a scoped JWT limited to the models and spend the task needs, not the full API key",
          "Stay under 200 concurrent requests per model. On 429 `engine_overloaded`, retry after a delay, or send `models` with up to four fallbacks",
          "Never inspect a JWT with `GET /v1/scoped-jwt?jwtoken=`, which puts the token in the URL. Keep credentials in the `Authorization` header"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 63
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 65,
          "payments": 20,
          "reliability": 70,
          "schema": 69,
          "security": 64,
          "transparency": 60
        },
        "provenanceScore": 74
      },
      "connect": {
        "install": "pip install openai   # or: npm install openai",
        "http": "curl \"https://api.deepinfra.com/v1/openai/chat/completions\" \\\n  -H \"Content-Type: application/json\" \\\n  -H \"Authorization: Bearer $DEEPINFRA_API_KEY\" \\\n  -d '{\"model\":\"deepseek-ai/DeepSeek-V4-Flash-0731\",\"messages\":[{\"role\":\"user\",\"content\":\"Hello\"}]}'"
      },
      "letme": {
        "capability": "https://letme.dev/inference.llm",
        "tool": "https://letme.dev/deepinfra"
      },
      "area": "models",
      "provenance": {
        "legalEntity": "Deep Infra Inc.",
        "domain": "deepinfra.com",
        "domainRegistered": "2017-12-08",
        "endpointOnVendorDomain": true,
        "terms": "https://deepinfra.com/terms",
        "privacy": "https://deepinfra.com/privacy",
        "statusPage": "https://status.deepinfra.com",
        "changelog": "",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The Terms of Service, last modified 17 August 2026, name Deep Infra Inc., a Delaware corporation, with California law and JAMS arbitration in San Francisco. They are written around Service Orders and govern the services, the API included.",
          "The Privacy Policy, last modified 15 August 2026, names DeepInfra, Inc., a Delaware corporation, covers the website and the APIs, and has a section on data sent to and returned by the inference service.",
          "No changelog or release notes page was found in the docs index or the site map.",
          "https://deepinfra.com/.well-known/security.txt returns 404.",
          "RDAP for deepinfra.com gives a registration date of 2017-12-08.",
          "The API answers at api.deepinfra.com, the docs at docs.deepinfra.com and the status page at status.deepinfra.com.",
          "No SLA or published DPA was found. The terms mention a DPA only where the parties execute one."
        ],
        "score": 74
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/deepinfra.json",
      "live": {
        "slug": "deepinfra",
        "probe": {
          "target": "https://api.deepinfra.com/v1/openai",
          "method": "get",
          "lastAt": "2026-10-09T11:46:26.647299241Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 371,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 377,
          "p95ms24h": 506,
          "samples24h": 44,
          "samples30d": 44,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 44,
              "ok": 44
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.deepinfra.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-09T07:57:47.5406023Z"
        },
        "updatedAt": "2026-10-09T11:46:26.647299241Z"
      }
    },
    "answer": "SambaCloud scores 66.4 (B) on agent readiness against DeepInfra's 63 (B), and leads in 4 of 7 scored categories. DeepInfra leads on agent ergonomics and security \u0026 auth.",
    "b": {
      "slug": "sambanova",
      "name": "SambaCloud",
      "vendor": "SambaNova Systems, Inc.",
      "vendorUrl": "https://sambanova.ai",
      "kind": "model",
      "category": "inference",
      "summary": "SambaCloud is SambaNova's hosted inference API for open-weight models running on its own RDU processors. It answers OpenAI-style chat, completions and Responses calls and Anthropic-style Messages calls, with Python and TypeScript SDKs.",
      "url": "https://www.anchorterminal.com/tools/sambanova",
      "markdownUrl": "https://www.anchorterminal.com/tools/sambanova.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/sambanova.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/sambanova.json",
      "repo": "https://github.com/sambanova/sambanova-python",
      "license": "Proprietary service under the SambaCloud Terms of Service. The SDKs and the OpenAPI document are Apache-2.0",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.sambanova.ai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "sambanova"
        },
        {
          "registry": "npm",
          "name": "sambanova"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve Bearer key from the SambaCloud console at https://cloud.sambanova.ai/apis after a browser sign-up. The Messages routes also take the same key in `x-api-key`. Up to 25 keys per user, shown once, with no scopes.",
      "pricing": "freemium",
      "pricingNotes": "Free tier with no payment method, at 20 requests a minute, 20 requests a day and 200,000 tokens a day per model. Linking a card moves the account to the Developer tier, billed per token through Stripe, from $0.22 in and $0.59 out per 1M tokens on gpt-oss-120b (https://cloud.sambanova.ai/plans/pricing, checked 2026-10-08).",
      "priceSummary": "from $0.22 / 1M in",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs index, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 2,
        "npmWeekly": 76,
        "pypiWeekly": 6129,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.sambanova.ai/docs/en/get-started/overview",
      "rateLimitsUrl": "https://docs.sambanova.ai/docs/en/models/rate-limits",
      "llmsTxt": "https://docs.sambanova.ai/docs/llms.txt",
      "openapi": "https://raw.githubusercontent.com/sambanova/sambanova-inference-api-spec/refs/heads/main/openapi.documented.json",
      "capabilities": [
        "inference.llm",
        "inference.open-weights",
        "inference.fast"
      ],
      "tags": [
        "hosted",
        "model",
        "open-weights",
        "free-tier",
        "no-card",
        "openapi",
        "llms-txt",
        "openai-compatible",
        "anthropic-compatible",
        "python",
        "typescript",
        "status-page",
        "soc2"
      ],
      "lastRelease": "2026-09-17",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 66.4,
        "grade": "B",
        "agentReady": false,
        "rank": 269,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 7,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 67,
          "maintenance": 77,
          "payments": 40,
          "reliability": 85,
          "schema": 81,
          "security": 58,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": -2,
        "negativeNotes": [
          "2026-10-08. The models and rate-limit pages list `MiniMax-M2.7` as a production model, and the July 2026 release notes say `Mistral-Large-3-675B-Instruct-2512` was promoted to production. Neither is on the pricing page or in `/v1/models`, and the deprecations page records neither. Two points, since a keyed call was not made to confirm the ids fail (https://docs.sambanova.ai/docs/en/models/sambacloud-models, https://api.sambanova.ai/v1/models)."
        ],
        "verdict": "Per-token prices and the model list are readable without a key at `/v1/models`, and a free tier needs no card. Production models get two to three weeks' notice before removal, 13 model ids left between March and June 2026, and the docs disagree with the live catalogue on MiniMax-M2.7.",
        "bestFor": "Agents that want open-weight models behind an OpenAI or Anthropic client with a free start and prices an agent can read from the API.",
        "strengths": [
          "Public OpenAPI 3.1.1 document for all 10 operations, plus `llms.txt` and a Markdown copy of every docs page",
          "`GET /v1/models` answers without a key and returns context length, output cap and per-token prices for each model",
          "Free tier with no payment method at 20 requests a minute, 20 requests a day and 200,000 tokens a day per model",
          "One key works with the OpenAI client, the Anthropic client (`/v1/messages`, `x-api-key`) and SambaNova's own SDKs",
          "Status page with a component per model and no incident posted since 25 June 2026"
        ],
        "weaknesses": [
          "Production models get a notice of two to three weeks, and 13 model ids were removed between 9 March and 9 June 2026",
          "The docs list `MiniMax-M2.7` as a production model, while the pricing page and `/v1/models` did not list it on 8 October 2026",
          "`strict: true` on a JSON schema is accepted and has no effect, and the July 2026 notes record failing structured output on DeepSeek-V3.2",
          "Keys carry no scopes or project limits, and no key rotation guidance was found in the reviewed documentation",
          "The privacy policy is dated 27 May 2023, and no DPA, retention period or sub-processor list was found on the legal page"
        ],
        "agentNotes": [
          "Call `GET https://api.sambanova.ai/v1/models` at start-up and use only ids it returns. The docs name models the endpoint no longer lists",
          "Read `x-ratelimit-remaining-requests` and `x-ratelimit-remaining-requests-day` on every response. The free tier allows 20 requests a day per model",
          "Treat 429 `queue_full` and 503 `maintenance` as retryable after a delay, and 410 `model_deprecated` as a signal to change model",
          "Validate JSON output yourself. Schema enforcement is best effort and `strict: true` changes nothing",
          "Check `max_completion_tokens` per model. DeepSeek-V3.1 caps output at 7,168 tokens and Llama 3.3 70B at 3,072"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 66.4
          }
        ],
        "editorialScores": {
          "ergonomics": 67,
          "maintenance": 77,
          "payments": 40,
          "reliability": 85,
          "schema": 81,
          "security": 58,
          "transparency": 44
        },
        "provenanceScore": 81
      },
      "connect": {
        "install": "pip install sambanova   # or: npm install sambanova",
        "http": "curl https://api.sambanova.ai/v1/chat/completions \\\n  -H \"Authorization: Bearer $SAMBANOVA_API_KEY\" -H \"Content-Type: application/json\" \\\n  -d '{\"model\":\"gpt-oss-120b\",\"messages\":[{\"role\":\"user\",\"content\":\"hello\"}]}'"
      },
      "letme": {
        "capability": "https://letme.dev/inference.llm",
        "tool": "https://letme.dev/sambanova"
      },
      "area": "models",
      "provenance": {
        "legalEntity": "SambaNova Systems, Inc.",
        "domain": "sambanova.ai",
        "domainRegistered": "2017-12-18",
        "endpointOnVendorDomain": true,
        "terms": "https://sambanova.ai/cloud-end-user-license-agreement",
        "privacy": "https://sambanova.ai/privacy-policy",
        "statusPage": "https://status.sambanova.ai",
        "changelog": "https://docs.sambanova.ai/docs/en/release-notes/sambacloud",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The SambaCloud Terms of Service name SambaNova Systems, Inc., a Delaware corporation at 2460 N. First St, Suite #100, San Jose, CA 95131, and define the Service as the SambaCloud platform. The page carries no date. The legal index says it was last updated on 5 February 2026.",
          "The privacy policy is the company's only one. Its last revision is dated 27 May 2023, it covers the website, communications and related services, and it does not name SambaCloud or API inputs.",
          "No DPA, SLA, acceptable use policy or sub-processor list is linked from https://sambanova.ai/legal-agreements.",
          "https://sambanova.ai/.well-known/security.txt returns 404.",
          "RDAP for sambanova.ai gives a registration date of 2017-12-18.",
          "The API answers at api.sambanova.ai and the console at cloud.sambanova.ai."
        ],
        "score": 81
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/sambanova.json",
      "live": {
        "slug": "sambanova",
        "probe": {
          "target": "https://api.sambanova.ai/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:46:39.537400402Z",
          "lastOk": true,
          "lastStatus": 405,
          "lastMs": 408,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 403,
          "p95ms24h": 1118,
          "samples24h": 44,
          "samples30d": 44,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 44,
              "ok": 44
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.sambanova.ai",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T11:40:38.07700378Z"
        },
        "updatedAt": "2026-10-09T11:46:39.537400402Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Deep Infra Inc.",
        "b": "SambaNova Systems, Inc.",
        "name": "Vendor"
      },
      {
        "a": "https://api.deepinfra.com/v1/openai",
        "b": "https://api.sambanova.ai/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under the DeepInfra Terms of Service. The Python and Node SDKs and the docs repository are MIT",
        "b": "Proprietary service under the SambaCloud Terms of Service. The SDKs and the OpenAPI document are Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-10-07",
        "b": "2026-09-17",
        "name": "Last release"
      },
      {
        "a": "2026-08-17",
        "b": "no date given",
        "name": "Terms last updated"
      },
      {
        "a": "2026-08-15",
        "b": "2023-05-27",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "21 stars, 1.2k npm/wk, 50 PyPI/wk",
        "b": "2 stars, 76 npm/wk, 6.1k PyPI/wk",
        "name": "Popularity"
      }
    ],
    "faq": [
      {
        "answer": "SambaCloud scores 66.4 (B) on agent readiness against DeepInfra's 63 (B), and leads in 4 of 7 scored categories. DeepInfra leads on agent ergonomics and security \u0026 auth.",
        "question": "Which is better for AI agents, DeepInfra or SambaCloud?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do DeepInfra and SambaCloud need an API key?"
      },
      {
        "answer": "Yes. DeepInfra has a hosted endpoint at https://api.deepinfra.com/v1/openai and SambaCloud at https://api.sambanova.ai/v1.",
        "question": "Can an agent call DeepInfra and SambaCloud without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Agent ergonomics, 77 against 67",
          "Security \u0026 auth, 64 against 58"
        ],
        "also": null,
        "goodFor": "Agents that want many open-weight models, embeddings, image and speech behind one OpenAI-style key at low per-token prices, with spend-capped tokens.",
        "slug": "deepinfra",
        "watchFor": "A deprecated model gets at least one week's notice, and requests are then forwarded to a replacement model under the old id"
      },
      {
        "aheadOn": [
          "Reliability, 85 against 70",
          "Schema \u0026 documentation, 81 against 69",
          "Payments \u0026 pricing, 40 against 20",
          "Maintenance \u0026 community, 77 against 65"
        ],
        "also": [
          "Free to start without a card"
        ],
        "goodFor": "Agents that want open-weight models behind an OpenAI or Anthropic client with a free start and prices an agent can read from the API.",
        "slug": "sambanova",
        "watchFor": "Production models get a notice of two to three weeks, and 13 model ids were removed between 9 March and 9 June 2026"
      }
    ],
    "job": {
      "capability": "inference.llm",
      "name": "LLM inference"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-deepinfra.json",
        "title": "Claude API vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-sambanova.json",
        "title": "Claude API vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-deepinfra.json",
        "title": "Antseed vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-sambanova.json",
        "title": "Antseed vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-deepinfra.json",
        "title": "BlockRun.AI vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-sambanova.json",
        "title": "BlockRun.AI vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-deepseek-api.json",
        "title": "DeepInfra vs DeepSeek API",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-deepseek-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-gemini-api.json",
        "title": "DeepInfra vs Gemini Developer API",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-gemini-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-groq.json",
        "title": "DeepInfra vs GroqCloud",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-groq"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-mistral-api.json",
        "title": "DeepInfra vs Mistral AI API",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-mistral-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-openai-api.json",
        "title": "DeepInfra vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-openrouter.json",
        "title": "DeepInfra vs OpenRouter",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-openrouter"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.json",
        "title": "DeepInfra vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepseek-api-vs-sambanova.json",
        "title": "DeepSeek API vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/deepseek-api-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-api-vs-sambanova.json",
        "title": "Gemini Developer API vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/gemini-api-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-vs-sambanova.json",
        "title": "GroqCloud vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/groq-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-api-vs-sambanova.json",
        "title": "Mistral AI API vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/mistral-api-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-api-vs-sambanova.json",
        "title": "OpenAI API vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/openai-api-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openrouter-vs-sambanova.json",
        "title": "OpenRouter vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/openrouter-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/prism-inference-vs-sambanova.json",
        "title": "Prism Inference vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/prism-inference-vs-sambanova"
      }
    ],
    "scores": [
      {
        "by": 15,
        "deepinfra": 70,
        "edge": "sambanova",
        "key": "reliability",
        "name": "Reliability",
        "sambanova": 85,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 12,
        "deepinfra": 69,
        "edge": "sambanova",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "sambanova": 81,
        "weight": 13
      },
      {
        "by": 10,
        "deepinfra": 77,
        "edge": "deepinfra",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "sambanova": 67,
        "weight": 13
      },
      {
        "by": 6,
        "deepinfra": 64,
        "edge": "deepinfra",
        "key": "security",
        "name": "Security \u0026 auth",
        "sambanova": 58,
        "weight": 14
      },
      {
        "by": 20,
        "deepinfra": 20,
        "edge": "sambanova",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "sambanova": 40,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 12,
        "deepinfra": 65,
        "edge": "sambanova",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "sambanova": 77,
        "weight": 7
      },
      {
        "by": 4,
        "deepinfra": 67,
        "edge": "deepinfra",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "sambanova": 63,
        "weight": 7
      }
    ],
    "summary": "SambaCloud scores 66.4 (B) on agent readiness against DeepInfra's 63 (B), and leads in 4 of 7 scored categories. DeepInfra leads on agent ergonomics and security \u0026 auth. Both do llm inference.",
    "verdicts": {
      "deepinfra": "The model list, context sizes and per-token prices are readable without a key, and keys can carry an IP allowlist, a monthly spending cap and model-limited JWTs. Deprecated models get one week's notice and are then redirected to another model, there is no changelog or SLA, and an account needs a card or prepayment before any call.",
      "sambanova": "Per-token prices and the model list are readable without a key at `/v1/models`, and a free tier needs no card. Production models get two to three weeks' notice before removal, 13 model ids left between March and June 2026, and the docs disagree with the live catalogue on MiniMax-M2.7."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/deepinfra-vs-sambanova",
    "json": "https://www.anchorterminal.com/compare/deepinfra-vs-sambanova.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/deepinfra-vs-sambanova.md",
    "slim": "https://www.anchorterminal.com/compare/deepinfra-vs-sambanova.min.md"
  },
  "markdown": "SambaCloud scores 66.4 (B) on agent readiness against DeepInfra's 63 (B), and leads in 4 of 7 scored categories. DeepInfra leads on agent ergonomics and security \u0026 auth. Both do llm inference.\n\n- DeepInfra: grade B, 63/100, rank #371 of 842. Markdown https://www.anchorterminal.com/tools/deepinfra.md · JSON https://www.anchorterminal.com/api/v1/tools/deepinfra.json\n- SambaCloud: grade B, 66.4/100, rank #269 of 842. Markdown https://www.anchorterminal.com/tools/sambanova.md · JSON https://www.anchorterminal.com/api/v1/tools/sambanova.json\n\n## Which one, for what\n\n### DeepInfra (B)\n\nGood for: Agents that want many open-weight models, embeddings, image and speech behind one OpenAI-style key at low per-token prices, with spend-capped tokens.\n\nAhead on:\n- Agent ergonomics, 77 against 67\n- Security \u0026 auth, 64 against 58\n\nWatch for: A deprecated model gets at least one week's notice, and requests are then forwarded to a replacement model under the old id\n\n### SambaCloud (B)\n\nGood for: Agents that want open-weight models behind an OpenAI or Anthropic client with a free start and prices an agent can read from the API.\n\nAhead on:\n- Reliability, 85 against 70\n- Schema \u0026 documentation, 81 against 69\n- Payments \u0026 pricing, 40 against 20\n- Maintenance \u0026 community, 77 against 65\n\nAlso in its favour:\n- Free to start without a card\n\nWatch for: Production models get a notice of two to three weeks, and 13 model ids were removed between 9 March and 9 June 2026\n\n\n## Score by category\n\n| Category | Weight | DeepInfra | SambaCloud | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 70 | 85 | SambaCloud +15 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 69 | 81 | SambaCloud +12 |\n| Agent ergonomics | 13% (16.2 this run) | 77 | 67 | DeepInfra +10 |\n| Security \u0026 auth | 14% (17.5 this run) | 64 | 58 | DeepInfra +6 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 40 | SambaCloud +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 65 | 77 | SambaCloud +12 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 67 | 63 | DeepInfra +4 |\n| Negative events | ≤15 | 0 | -2 | |\n| **Total** | | **63 · B** | **66.4 · B** | |\n\n## Facts side by side\n\n| Fact | DeepInfra | SambaCloud |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Deep Infra Inc. | SambaNova Systems, Inc. |\n| Hosted endpoint | `https://api.deepinfra.com/v1/openai` | `https://api.sambanova.ai/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Freemium |\n| x402 | no | no |\n| Licence | Proprietary service under the DeepInfra Terms of Service. The Python and Node SDKs and the docs repository are MIT | Proprietary service under the SambaCloud Terms of Service. The SDKs and the OpenAPI document are Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-10-07 | 2026-09-17 |\n| Terms last updated | 2026-08-17 | no date given |\n| Privacy policy last updated | 2026-08-15 | 2023-05-27 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | not found in the text | not found in the text |\n| Terms restrict benchmarking | yes | yes |\n| Terms or service can change without notice | not found in the text | not found in the text |\n| Arbitration or class-action waiver | yes | not found in the text |\n| Popularity | 21 stars, 1.2k npm/wk, 50 PyPI/wk | 2 stars, 76 npm/wk, 6.1k PyPI/wk |\n\n## Verdicts\n\n**DeepInfra.** The model list, context sizes and per-token prices are readable without a key, and keys can carry an IP allowlist, a monthly spending cap and model-limited JWTs. Deprecated models get one week's notice and are then redirected to another model, there is no changelog or SLA, and an account needs a card or prepayment before any call.\n\n**SambaCloud.** Per-token prices and the model list are readable without a key at `/v1/models`, and a free tier needs no card. Production models get two to three weeks' notice before removal, 13 model ids left between March and June 2026, and the docs disagree with the live catalogue on MiniMax-M2.7.\n\n## Before you call either\n\n### DeepInfra\n\n1. Call `GET https://api.deepinfra.com/v1/openai/models` at start-up for ids, context sizes and prices. No key is needed\n2. Check the `model` field of each response. After a deprecation date, requests to the old id are served by a replacement model\n3. Ask the account owner for a scoped JWT limited to the models and spend the task needs, not the full API key\n4. Stay under 200 concurrent requests per model. On 429 `engine_overloaded`, retry after a delay, or send `models` with up to four fallbacks\n5. Never inspect a JWT with `GET /v1/scoped-jwt?jwtoken=`, which puts the token in the URL. Keep credentials in the `Authorization` header\n\n### SambaCloud\n\n1. Call `GET https://api.sambanova.ai/v1/models` at start-up and use only ids it returns. The docs name models the endpoint no longer lists\n2. Read `x-ratelimit-remaining-requests` and `x-ratelimit-remaining-requests-day` on every response. The free tier allows 20 requests a day per model\n3. Treat 429 `queue_full` and 503 `maintenance` as retryable after a delay, and 410 `model_deprecated` as a signal to change model\n4. Validate JSON output yourself. Schema enforcement is best effort and `strict: true` changes nothing\n5. Check `max_completion_tokens` per model. DeepSeek-V3.1 caps output at 7,168 tokens and Llama 3.3 70B at 3,072\n\n## Questions\n\n### Which is better for AI agents, DeepInfra or SambaCloud?\n\nSambaCloud scores 66.4 (B) on agent readiness against DeepInfra's 63 (B), and leads in 4 of 7 scored categories. DeepInfra leads on agent ergonomics and security \u0026 auth.\n\n### Do DeepInfra and SambaCloud need an API key?\n\nBoth need an API key.\n\n### Can an agent call DeepInfra and SambaCloud without installing anything?\n\nYes. DeepInfra has a hosted endpoint at https://api.deepinfra.com/v1/openai and SambaCloud at https://api.sambanova.ai/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/deepinfra-vs-sambanova.json, and with the fewest tokens: https://www.anchorterminal.com/compare/deepinfra-vs-sambanova.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"deepinfra\", \"b\": \"sambanova\"}`. From a terminal: `anchor compare deepinfra sambanova`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/deepinfra.json and https://www.anchorterminal.com/api/v1/tools/sambanova.json\n\n## Other comparisons with DeepInfra or SambaCloud\n\n- [Claude API vs DeepInfra](https://www.anchorterminal.com/compare/anthropic-api-vs-deepinfra.md)\n- [Claude API vs SambaCloud](https://www.anchorterminal.com/compare/anthropic-api-vs-sambanova.md)\n- [Antseed vs DeepInfra](https://www.anchorterminal.com/compare/antseed-vs-deepinfra.md)\n- [Antseed vs SambaCloud](https://www.anchorterminal.com/compare/antseed-vs-sambanova.md)\n- [BlockRun.AI vs DeepInfra](https://www.anchorterminal.com/compare/blockrun-ai-vs-deepinfra.md)\n- [BlockRun.AI vs SambaCloud](https://www.anchorterminal.com/compare/blockrun-ai-vs-sambanova.md)\n- [DeepInfra vs DeepSeek API](https://www.anchorterminal.com/compare/deepinfra-vs-deepseek-api.md)\n- [DeepInfra vs Gemini Developer API](https://www.anchorterminal.com/compare/deepinfra-vs-gemini-api.md)\n- [DeepInfra vs GroqCloud](https://www.anchorterminal.com/compare/deepinfra-vs-groq.md)\n- [DeepInfra vs Mistral AI API](https://www.anchorterminal.com/compare/deepinfra-vs-mistral-api.md)\n- [DeepInfra vs OpenAI API](https://www.anchorterminal.com/compare/deepinfra-vs-openai-api.md)\n- [DeepInfra vs OpenRouter](https://www.anchorterminal.com/compare/deepinfra-vs-openrouter.md)\n- [DeepInfra vs Prism Inference](https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.md)\n- [DeepSeek API vs SambaCloud](https://www.anchorterminal.com/compare/deepseek-api-vs-sambanova.md)\n- [Gemini Developer API vs SambaCloud](https://www.anchorterminal.com/compare/gemini-api-vs-sambanova.md)\n- [GroqCloud vs SambaCloud](https://www.anchorterminal.com/compare/groq-vs-sambanova.md)\n- [Mistral AI API vs SambaCloud](https://www.anchorterminal.com/compare/mistral-api-vs-sambanova.md)\n- [OpenAI API vs SambaCloud](https://www.anchorterminal.com/compare/openai-api-vs-sambanova.md)\n- [OpenRouter vs SambaCloud](https://www.anchorterminal.com/compare/openrouter-vs-sambanova.md)\n- [Prism Inference vs SambaCloud](https://www.anchorterminal.com/compare/prism-inference-vs-sambanova.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "DeepInfra vs SambaCloud",
        "url": ""
      }
    ],
    "description": "SambaCloud scores 66.4 (B) on agent readiness against DeepInfra's 63 (B), and leads in 4 of 7 scored categories. DeepInfra leads on agent ergonomics and security \u0026 auth. Both do llm inference. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "DeepInfra B 63",
      "SambaCloud B 66.4",
      "scores"
    ],
    "h1": "DeepInfra vs SambaCloud",
    "image": "https://www.anchorterminal.com/assets/og/compare-deepinfra-vs-sambanova.png",
    "path": "/compare/deepinfra-vs-sambanova",
    "published": "2026-10-01",
    "section": "tools",
    "title": "DeepInfra vs SambaCloud for AI agents, B 63 vs B 66.4",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/deepinfra-vs-sambanova"
  },
  "tokens": {
    "markdown": 2350,
    "slim": 630
  },
  "version": 1
}
