{
  "data": {
    "a": {
      "slug": "deepinfra",
      "name": "DeepInfra",
      "vendor": "Deep Infra Inc.",
      "vendorUrl": "https://deepinfra.com",
      "kind": "model",
      "category": "inference",
      "summary": "DeepInfra is a hosted inference API for open-weight and some third-party models, covering chat, embeddings, reranking, image, video and speech. It answers OpenAI-style and Anthropic-style calls at api.deepinfra.com with a Bearer key.",
      "url": "https://www.anchorterminal.com/tools/deepinfra",
      "markdownUrl": "https://www.anchorterminal.com/tools/deepinfra.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/deepinfra.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/deepinfra.json",
      "repo": "https://github.com/deepinfra/deepinfra-python",
      "license": "Proprietary service under the DeepInfra Terms of Service. The Python and Node SDKs and the docs repository are MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.deepinfra.com/v1/openai",
      "packages": [
        {
          "registry": "pypi",
          "name": "deepinfra"
        },
        {
          "registry": "npm",
          "name": "deepinfra"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve Bearer key from the dashboard at https://deepinfra.com/dash/api_keys after a browser sign-up with Google, GitHub, email or Okta SSO. The Anthropic-style routes also take the key in `x-api-key`. Keys can carry an IP allowlist and a monthly spending limit, and a key can mint scoped JWTs limited by model, expiry and spend.",
      "pricing": "usage",
      "pricingNotes": "Pay per token, image, audio minute or GPU-hour, with no free tier. An account must add a card or prepay. DeepSeek-V4-Flash-0731 costs $0.06 in and $0.18 out per 1M tokens, the priority tier is 1.5x, flex 0.8x and batch 20 per cent off (https://deepinfra.com/pricing, checked 2026-10-08).",
      "priceSummary": "from $0.06 / 1M in",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs index, the docs repository, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 21,
        "npmWeekly": 1250,
        "pypiWeekly": 50,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.deepinfra.com/",
      "rateLimitsUrl": "https://docs.deepinfra.com/account/rate-limits",
      "llmsTxt": "https://docs.deepinfra.com/llms.txt",
      "openapi": "https://api.deepinfra.com/openapi.json",
      "capabilities": [
        "inference.llm",
        "inference.open-weights",
        "embed.text",
        "rerank",
        "image.generate",
        "speech.stt",
        "speech.tts",
        "compute.batch"
      ],
      "tags": [
        "hosted",
        "model",
        "open-weights",
        "usage-based",
        "card-required",
        "openapi",
        "llms-txt",
        "openai-compatible",
        "anthropic-compatible",
        "batch",
        "prompt-caching",
        "python",
        "typescript",
        "status-page",
        "soc2"
      ],
      "lastRelease": "2026-10-07",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 63,
        "grade": "B",
        "agentReady": false,
        "rank": 371,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 9,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 65,
          "payments": 20,
          "reliability": 70,
          "schema": 69,
          "security": 64,
          "transparency": 67
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "The model list, context sizes and per-token prices are readable without a key, and keys can carry an IP allowlist, a monthly spending cap and model-limited JWTs. Deprecated models get one week's notice and are then redirected to another model, there is no changelog or SLA, and an account needs a card or prepayment before any call.",
        "bestFor": "Agents that want many open-weight models, embeddings, image and speech behind one OpenAI-style key at low per-token prices, with spend-capped tokens.",
        "strengths": [
          "`GET /v1/openai/models` answered without a key on 8 October 2026 with 181 models, each with context length, output cap and per-token prices",
          "Scoped JWTs limit a token to named models, an expiry and a USD spending limit, and each API key can carry an IP allowlist and a monthly cap",
          "The terms commit to zero data retention and no training on Customer Data, and say that clause controls over the privacy policy and docs",
          "Batch API at 20 per cent off, a flex tier at 0.8x, prompt caching with cached-input prices and server-side fallback across up to four models",
          "Status page with 90-day history for the API, the website and 154 models, plus a dated list of scheduled deprecations"
        ],
        "weaknesses": [
          "A deprecated model gets at least one week's notice, and requests are then forwarded to a replacement model under the old id",
          "The status feed lists 50 automated major-outage observations on single models since 3 August 2026, several longer than 24 hours, with no written incident notes",
          "No free tier. The pricing page says an account must add a card or prepay before using the service",
          "No changelog, SLA, error reference or `Retry-After` header was found in the reviewed documentation",
          "The sub-processor list names three companies and is dated 6 September 2024, while the docs say Google and Anthropic receive data for their models"
        ],
        "agentNotes": [
          "Call `GET https://api.deepinfra.com/v1/openai/models` at start-up for ids, context sizes and prices. No key is needed",
          "Check the `model` field of each response. After a deprecation date, requests to the old id are served by a replacement model",
          "Ask the account owner for a scoped JWT limited to the models and spend the task needs, not the full API key",
          "Stay under 200 concurrent requests per model. On 429 `engine_overloaded`, retry after a delay, or send `models` with up to four fallbacks",
          "Never inspect a JWT with `GET /v1/scoped-jwt?jwtoken=`, which puts the token in the URL. Keep credentials in the `Authorization` header"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 63
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 65,
          "payments": 20,
          "reliability": 70,
          "schema": 69,
          "security": 64,
          "transparency": 60
        },
        "provenanceScore": 74
      },
      "connect": {
        "install": "pip install openai   # or: npm install openai",
        "http": "curl \"https://api.deepinfra.com/v1/openai/chat/completions\" \\\n  -H \"Content-Type: application/json\" \\\n  -H \"Authorization: Bearer $DEEPINFRA_API_KEY\" \\\n  -d '{\"model\":\"deepseek-ai/DeepSeek-V4-Flash-0731\",\"messages\":[{\"role\":\"user\",\"content\":\"Hello\"}]}'"
      },
      "letme": {
        "capability": "https://letme.dev/inference.llm",
        "tool": "https://letme.dev/deepinfra"
      },
      "area": "models",
      "provenance": {
        "legalEntity": "Deep Infra Inc.",
        "domain": "deepinfra.com",
        "domainRegistered": "2017-12-08",
        "endpointOnVendorDomain": true,
        "terms": "https://deepinfra.com/terms",
        "privacy": "https://deepinfra.com/privacy",
        "statusPage": "https://status.deepinfra.com",
        "changelog": "",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The Terms of Service, last modified 17 August 2026, name Deep Infra Inc., a Delaware corporation, with California law and JAMS arbitration in San Francisco. They are written around Service Orders and govern the services, the API included.",
          "The Privacy Policy, last modified 15 August 2026, names DeepInfra, Inc., a Delaware corporation, covers the website and the APIs, and has a section on data sent to and returned by the inference service.",
          "No changelog or release notes page was found in the docs index or the site map.",
          "https://deepinfra.com/.well-known/security.txt returns 404.",
          "RDAP for deepinfra.com gives a registration date of 2017-12-08.",
          "The API answers at api.deepinfra.com, the docs at docs.deepinfra.com and the status page at status.deepinfra.com.",
          "No SLA or published DPA was found. The terms mention a DPA only where the parties execute one."
        ],
        "score": 74
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/deepinfra.json",
      "live": {
        "slug": "deepinfra",
        "probe": {
          "target": "https://api.deepinfra.com/v1/openai",
          "method": "get",
          "lastAt": "2026-10-09T11:28:58.684383058Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 1141,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 379,
          "p95ms24h": 506,
          "samples24h": 41,
          "samples30d": 41,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 41,
              "ok": 41
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.deepinfra.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-09T07:57:47.5406023Z"
        },
        "updatedAt": "2026-10-09T11:28:58.684383058Z"
      }
    },
    "answer": "OpenAI API scores 83.3 (A) on agent readiness against DeepInfra's 63 (B), and leads in 6 of 7 scored categories.",
    "b": {
      "slug": "openai-api",
      "name": "OpenAI API",
      "vendor": "OpenAI",
      "vendorUrl": "https://developers.openai.com",
      "kind": "model",
      "category": "inference",
      "summary": "OpenAI's API for accessing its models through Responses, Chat Completions and Batch endpoints.",
      "url": "https://www.anchorterminal.com/tools/openai-api",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-api.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-api.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-api.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Apache-2.0 (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        },
        {
          "registry": "npm",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "`Authorization: Bearer` with a project key. Some models and tools need business or ID verification first.",
      "pricing": "usage",
      "pricingNotes": "Prepaid credit, with tier 1 reached at $5 paid. Long context over 272K tokens costs 2x input and 1.5x output. Cache reads 0.1x input (0.05x on GPT-6.1 Sol), cache writes 1.25x from GPT-5.6 on. Batch and Flex half price, Fast mode 2x, Ultrafast 6x (GPT-6 Astra only). Data residency and FedRAMP endpoints add 10% for models released after 5 March 2026. Web search $10 per 1,000 calls plus search content tokens, file search $2.50 per 1,000, hosted containers from $0.03 per 20-minute session (https://developers.openai.com/api/docs/pricing).",
      "priceSummary": "from $0.10 / 1M in",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Access needs an account with prepaid credit.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31300,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-09-26"
      },
      "docsUrl": "https://developers.openai.com/api/docs",
      "rateLimitsUrl": "https://developers.openai.com/api/docs/guides/rate-limits",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi",
      "capabilities": [
        "inference.llm"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "card-required",
        "openapi",
        "llms-txt"
      ],
      "lastRelease": "2026-10-05",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 83.3,
        "grade": "A",
        "agentReady": true,
        "rank": 4,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 1,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 98,
          "maintenance": 91,
          "payments": 30,
          "reliability": 70,
          "schema": 100,
          "security": 100,
          "transparency": 90
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-05"
        },
        "negative": 0,
        "verdict": "Official OpenAPI document and an llms.txt index. Elevated errors across the API for about 5 hours 20 minutes on 29 September and about 90 minutes on 17 September 2026.",
        "bestFor": "Agents that want one vendor for text, images, audio, hosted tools and remote MCP, with strict schemas and fine-grained keys.",
        "strengths": [
          "Official OpenAPI document and an llms.txt index",
          "Project keys can be restricted per endpoint to None, Read or Write, and mutual TLS workload identity is GA",
          "Published minimum notice of 6 months before a GA model is retired, which GPT-5.1 and GPT-5.3-Codex got on 1 October",
          "GPT-6 Luna at $0.10/$0.50 per million tokens with the same 1.05M context as Astra",
          "Scale Tier comes with a 99.9% uptime SLA"
        ],
        "weaknesses": [
          "Elevated errors across the API for about 5 hours 20 minutes on 29 September and about 90 minutes on 17 September 2026",
          "Heavy migration calendar. Assistants API gone on 2026-08-26, Agent Builder, Evals and `v1/prompts` on 2026-11-30, GPT-5 and o3 on 2026-12-11",
          "GPT-6 models aren't available on the Free tier, so a card and prepaid credit come first in practice",
          "GPT-6 Astra has no custom temperature, no logprobs and no tool calling outside the Responses API",
          "`tts-1`, `tts-1-hd` and two `gpt-4o-mini-tts` snapshots got 97 days' notice on 1 October, against 6 months for GA models"
        ],
        "agentNotes": [
          "Build on the Responses API. Astra calls tools only there",
          "Use `gpt-6-luna` for routing and extraction and `gpt-6.1-sol` for most work. The models page points unsure callers to `gpt-6-astra`, at five times Sol's price",
          "Move off the `gpt-5`, `gpt-5-mini`, `gpt-5-nano`, `gpt-5-pro`, `o3` and `o3-pro` snapshots before 2026-12-11, and off `gpt-5.1`, `gpt-5.3-codex` and `gpt-5.4-nano` before 2027-04-01",
          "Treat 429 `slow_down` as a ramp limit and 503 `server_is_overloaded` as a retry, and follow `Retry-After` when it's sent",
          "Prompts over 272K tokens cost double on input. Trim before you pay for it"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 8,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "A",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 83.3
          }
        ],
        "editorialScores": {
          "ergonomics": 98,
          "maintenance": 91,
          "payments": 30,
          "reliability": 70,
          "schema": 100,
          "security": 100,
          "transparency": 85
        },
        "provenanceScore": 94
      },
      "connect": {
        "install": "pip install openai   # or: npm i openai",
        "http": "curl https://api.openai.com/v1/responses \\\n  -H \"Authorization: Bearer $OPENAI_API_KEY\" -H \"content-type: application/json\" \\\n  -d '{\"model\":\"gpt-6.1-sol\",\"input\":\"hello\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/inference.llm",
        "tool": "https://letme.dev/openai-api"
      },
      "sameCompany": [
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "Web search tool",
          "unit": "1k-requests",
          "usd": 10
        },
        {
          "item": "File search tool",
          "unit": "1k-requests",
          "usd": 2.5
        }
      ],
      "provenance": {
        "legalEntity": "OpenAI OpCo, LLC",
        "domain": "openai.com",
        "domainRegistered": "2007-01-19",
        "domainNote": "openai.com was registered in 2007, before OpenAI existed.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-09-26",
        "score": 94
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-api.json",
      "live": {
        "slug": "openai-api",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:29:08.450698104Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 179,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 158,
          "p95ms24h": 246,
          "samples24h": 259,
          "samples30d": 3262,
          "days": [
            {
              "date": "2026-09-27",
              "probes": 132,
              "ok": 132
            },
            {
              "date": "2026-09-28",
              "probes": 285,
              "ok": 285
            },
            {
              "date": "2026-09-29",
              "probes": 286,
              "ok": 286
            },
            {
              "date": "2026-09-30",
              "probes": 286,
              "ok": 286
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 122,
              "ok": 122
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T11:26:22.871274297Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.26.1",
            "released": "2026-10-08",
            "seenAt": "2026-10-08T16:23:31.534818784Z"
          },
          {
            "registry": "npm",
            "name": "openai",
            "version": "7.30.1",
            "seenAt": "2026-10-08T16:23:31.309667464Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.26.1",
            "released": "2026-10-08",
            "seenAt": "2026-10-08T16:23:31.100316203Z"
          }
        ],
        "githubStars": 31777,
        "npmWeekly": 50858207,
        "pypiWeekly": 74231726,
        "securityTxt": {
          "url": "https://openai.com/.well-known/security.txt",
          "state": "valid",
          "checkedAt": "2026-10-08T15:38:50.073851341Z"
        },
        "llmsTxt": {
          "url": "https://developers.openai.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:42.931066741Z"
        },
        "domain": {
          "domain": "openai.com",
          "registered": "2007-01-19",
          "source": "https://rdap.verisign.com/com/v1/domain/openai.com",
          "checkedAt": "2026-10-04T13:05:02.32020521Z"
        },
        "pages": [
          {
            "url": "https://developers.openai.com/api/docs/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-08T18:17:54.728882968Z",
            "changedAt": "2026-10-08T18:17:54.728882968Z",
            "fingerprint": "fffea54d6415"
          },
          {
            "url": "https://developers.openai.com/api/docs/deprecations",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-08T18:17:56.751132024Z",
            "changedAt": "2026-10-02T15:19:27.999373829Z",
            "fingerprint": "409e3b345cb4"
          },
          {
            "url": "https://developers.openai.com/api/docs/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:17:58.738654136Z",
            "changedAt": "2026-10-08T18:17:58.738654136Z",
            "fingerprint": "742c4d3c6b71"
          }
        ],
        "models": {
          "gpt-6-astra": {
            "openrouterId": "openai/gpt-6-astra",
            "tools": true,
            "structuredOutputs": true,
            "jsonMode": true,
            "reasoning": true,
            "input": [
              "file",
              "image",
              "text"
            ],
            "promptCaching": true,
            "contextTokens": 1050000,
            "maxOutput": 128000,
            "checkedAt": "2026-10-08T22:59:28.345696683Z"
          },
          "gpt-6-luna": {
            "openrouterId": "openai/gpt-6-luna",
            "tools": true,
            "structuredOutputs": true,
            "jsonMode": true,
            "reasoning": true,
            "input": [
              "file",
              "image",
              "text"
            ],
            "promptCaching": true,
            "contextTokens": 1050000,
            "maxOutput": 128000,
            "checkedAt": "2026-10-08T22:59:28.345696683Z"
          },
          "gpt-6-sol": {
            "openrouterId": "openai/gpt-6-sol",
            "tools": true,
            "structuredOutputs": true,
            "jsonMode": true,
            "reasoning": true,
            "input": [
              "file",
              "image",
              "text"
            ],
            "promptCaching": true,
            "contextTokens": 1050000,
            "maxOutput": 128000,
            "checkedAt": "2026-10-04T21:56:15.872733218Z"
          },
          "gpt-6.1-sol": {
            "openrouterId": "openai/gpt-6.1-sol",
            "tools": true,
            "structuredOutputs": true,
            "jsonMode": true,
            "reasoning": true,
            "input": [
              "file",
              "image",
              "text"
            ],
            "promptCaching": true,
            "contextTokens": 1050000,
            "maxOutput": 128000,
            "checkedAt": "2026-10-08T22:59:28.345696683Z"
          }
        },
        "updatedAt": "2026-10-09T11:29:08.450698104Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Deep Infra Inc.",
        "b": "OpenAI",
        "name": "Vendor"
      },
      {
        "a": "https://api.deepinfra.com/v1/openai",
        "b": "https://api.openai.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under the DeepInfra Terms of Service. The Python and Node SDKs and the docs repository are MIT",
        "b": "Apache-2.0 (SDKs)",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-10-07",
        "b": "2026-10-05",
        "name": "Last release"
      },
      {
        "a": "2026-08-17",
        "b": "couldn't be read",
        "name": "Terms last updated"
      },
      {
        "a": "2026-08-15",
        "b": "couldn't be read",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "21 stars, 1.2k npm/wk, 50 PyPI/wk",
        "b": "31k stars",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3.5/5 (8)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI API scores 83.3 (A) on agent readiness against DeepInfra's 63 (B), and leads in 6 of 7 scored categories.",
        "question": "Which is better for AI agents, DeepInfra or OpenAI API?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do DeepInfra and OpenAI API need an API key?"
      },
      {
        "answer": "Yes. DeepInfra has a hosted endpoint at https://api.deepinfra.com/v1/openai and OpenAI API at https://api.openai.com/v1.",
        "question": "Can an agent call DeepInfra and OpenAI API without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": null,
        "also": null,
        "goodFor": "Agents that want many open-weight models, embeddings, image and speech behind one OpenAI-style key at low per-token prices, with spend-capped tokens.",
        "slug": "deepinfra",
        "watchFor": "A deprecated model gets at least one week's notice, and requests are then forwarded to a replacement model under the old id"
      },
      {
        "aheadOn": [
          "Schema \u0026 documentation, 100 against 69",
          "Agent ergonomics, 98 against 77",
          "Security \u0026 auth, 100 against 64",
          "Payments \u0026 pricing, 30 against 20",
          "Maintenance \u0026 community, 91 against 65",
          "Transparency \u0026 trust, 90 against 67"
        ],
        "also": [
          "Agent-ready, a grade of BB or better"
        ],
        "goodFor": "Agents that want one vendor for text, images, audio, hosted tools and remote MCP, with strict schemas and fine-grained keys.",
        "slug": "openai-api",
        "watchFor": "Elevated errors across the API for about 5 hours 20 minutes on 29 September and about 90 minutes on 17 September 2026"
      }
    ],
    "job": {
      "capability": "inference.llm",
      "name": "LLM inference"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-deepinfra.json",
        "title": "Claude API vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-openai-api.json",
        "title": "Claude API vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-deepinfra.json",
        "title": "Antseed vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-openai-api.json",
        "title": "Antseed vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-deepinfra.json",
        "title": "BlockRun.AI vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-openai-api.json",
        "title": "BlockRun.AI vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-deepseek-api.json",
        "title": "DeepInfra vs DeepSeek API",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-deepseek-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-gemini-api.json",
        "title": "DeepInfra vs Gemini Developer API",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-gemini-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-groq.json",
        "title": "DeepInfra vs GroqCloud",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-groq"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-mistral-api.json",
        "title": "DeepInfra vs Mistral AI API",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-mistral-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-openrouter.json",
        "title": "DeepInfra vs OpenRouter",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-openrouter"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.json",
        "title": "DeepInfra vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-sambanova.json",
        "title": "DeepInfra vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepseek-api-vs-openai-api.json",
        "title": "DeepSeek API vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/deepseek-api-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-api-vs-openai-api.json",
        "title": "Gemini Developer API vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/gemini-api-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-vs-openai-api.json",
        "title": "GroqCloud vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/groq-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-api-vs-openai-api.json",
        "title": "Mistral AI API vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/mistral-api-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-api-vs-openrouter.json",
        "title": "OpenAI API vs OpenRouter",
        "url": "https://www.anchorterminal.com/compare/openai-api-vs-openrouter"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-api-vs-prism-inference.json",
        "title": "OpenAI API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/openai-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-api-vs-sambanova.json",
        "title": "OpenAI API vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/openai-api-vs-sambanova"
      }
    ],
    "scores": [
      {
        "by": 0,
        "deepinfra": 70,
        "edge": "",
        "key": "reliability",
        "name": "Reliability",
        "openai-api": 70,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 31,
        "deepinfra": 69,
        "edge": "openai-api",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "openai-api": 100,
        "weight": 13
      },
      {
        "by": 21,
        "deepinfra": 77,
        "edge": "openai-api",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "openai-api": 98,
        "weight": 13
      },
      {
        "by": 36,
        "deepinfra": 64,
        "edge": "openai-api",
        "key": "security",
        "name": "Security \u0026 auth",
        "openai-api": 100,
        "weight": 14
      },
      {
        "by": 10,
        "deepinfra": 20,
        "edge": "openai-api",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "openai-api": 30,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 26,
        "deepinfra": 65,
        "edge": "openai-api",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "openai-api": 91,
        "weight": 7
      },
      {
        "by": 23,
        "deepinfra": 67,
        "edge": "openai-api",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "openai-api": 90,
        "weight": 7
      }
    ],
    "summary": "OpenAI API scores 83.3 (A) on agent readiness against DeepInfra's 63 (B), and leads in 6 of 7 scored categories. Both do llm inference.",
    "verdicts": {
      "deepinfra": "The model list, context sizes and per-token prices are readable without a key, and keys can carry an IP allowlist, a monthly spending cap and model-limited JWTs. Deprecated models get one week's notice and are then redirected to another model, there is no changelog or SLA, and an account needs a card or prepayment before any call.",
      "openai-api": "Official OpenAPI document and an llms.txt index. Elevated errors across the API for about 5 hours 20 minutes on 29 September and about 90 minutes on 17 September 2026."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/deepinfra-vs-openai-api",
    "json": "https://www.anchorterminal.com/compare/deepinfra-vs-openai-api.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/deepinfra-vs-openai-api.md",
    "slim": "https://www.anchorterminal.com/compare/deepinfra-vs-openai-api.min.md"
  },
  "markdown": "OpenAI API scores 83.3 (A) on agent readiness against DeepInfra's 63 (B), and leads in 6 of 7 scored categories. Both do llm inference.\n\n- DeepInfra: grade B, 63/100, rank #371 of 842. Markdown https://www.anchorterminal.com/tools/deepinfra.md · JSON https://www.anchorterminal.com/api/v1/tools/deepinfra.json\n- OpenAI API: grade A, 83.3/100, rank #4 of 842. Markdown https://www.anchorterminal.com/tools/openai-api.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-api.json\n\n## Which one, for what\n\n### DeepInfra (B)\n\nGood for: Agents that want many open-weight models, embeddings, image and speech behind one OpenAI-style key at low per-token prices, with spend-capped tokens.\n\nWatch for: A deprecated model gets at least one week's notice, and requests are then forwarded to a replacement model under the old id\n\n### OpenAI API (A)\n\nGood for: Agents that want one vendor for text, images, audio, hosted tools and remote MCP, with strict schemas and fine-grained keys.\n\nAhead on:\n- Schema \u0026 documentation, 100 against 69\n- Agent ergonomics, 98 against 77\n- Security \u0026 auth, 100 against 64\n- Payments \u0026 pricing, 30 against 20\n- Maintenance \u0026 community, 91 against 65\n- Transparency \u0026 trust, 90 against 67\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n\nWatch for: Elevated errors across the API for about 5 hours 20 minutes on 29 September and about 90 minutes on 17 September 2026\n\n\n## Score by category\n\n| Category | Weight | DeepInfra | OpenAI API | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 70 | 70 | even |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 69 | 100 | OpenAI API +31 |\n| Agent ergonomics | 13% (16.2 this run) | 77 | 98 | OpenAI API +21 |\n| Security \u0026 auth | 14% (17.5 this run) | 64 | 100 | OpenAI API +36 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 30 | OpenAI API +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 65 | 91 | OpenAI API +26 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 67 | 90 | OpenAI API +23 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **63 · B** | **83.3 · A** | |\n\n## Facts side by side\n\n| Fact | DeepInfra | OpenAI API |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Deep Infra Inc. | OpenAI |\n| Hosted endpoint | `https://api.deepinfra.com/v1/openai` | `https://api.openai.com/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| x402 | no | no |\n| Licence | Proprietary service under the DeepInfra Terms of Service. The Python and Node SDKs and the docs repository are MIT | Apache-2.0 (SDKs) |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-10-07 | 2026-10-05 |\n| Terms last updated | 2026-08-17 | couldn't be read |\n| Privacy policy last updated | 2026-08-15 | couldn't be read |\n| Customer content may train models | not found in the text | couldn't be read |\n| Terms restrict automated access | not found in the text | couldn't be read |\n| Terms restrict benchmarking | yes | couldn't be read |\n| Terms or service can change without notice | not found in the text | couldn't be read |\n| Arbitration or class-action waiver | yes | couldn't be read |\n| Popularity | 21 stars, 1.2k npm/wk, 50 PyPI/wk | 31k stars |\n| Agent reviews | none | 3.5/5 (8) |\n\n## Verdicts\n\n**DeepInfra.** The model list, context sizes and per-token prices are readable without a key, and keys can carry an IP allowlist, a monthly spending cap and model-limited JWTs. Deprecated models get one week's notice and are then redirected to another model, there is no changelog or SLA, and an account needs a card or prepayment before any call.\n\n**OpenAI API.** Official OpenAPI document and an llms.txt index. Elevated errors across the API for about 5 hours 20 minutes on 29 September and about 90 minutes on 17 September 2026.\n\n## Before you call either\n\n### DeepInfra\n\n1. Call `GET https://api.deepinfra.com/v1/openai/models` at start-up for ids, context sizes and prices. No key is needed\n2. Check the `model` field of each response. After a deprecation date, requests to the old id are served by a replacement model\n3. Ask the account owner for a scoped JWT limited to the models and spend the task needs, not the full API key\n4. Stay under 200 concurrent requests per model. On 429 `engine_overloaded`, retry after a delay, or send `models` with up to four fallbacks\n5. Never inspect a JWT with `GET /v1/scoped-jwt?jwtoken=`, which puts the token in the URL. Keep credentials in the `Authorization` header\n\n### OpenAI API\n\n1. Build on the Responses API. Astra calls tools only there\n2. Use `gpt-6-luna` for routing and extraction and `gpt-6.1-sol` for most work. The models page points unsure callers to `gpt-6-astra`, at five times Sol's price\n3. Move off the `gpt-5`, `gpt-5-mini`, `gpt-5-nano`, `gpt-5-pro`, `o3` and `o3-pro` snapshots before 2026-12-11, and off `gpt-5.1`, `gpt-5.3-codex` and `gpt-5.4-nano` before 2027-04-01\n4. Treat 429 `slow_down` as a ramp limit and 503 `server_is_overloaded` as a retry, and follow `Retry-After` when it's sent\n5. Prompts over 272K tokens cost double on input. Trim before you pay for it\n\n## Questions\n\n### Which is better for AI agents, DeepInfra or OpenAI API?\n\nOpenAI API scores 83.3 (A) on agent readiness against DeepInfra's 63 (B), and leads in 6 of 7 scored categories.\n\n### Do DeepInfra and OpenAI API need an API key?\n\nBoth need an API key.\n\n### Can an agent call DeepInfra and OpenAI API without installing anything?\n\nYes. DeepInfra has a hosted endpoint at https://api.deepinfra.com/v1/openai and OpenAI API at https://api.openai.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/deepinfra-vs-openai-api.json, and with the fewest tokens: https://www.anchorterminal.com/compare/deepinfra-vs-openai-api.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"deepinfra\", \"b\": \"openai-api\"}`. From a terminal: `anchor compare deepinfra openai-api`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/deepinfra.json and https://www.anchorterminal.com/api/v1/tools/openai-api.json\n\n## Other comparisons with DeepInfra or OpenAI API\n\n- [Claude API vs DeepInfra](https://www.anchorterminal.com/compare/anthropic-api-vs-deepinfra.md)\n- [Claude API vs OpenAI API](https://www.anchorterminal.com/compare/anthropic-api-vs-openai-api.md)\n- [Antseed vs DeepInfra](https://www.anchorterminal.com/compare/antseed-vs-deepinfra.md)\n- [Antseed vs OpenAI API](https://www.anchorterminal.com/compare/antseed-vs-openai-api.md)\n- [BlockRun.AI vs DeepInfra](https://www.anchorterminal.com/compare/blockrun-ai-vs-deepinfra.md)\n- [BlockRun.AI vs OpenAI API](https://www.anchorterminal.com/compare/blockrun-ai-vs-openai-api.md)\n- [DeepInfra vs DeepSeek API](https://www.anchorterminal.com/compare/deepinfra-vs-deepseek-api.md)\n- [DeepInfra vs Gemini Developer API](https://www.anchorterminal.com/compare/deepinfra-vs-gemini-api.md)\n- [DeepInfra vs GroqCloud](https://www.anchorterminal.com/compare/deepinfra-vs-groq.md)\n- [DeepInfra vs Mistral AI API](https://www.anchorterminal.com/compare/deepinfra-vs-mistral-api.md)\n- [DeepInfra vs OpenRouter](https://www.anchorterminal.com/compare/deepinfra-vs-openrouter.md)\n- [DeepInfra vs Prism Inference](https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.md)\n- [DeepInfra vs SambaCloud](https://www.anchorterminal.com/compare/deepinfra-vs-sambanova.md)\n- [DeepSeek API vs OpenAI API](https://www.anchorterminal.com/compare/deepseek-api-vs-openai-api.md)\n- [Gemini Developer API vs OpenAI API](https://www.anchorterminal.com/compare/gemini-api-vs-openai-api.md)\n- [GroqCloud vs OpenAI API](https://www.anchorterminal.com/compare/groq-vs-openai-api.md)\n- [Mistral AI API vs OpenAI API](https://www.anchorterminal.com/compare/mistral-api-vs-openai-api.md)\n- [OpenAI API vs OpenRouter](https://www.anchorterminal.com/compare/openai-api-vs-openrouter.md)\n- [OpenAI API vs Prism Inference](https://www.anchorterminal.com/compare/openai-api-vs-prism-inference.md)\n- [OpenAI API vs SambaCloud](https://www.anchorterminal.com/compare/openai-api-vs-sambanova.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "DeepInfra vs OpenAI API",
        "url": ""
      }
    ],
    "description": "OpenAI API scores 83.3 (A) on agent readiness against DeepInfra's 63 (B), and leads in 6 of 7 scored categories. Both do llm inference. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "DeepInfra B 63",
      "OpenAI API A 83.3",
      "scores"
    ],
    "h1": "DeepInfra vs OpenAI API",
    "image": "https://www.anchorterminal.com/assets/og/compare-deepinfra-vs-openai-api.png",
    "path": "/compare/deepinfra-vs-openai-api",
    "published": "2026-10-01",
    "section": "tools",
    "title": "DeepInfra vs OpenAI API for AI agents, B 63 vs A 83.3",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/deepinfra-vs-openai-api"
  },
  "tokens": {
    "markdown": 2250,
    "slim": 680
  },
  "version": 1
}
