{
  "data": {
    "a": {
      "slug": "deepgram-stt",
      "name": "Deepgram Speech-to-Text (Nova-3, Flux)",
      "vendor": "Deepgram",
      "vendorUrl": "https://deepgram.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Deepgram's speech-to-text API for recorded audio and live streams, including turn detection for voice agents.",
      "url": "https://www.anchorterminal.com/tools/deepgram-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/deepgram-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/deepgram-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/deepgram-stt.json",
      "repo": "https://github.com/deepgram/deepgram-python-sdk",
      "license": "MIT (SDKs)",
      "transports": [
        "http",
        "streamable-http",
        "stdio",
        "sse"
      ],
      "remoteUrl": "https://api.deepgram.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "@deepgram/sdk"
        },
        {
          "registry": "pypi",
          "name": "deepgram-sdk"
        },
        {
          "registry": "pypi",
          "name": "deepctl"
        }
      ],
      "auth": "api-key",
      "authNotes": "`Authorization: Token \u003ckey\u003e` header on REST and WebSocket calls. Short-lived JWTs (30-second TTL) from the token endpoint for browsers. The `dg` CLI MCP server uses `dg login` credentials or `DEEPGRAM_API_KEY`. The docs MCP needs no key.",
      "pricing": "usage",
      "pricingNotes": "$200 free credit with no card, then pay as you go, or Growth from $4,000 a year prepaid for up to 20 per cent off. Nova-3 pre-recorded $0.0043 a minute (multilingual $0.0052). Streaming Nova-3 is on a promotional $0.0048 a minute (regular $0.0077), multilingual $0.0058 (regular $0.0092). Flux English $0.0065 promotional (regular $0.0077), Flux Multilingual $0.0078. Streaming diarisation adds $0.0020 a minute, redaction $0.0020 and keyterm prompting $0.0013 (https://deepgram.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Card or prepaid credits only (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 468,
        "npmWeekly": 1123798,
        "pypiWeekly": 805026,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://developers.deepgram.com/docs/models-languages-overview",
      "llmsTxt": "https://developers.deepgram.com/llms.txt",
      "openapi": "https://developers.deepgram.com/openapi.json",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "hosted",
        "no-card",
        "closed-source",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "mcp",
        "streaming",
        "batch",
        "webhooks",
        "enterprise",
        "self-hosted"
      ],
      "lastRelease": "2026-09-29",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 70.3,
        "grade": "BB",
        "agentReady": true,
        "rank": 157,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 5,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 40,
          "reliability": 65,
          "schema": 95,
          "security": 65,
          "transparency": 72
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Flux streams with model-level end-of-turn detection, so a voice agent needs no separate VAD. Training on audio is the default and the opt-out is a per-request flag.",
        "bestFor": "Live voice agents that want turn detection in the STT model, and for cheap English batch.",
        "strengths": [
          "Flux streams with model-level end-of-turn detection, so a voice agent needs no separate VAD",
          "Nova-3 pre-recorded at $0.0043 a minute with diarisation included",
          "OpenAPI and AsyncAPI files, an llms.txt and SDKs in six languages",
          "Keys can carry a role and an expiry date",
          "$200 free credit with no card"
        ],
        "weaknesses": [
          "Training on audio is the default and the opt-out is a per-request flag",
          "Two incidents over 2 hours in July 2026, on Flux streaming and batch",
          "No SLA published for self-serve plans",
          "The privacy policy dates from October 2021 and doesn't mention the Model Improvement Program",
          "Streaming prices are promotional and may rise to the regular rate"
        ],
        "agentNotes": [
          "Add `mip_opt_out=true` to every request that carries customer audio",
          "Use Flux (`flux-general-en`) on `/v2/listen` for live agents and Nova-3 for files",
          "Back off exponentially on 429. The concurrency limit is per project",
          "Pass `callback` for long files so the request doesn't hit the 10-minute processing timeout",
          "Mint keys with an expiry for short-lived jobs"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 70.3
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 40,
          "reliability": 65,
          "schema": 95,
          "security": 65,
          "transparency": 60
        },
        "provenanceScore": 84
      },
      "connect": {
        "http": "curl -X POST \"https://api.deepgram.com/v1/listen?model=nova-3\u0026smart_format=true\u0026diarize=true\" \\\n  -H \"Authorization: Token $DEEPGRAM_API_KEY\" -H \"content-type: application/json\" \\\n  -d '{\"url\":\"https://dpgr.am/spacewalk.wav\"}'",
        "claudeCode": "claude mcp add deepgram-docs --transport http https://api.dx.deepgram.com/kapa/mcp",
        "config": {
          "mcpServers": {
            "deepgram": {
              "args": [
                "mcp"
              ],
              "command": "dg",
              "env": {
                "DEEPGRAM_API_KEY": "${DEEPGRAM_API_KEY}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/deepgram-stt"
      },
      "sameCompany": [
        "deepgram-tts",
        "deepgram-voice-agent"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Nova-3 pre-recorded",
          "unit": "audio-minute",
          "usd": 0.0043,
          "note": "pay as you go, diarisation included"
        },
        {
          "item": "Nova-3 Multilingual pre-recorded",
          "unit": "audio-minute",
          "usd": 0.0052
        },
        {
          "item": "Nova-3 streaming",
          "unit": "audio-minute",
          "usd": 0.0048,
          "note": "promotional, regular $0.0077"
        },
        {
          "item": "Nova-3 Multilingual streaming",
          "unit": "audio-minute",
          "usd": 0.0058,
          "note": "promotional, regular $0.0092"
        },
        {
          "item": "Flux English streaming",
          "unit": "audio-minute",
          "usd": 0.0065,
          "note": "promotional, regular $0.0077"
        },
        {
          "item": "Flux Multilingual streaming",
          "unit": "audio-minute",
          "usd": 0.0078
        },
        {
          "item": "Streaming diarisation add-on",
          "unit": "audio-minute",
          "usd": 0.002
        }
      ],
      "provenance": {
        "legalEntity": "Deepgram, Inc.",
        "domain": "deepgram.com",
        "domainRegistered": "2016-01-28",
        "endpointOnVendorDomain": true,
        "terms": "https://deepgram.com/terms",
        "privacy": "https://deepgram.com/privacy",
        "statusPage": "https://status.deepgram.com",
        "changelog": "https://developers.deepgram.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "score": 84
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/deepgram-stt.json",
      "live": {
        "slug": "deepgram-stt",
        "probe": {
          "target": "https://api.deepgram.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T02:54:52.701474148Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 408,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 411,
          "p95ms24h": 554,
          "samples24h": 249,
          "samples30d": 2466,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.deepgram.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-10T02:50:11.928707251Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "deepgram/deepgram-python-sdk",
            "version": "v7.12.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-09T16:49:03.100560025Z"
          },
          {
            "registry": "npm",
            "name": "@deepgram/sdk",
            "version": "5.14.0",
            "seenAt": "2026-10-09T16:49:00.559079153Z"
          },
          {
            "registry": "pypi",
            "name": "deepctl",
            "version": "0.3.2",
            "released": "2026-10-05",
            "seenAt": "2026-10-09T16:49:01.214360086Z"
          },
          {
            "registry": "pypi",
            "name": "deepgram-sdk",
            "version": "7.12.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-09T16:49:01.018044441Z"
          }
        ],
        "githubStars": 469,
        "npmWeekly": 913554,
        "pypiWeekly": 787283,
        "securityTxt": {
          "url": "https://deepgram.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-09T15:40:20.111519289Z"
        },
        "llmsTxt": {
          "url": "https://developers.deepgram.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:01:47.341051037Z"
        },
        "domain": {
          "domain": "deepgram.com",
          "registered": "2016-01-28",
          "source": "https://rdap.verisign.com/com/v1/domain/deepgram.com",
          "checkedAt": "2026-10-04T13:03:28.939824686Z"
        },
        "pages": [
          {
            "url": "https://developers.deepgram.com/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:35:34.851707978Z",
            "changedAt": "2026-10-09T18:35:34.851707978Z",
            "fingerprint": "e5a0ca096c74"
          },
          {
            "url": "https://deepgram.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:34:45.883547785Z",
            "changedAt": "2026-10-09T18:34:45.883547785Z",
            "fingerprint": "2ac9babde538"
          },
          {
            "url": "https://deepgram.com/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:34:48.804992391Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "89846510fbe6"
          },
          {
            "url": "https://deepgram.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:34:50.091846483Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "7042d107fca1"
          }
        ],
        "updatedAt": "2026-10-10T02:54:52.701474148Z"
      }
    },
    "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Deepgram Speech-to-Text (Nova-3, Flux)'s 70.3 (BB), and leads in 3 of 7 scored categories. Deepgram Speech-to-Text (Nova-3, Flux) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust.",
    "b": {
      "slug": "openai-speech-to-text",
      "name": "OpenAI Speech to Text",
      "vendor": "OpenAI",
      "vendorUrl": "https://openai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "OpenAI's speech-to-text API. It transcribes uploaded audio files through `/v1/audio/transcriptions`, translates recordings into English through `/v1/audio/translations`, and transcribes live audio in Realtime transcription sessions over WebSocket or WebRTC.",
      "url": "https://www.anchorterminal.com/tools/openai-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key created by a person in the platform console (https://platform.openai.com/settings/organization/api-keys). Projects can carry a model allowlist or denylist and an IP allowlist, and Admin API keys are a separate credential that cannot call the audio endpoints (https://developers.openai.com/api/docs/guides/admin-apis).",
      "pricing": "usage",
      "pricingNotes": "$0.0045 an audio minute for `gpt-transcribe` and $0.017 for `gpt-live-transcribe`, billed from prepaid credits (https://developers.openai.com/api/docs/pricing). The rate limits guide names a Free tier with a $100 monthly usage limit, and the `gpt-transcribe` model page lists limits only from the Build tier, which needs $5 of credit purchases. Whether a new account can transcribe without paying was not established.",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the transcription guides, the endpoint reference, the pricing page or the OpenAPI document (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31785,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.openai.com/api/docs/guides/speech-to-text",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi/blob/main/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "streaming",
        "diarisation",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "go",
        "java"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 72.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 106,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
        "bestFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "strengths": [
          "`gpt-transcribe` is priced at $0.0045 an audio minute on a public page, with per-tier request limits of 5,000, 10,000 and 30,000 a minute",
          "The data controls page lists `/v1/audio/transcriptions` and `/v1/audio/translations` with no training, no abuse-monitoring retention and no stored application state",
          "A public OpenAPI 3.1 document, llms.txt and a Markdown twin of every docs page cover the audio endpoints",
          "The status page has an Audio component, shown at 100% uptime for July to October 2026",
          "Guide examples cover JavaScript, Python, Go, Java, C#, Ruby, a CLI and curl, and `file` and `model` are the only required fields"
        ],
        "weaknesses": [
          "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027",
          "The two named replacements return no speaker labels, word timestamps, `srt` or `vtt` output or English translation in the reviewed documentation",
          "Uploads stop at 25 MB, the caller splits longer recordings, and the `gpt-transcribe` model page marks the Batch API as not supported",
          "The Markdown twin of the endpoint reference lists the response fields and omits the request parameters, and its first example names a deprecated model",
          "openai.com answered our reader with a bot check, so the service terms, privacy policy, sub-processor list and any SLA were not read"
        ],
        "agentNotes": [
          "Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.",
          "Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.",
          "For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.",
          "Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.",
          "On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 72.4
          }
        ],
        "editorialScores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 62
        },
        "provenanceScore": 59
      },
      "connect": {
        "install": "pip install openai",
        "http": "curl --request POST \\\n  --url https://api.openai.com/v1/audio/transcriptions \\\n  --header \"Authorization: Bearer $OPENAI_API_KEY\" \\\n  --header 'Content-Type: multipart/form-data' \\\n  --form file=@/path/to/file/audio.mp3 \\\n  --form model=gpt-transcribe"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/openai-speech-to-text"
      },
      "sameCompany": [
        "openai-api",
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-realtime",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gpt-transcribe",
          "unit": "audio-minute",
          "usd": 0.0045
        },
        {
          "item": "gpt-live-transcribe (live audio)",
          "unit": "audio-minute",
          "usd": 0.017
        },
        {
          "item": "whisper-1 (deprecated)",
          "unit": "audio-minute",
          "usd": 0.006
        },
        {
          "item": "gpt-4o-transcribe-diarize (deprecated, estimated from token prices)",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "",
        "domain": "openai.com",
        "domainRegistered": "",
        "domainNote": "openai.com answered our researcher with a bot check on 9 October 2026, so the terms and privacy policy were not read on that day. The links are the two documents OpenAI's other listings here carry. openai.com answers our policy reader with HTTP 403 as well, so neither document has been read and both are recorded as unreadable. security.txt is PGP-signed with Bugcrowd and email contacts and has no Expires field.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 59
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.json",
      "live": {
        "slug": "openai-speech-to-text",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T02:55:04.59384396Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 120,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 138,
          "p95ms24h": 172,
          "samples24h": 115,
          "samples30d": 115,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "minor",
          "summary": "Partial System Degradation",
          "checkedAt": "2026-10-10T02:50:38.239143757Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.877366615Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.73110612Z"
          }
        ],
        "githubStars": 31787,
        "pypiWeekly": 74761714,
        "updatedAt": "2026-10-10T02:55:04.59384396Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Deepgram",
        "b": "OpenAI",
        "name": "Vendor"
      },
      {
        "a": "https://api.deepgram.com/v1",
        "b": "https://api.openai.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, Streamable HTTP, stdio, SSE (legacy)",
        "b": "HTTP, websocket",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT (SDKs)",
        "b": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-29",
        "b": "2026-08-26",
        "name": "Last release"
      },
      {
        "a": "2026-08-06",
        "b": "couldn't be read",
        "name": "Terms last updated"
      },
      {
        "a": "2021-10-26",
        "b": "couldn't be read",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes, with an opt-out",
        "b": "couldn't be read",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "468 stars, 1.1M npm/wk, 805k PyPI/wk",
        "b": "32k stars",
        "name": "Popularity"
      },
      {
        "a": "3.5/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Deepgram Speech-to-Text (Nova-3, Flux)'s 70.3 (BB), and leads in 3 of 7 scored categories. Deepgram Speech-to-Text (Nova-3, Flux) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust.",
        "question": "Which is better for AI agents, Deepgram Speech-to-Text (Nova-3, Flux) or OpenAI Speech to Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Deepgram Speech-to-Text (Nova-3, Flux) and OpenAI Speech to Text need an API key?"
      },
      {
        "answer": "Yes. Deepgram Speech-to-Text (Nova-3, Flux) has a hosted endpoint at https://api.deepgram.com/v1 and OpenAI Speech to Text at https://api.openai.com/v1.",
        "question": "Can an agent call Deepgram Speech-to-Text (Nova-3, Flux) and OpenAI Speech to Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Schema \u0026 documentation, 95 against 88",
          "Payments \u0026 pricing, 40 against 20",
          "Maintenance \u0026 community, 80 against 75",
          "Transparency \u0026 trust, 72 against 61"
        ],
        "also": [
          "Runs on your own machine",
          "Free to start without a card"
        ],
        "goodFor": "Live voice agents that want turn detection in the STT model, and for cheap English batch.",
        "slug": "deepgram-stt",
        "watchFor": "Training on audio is the default and the opt-out is a per-request flag"
      },
      {
        "aheadOn": [
          "Reliability, 80 against 65",
          "Security \u0026 auth, 86 against 65"
        ],
        "also": null,
        "goodFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "slug": "openai-speech-to-text",
        "watchFor": "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-deepgram-stt.json",
        "title": "Amazon Transcribe vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.json",
        "title": "Amazon Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-deepgram-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt.json",
        "title": "Azure AI Speech speech-to-text vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-deepgram-stt.json",
        "title": "Cartesia Ink vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-elevenlabs-scribe.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-gladia-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Gladia Speech-to-Text API + MCP",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-gladia-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-rev-ai-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Groq Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json",
        "title": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.json",
        "title": "OpenAI Speech to Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json",
        "title": "OpenAI Speech to Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json",
        "title": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 15,
        "deepgram-stt": 65,
        "edge": "openai-speech-to-text",
        "key": "reliability",
        "name": "Reliability",
        "openai-speech-to-text": 80,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 7,
        "deepgram-stt": 95,
        "edge": "deepgram-stt",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "openai-speech-to-text": 88,
        "weight": 13
      },
      {
        "by": 3,
        "deepgram-stt": 75,
        "edge": "openai-speech-to-text",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "openai-speech-to-text": 78,
        "weight": 13
      },
      {
        "by": 21,
        "deepgram-stt": 65,
        "edge": "openai-speech-to-text",
        "key": "security",
        "name": "Security \u0026 auth",
        "openai-speech-to-text": 86,
        "weight": 14
      },
      {
        "by": 20,
        "deepgram-stt": 40,
        "edge": "deepgram-stt",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "openai-speech-to-text": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 5,
        "deepgram-stt": 80,
        "edge": "deepgram-stt",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "openai-speech-to-text": 75,
        "weight": 7
      },
      {
        "by": 11,
        "deepgram-stt": 72,
        "edge": "deepgram-stt",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "openai-speech-to-text": 61,
        "weight": 7
      }
    ],
    "summary": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Deepgram Speech-to-Text (Nova-3, Flux)'s 70.3 (BB), and leads in 3 of 7 scored categories. Deepgram Speech-to-Text (Nova-3, Flux) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust. Both do speech-to-text.",
    "verdicts": {
      "deepgram-stt": "Flux streams with model-level end-of-turn detection, so a voice agent needs no separate VAD. Training on audio is the default and the opt-out is a per-request flag.",
      "openai-speech-to-text": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.min.md"
  },
  "markdown": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Deepgram Speech-to-Text (Nova-3, Flux)'s 70.3 (BB), and leads in 3 of 7 scored categories. Deepgram Speech-to-Text (Nova-3, Flux) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust. Both do speech-to-text.\n\n- Deepgram Speech-to-Text (Nova-3, Flux): grade BB, 70.3/100, rank #157 of 950. Markdown https://www.anchorterminal.com/tools/deepgram-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/deepgram-stt.json\n- OpenAI Speech to Text: grade BB, 72.4/100, rank #106 of 950. Markdown https://www.anchorterminal.com/tools/openai-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### Deepgram Speech-to-Text (Nova-3, Flux) (BB)\n\nGood for: Live voice agents that want turn detection in the STT model, and for cheap English batch.\n\nAhead on:\n- Schema \u0026 documentation, 95 against 88\n- Payments \u0026 pricing, 40 against 20\n- Maintenance \u0026 community, 80 against 75\n- Transparency \u0026 trust, 72 against 61\n\nAlso in its favour:\n- Runs on your own machine\n- Free to start without a card\n\nWatch for: Training on audio is the default and the opt-out is a per-request flag\n\n### OpenAI Speech to Text (BB)\n\nGood for: Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.\n\nAhead on:\n- Reliability, 80 against 65\n- Security \u0026 auth, 86 against 65\n\nWatch for: `whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027\n\n\n## Score by category\n\n| Category | Weight | Deepgram Speech-to-Text (Nova-3, Flux) | OpenAI Speech to Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 65 | 80 | OpenAI Speech to Text +15 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 95 | 88 | Deepgram Speech-to-Text (Nova-3, Flux) +7 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 78 | OpenAI Speech to Text +3 |\n| Security \u0026 auth | 14% (17.5 this run) | 65 | 86 | OpenAI Speech to Text +21 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | Deepgram Speech-to-Text (Nova-3, Flux) +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 80 | 75 | Deepgram Speech-to-Text (Nova-3, Flux) +5 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 72 | 61 | Deepgram Speech-to-Text (Nova-3, Flux) +11 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **70.3 · BB** | **72.4 · BB** | |\n\n## Facts side by side\n\n| Fact | Deepgram Speech-to-Text (Nova-3, Flux) | OpenAI Speech to Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Deepgram | OpenAI |\n| Hosted endpoint | `https://api.deepgram.com/v1` | `https://api.openai.com/v1` |\n| Transports | HTTP, Streamable HTTP, stdio, SSE (legacy) | HTTP, websocket |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| x402 | no | no |\n| Licence | MIT (SDKs) | Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-29 | 2026-08-26 |\n| Terms last updated | 2026-08-06 | couldn't be read |\n| Privacy policy last updated | 2021-10-26 | couldn't be read |\n| Customer content may train models | yes, with an opt-out | couldn't be read |\n| Terms restrict automated access | not found in the text | couldn't be read |\n| Terms restrict benchmarking | yes | couldn't be read |\n| Terms or service can change without notice | yes | couldn't be read |\n| Arbitration or class-action waiver | yes | couldn't be read |\n| Popularity | 468 stars, 1.1M npm/wk, 805k PyPI/wk | 32k stars |\n| Agent reviews | 3.5/5 (2) | none |\n\n## Verdicts\n\n**Deepgram Speech-to-Text (Nova-3, Flux).** Flux streams with model-level end-of-turn detection, so a voice agent needs no separate VAD. Training on audio is the default and the opt-out is a per-request flag.\n\n**OpenAI Speech to Text.** `gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.\n\n## Before you call either\n\n### Deepgram Speech-to-Text (Nova-3, Flux)\n\n1. Add `mip_opt_out=true` to every request that carries customer audio\n2. Use Flux (`flux-general-en`) on `/v2/listen` for live agents and Nova-3 for files\n3. Back off exponentially on 429. The concurrency limit is per project\n4. Pass `callback` for long files so the request doesn't hit the 10-minute processing timeout\n5. Mint keys with an expiry for short-lived jobs\n\n### OpenAI Speech to Text\n\n1. Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.\n2. Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.\n3. For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.\n4. Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.\n5. On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors.\n\n## Questions\n\n### Which is better for AI agents, Deepgram Speech-to-Text (Nova-3, Flux) or OpenAI Speech to Text?\n\nOpenAI Speech to Text scores 72.4 (BB) on agent readiness against Deepgram Speech-to-Text (Nova-3, Flux)'s 70.3 (BB), and leads in 3 of 7 scored categories. Deepgram Speech-to-Text (Nova-3, Flux) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust.\n\n### Do Deepgram Speech-to-Text (Nova-3, Flux) and OpenAI Speech to Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call Deepgram Speech-to-Text (Nova-3, Flux) and OpenAI Speech to Text without installing anything?\n\nYes. Deepgram Speech-to-Text (Nova-3, Flux) has a hosted endpoint at https://api.deepgram.com/v1 and OpenAI Speech to Text at https://api.openai.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json, and with the fewest tokens: https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"deepgram-stt\", \"b\": \"openai-speech-to-text\"}`. From a terminal: `anchor compare deepgram-stt openai-speech-to-text`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/deepgram-stt.json and https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n\n## Other comparisons with Deepgram Speech-to-Text (Nova-3, Flux) or OpenAI Speech to Text\n\n- [Amazon Transcribe vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/amazon-transcribe-vs-deepgram-stt.md)\n- [Amazon Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/assemblyai-stt-vs-deepgram-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt.md)\n- [Azure AI Speech speech-to-text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-deepgram-stt.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/deepgram-stt-vs-elevenlabs-scribe.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/deepgram-stt-vs-gladia-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/deepgram-stt-vs-rev-ai-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md)\n- [Groq Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.md)\n- [Mistral Voxtral Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.md)\n- [OpenAI Speech to Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.md)\n- [OpenAI Speech to Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.md)\n- [OpenAI Speech to Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
        "url": ""
      }
    ],
    "description": "OpenAI Speech to Text scores 72.4 (BB) to Deepgram Speech-to-Text's 70.3 (BB) for speech-to-text. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "Deepgram Speech-to-Text (Nova-3, Flux) BB 70.3",
      "OpenAI Speech to Text BB 72.4",
      "scores"
    ],
    "h1": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-deepgram-stt-vs-openai-speech-to-text.png",
    "path": "/compare/deepgram-stt-vs-openai-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Deepgram Speech-to-Text vs OpenAI Speech to Text for AI agents (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text"
  },
  "tokens": {
    "markdown": 2900,
    "slim": 780
  },
  "version": 1
}
