{
  "data": {
    "a": {
      "slug": "openai-speech-to-text",
      "name": "OpenAI Speech to Text",
      "vendor": "OpenAI",
      "vendorUrl": "https://openai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "OpenAI's speech-to-text API. It transcribes uploaded audio files through `/v1/audio/transcriptions`, translates recordings into English through `/v1/audio/translations`, and transcribes live audio in Realtime transcription sessions over WebSocket or WebRTC.",
      "url": "https://www.anchorterminal.com/tools/openai-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key created by a person in the platform console (https://platform.openai.com/settings/organization/api-keys). Projects can carry a model allowlist or denylist and an IP allowlist, and Admin API keys are a separate credential that cannot call the audio endpoints (https://developers.openai.com/api/docs/guides/admin-apis).",
      "pricing": "usage",
      "pricingNotes": "$0.0045 an audio minute for `gpt-transcribe` and $0.017 for `gpt-live-transcribe`, billed from prepaid credits (https://developers.openai.com/api/docs/pricing). The rate limits guide names a Free tier with a $100 monthly usage limit, and the `gpt-transcribe` model page lists limits only from the Build tier, which needs $5 of credit purchases. Whether a new account can transcribe without paying was not established.",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the transcription guides, the endpoint reference, the pricing page or the OpenAPI document (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31785,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.openai.com/api/docs/guides/speech-to-text",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi/blob/main/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "streaming",
        "diarisation",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "go",
        "java"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 72.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 106,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
        "bestFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "strengths": [
          "`gpt-transcribe` is priced at $0.0045 an audio minute on a public page, with per-tier request limits of 5,000, 10,000 and 30,000 a minute",
          "The data controls page lists `/v1/audio/transcriptions` and `/v1/audio/translations` with no training, no abuse-monitoring retention and no stored application state",
          "A public OpenAPI 3.1 document, llms.txt and a Markdown twin of every docs page cover the audio endpoints",
          "The status page has an Audio component, shown at 100% uptime for July to October 2026",
          "Guide examples cover JavaScript, Python, Go, Java, C#, Ruby, a CLI and curl, and `file` and `model` are the only required fields"
        ],
        "weaknesses": [
          "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027",
          "The two named replacements return no speaker labels, word timestamps, `srt` or `vtt` output or English translation in the reviewed documentation",
          "Uploads stop at 25 MB, the caller splits longer recordings, and the `gpt-transcribe` model page marks the Batch API as not supported",
          "The Markdown twin of the endpoint reference lists the response fields and omits the request parameters, and its first example names a deprecated model",
          "openai.com answered our reader with a bot check, so the service terms, privacy policy, sub-processor list and any SLA were not read"
        ],
        "agentNotes": [
          "Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.",
          "Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.",
          "For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.",
          "Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.",
          "On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 72.4
          }
        ],
        "editorialScores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 62
        },
        "provenanceScore": 59
      },
      "connect": {
        "install": "pip install openai",
        "http": "curl --request POST \\\n  --url https://api.openai.com/v1/audio/transcriptions \\\n  --header \"Authorization: Bearer $OPENAI_API_KEY\" \\\n  --header 'Content-Type: multipart/form-data' \\\n  --form file=@/path/to/file/audio.mp3 \\\n  --form model=gpt-transcribe"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/openai-speech-to-text"
      },
      "sameCompany": [
        "openai-api",
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-realtime",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gpt-transcribe",
          "unit": "audio-minute",
          "usd": 0.0045
        },
        {
          "item": "gpt-live-transcribe (live audio)",
          "unit": "audio-minute",
          "usd": 0.017
        },
        {
          "item": "whisper-1 (deprecated)",
          "unit": "audio-minute",
          "usd": 0.006
        },
        {
          "item": "gpt-4o-transcribe-diarize (deprecated, estimated from token prices)",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "",
        "domain": "openai.com",
        "domainRegistered": "",
        "domainNote": "openai.com answered our researcher with a bot check on 9 October 2026, so the terms and privacy policy were not read on that day. The links are the two documents OpenAI's other listings here carry. openai.com answers our policy reader with HTTP 403 as well, so neither document has been read and both are recorded as unreadable. security.txt is PGP-signed with Bugcrowd and email contacts and has no Expires field.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 59
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.json",
      "live": {
        "slug": "openai-speech-to-text",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T02:55:04.59384396Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 120,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 138,
          "p95ms24h": 172,
          "samples24h": 115,
          "samples30d": 115,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "minor",
          "summary": "Partial System Degradation",
          "checkedAt": "2026-10-10T02:50:38.239143757Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.877366615Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.73110612Z"
          }
        ],
        "githubStars": 31787,
        "pypiWeekly": 74761714,
        "updatedAt": "2026-10-10T02:55:04.59384396Z"
      }
    },
    "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on payments \u0026 pricing and transparency \u0026 trust.",
    "b": {
      "slug": "speechmatics-stt",
      "name": "Speechmatics Speech-to-Text",
      "vendor": "Speechmatics",
      "vendorUrl": "https://www.speechmatics.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Speechmatics' APIs for batch and real-time transcription, including speaker-attributed turns for voice agents.",
      "url": "https://www.anchorterminal.com/tools/speechmatics-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/speechmatics-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/speechmatics-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/speechmatics-stt.json",
      "repo": "https://github.com/speechmatics/speechmatics-python-sdk",
      "license": "MIT (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://eu1.asr.api.speechmatics.com/v2",
      "packages": [
        {
          "registry": "npm",
          "name": "@speechmatics/batch-client"
        },
        {
          "registry": "npm",
          "name": "@speechmatics/real-time-client"
        },
        {
          "registry": "pypi",
          "name": "speechmatics-batch"
        },
        {
          "registry": "pypi",
          "name": "speechmatics-rt"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key as a Bearer token. Short-lived JWTs for browser and realtime clients, passed as `?jwt=` on the WebSocket URL. A separate management token covers project and key administration.",
      "pricing": "usage",
      "pricingNotes": "$100 free credit with no card, then pay as you go per hour of audio, billed to the second. Batch Melia 1 $0.13, Batch Standard $0.24, Batch Enhanced $0.40, Realtime Standard $0.24, Realtime Enhanced $0.43, Linden 1 (Agent STT) $0.16 (was $0.21). Translation adds $0.65 an hour. 20 per cent off usage over 500 hours a month per model, and 33 per cent off if you opt in to model training (https://www.speechmatics.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Card or prepaid credits only (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 20,
        "npmWeekly": 58050,
        "pypiWeekly": 48085,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.speechmatics.com",
      "llmsTxt": "https://docs.speechmatics.com/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "no-card",
        "closed-source",
        "python",
        "typescript",
        "llms-txt",
        "streaming",
        "batch",
        "webhooks",
        "async-jobs",
        "enterprise",
        "self-hosted"
      ],
      "lastRelease": "2026-09-22",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 67.1,
        "grade": "B",
        "agentReady": false,
        "rank": 262,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 10,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 80,
          "maintenance": 75,
          "payments": 40,
          "reliability": 70,
          "schema": 65,
          "security": 65,
          "transparency": 75
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Training is opt-in and real-time audio is not stored. Enhanced transcription costs $0.40 to $0.43 an hour.",
        "bestFor": "Regulated or privacy-sensitive audio, multilingual batch with Melia 1, and voice agents that want speaker-attributed turns.",
        "strengths": [
          "Training is opt-in only, and realtime audio isn't stored",
          "ISO/IEC 27001:2022 and SOC 2 Type II",
          "Transcripts fetched as plain text, JSON or SRT",
          "Ten dated changelog entries in September 2026",
          "$100 credit with no card"
        ],
        "weaknesses": [
          "Enhanced costs $0.40 to $0.43 an hour, above most rivals",
          "No OpenAPI or AsyncAPI file linked from the docs",
          "429s carry a reason but no Retry-After or backoff guidance",
          "Realtime JWTs travel in the WebSocket URL",
          "Free plan allows 2 realtime sessions"
        ],
        "agentNotes": [
          "Set `\"model\": \"enhanced\"` explicitly. The default is `standard`",
          "Use notifications instead of polling. Polling waits 5 seconds by default since the 23 September 2026 change, and `wait=0` turns that off",
          "Fetch batch transcripts within 7 days. After that the API returns 404 `expired`",
          "Pass a `fetch_data` URL for files over 1 GB",
          "Use `/v2/agent` with `linden-1` for live agents instead of the plain realtime path"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 67.1
          }
        ],
        "editorialScores": {
          "ergonomics": 80,
          "maintenance": 75,
          "payments": 40,
          "reliability": 70,
          "schema": 65,
          "security": 65,
          "transparency": 65
        },
        "provenanceScore": 85
      },
      "connect": {
        "http": "curl -X POST https://eu1.asr.api.speechmatics.com/v2/jobs/ -H \"Authorization: Bearer $SPEECHMATICS_API_KEY\" \\\n  -F data_file=@call.wav \\\n  -F config='{\"type\":\"transcription\",\"transcription_config\":{\"language\":\"en\",\"model\":\"enhanced\",\"diarization\":\"speaker\"}}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/speechmatics-stt"
      },
      "area": "voice",
      "unitPrices": [
        {
          "item": "Batch Melia 1",
          "unit": "audio-minute",
          "usd": 0.0022,
          "note": "published as $0.13 an hour"
        },
        {
          "item": "Batch Standard",
          "unit": "audio-minute",
          "usd": 0.004,
          "note": "published as $0.24 an hour"
        },
        {
          "item": "Batch Enhanced",
          "unit": "audio-minute",
          "usd": 0.0067,
          "note": "published as $0.40 an hour"
        },
        {
          "item": "Realtime Standard",
          "unit": "audio-minute",
          "usd": 0.004,
          "note": "published as $0.24 an hour"
        },
        {
          "item": "Realtime Enhanced",
          "unit": "audio-minute",
          "usd": 0.0072,
          "note": "published as $0.43 an hour"
        },
        {
          "item": "Linden 1 Agent STT",
          "unit": "audio-minute",
          "usd": 0.0027,
          "note": "published as $0.16 an hour, reduced from $0.21"
        },
        {
          "item": "Translation add-on",
          "unit": "audio-minute",
          "usd": 0.0108,
          "note": "published as $0.65 an hour"
        }
      ],
      "provenance": {
        "legalEntity": "Cantab Research Ltd",
        "domain": "speechmatics.com",
        "domainRegistered": "2006-05-10",
        "endpointOnVendorDomain": true,
        "terms": "https://www.speechmatics.com/legal/terms-of-service",
        "privacy": "https://www.speechmatics.com/legal/privacy-policy",
        "statusPage": "https://status.speechmatics.com",
        "changelog": "https://speechmatics.featurebase.app/en/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Trades as Speechmatics, company number 05697423 in England and Wales. US customers contract with Speechmatics (USA) Inc., a Delaware company"
        ],
        "score": 85
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/speechmatics-stt.json",
      "live": {
        "slug": "speechmatics-stt",
        "probe": {
          "target": "https://eu1.asr.api.speechmatics.com/v2",
          "method": "get",
          "lastAt": "2026-10-10T02:55:11.457847081Z",
          "lastOk": true,
          "lastStatus": 200,
          "lastMs": 54,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 73,
          "p95ms24h": 185,
          "samples24h": 249,
          "samples30d": 2466,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.speechmatics.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-10T02:50:47.488319972Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "speechmatics/speechmatics-python-sdk",
            "version": "rt/v1.2.1",
            "released": "2026-10-05",
            "seenAt": "2026-10-09T17:21:19.573197737Z"
          },
          {
            "registry": "npm",
            "name": "@speechmatics/batch-client",
            "version": "5.4.2",
            "seenAt": "2026-10-09T17:21:15.241924919Z"
          },
          {
            "registry": "npm",
            "name": "@speechmatics/real-time-client",
            "version": "8.5.1",
            "seenAt": "2026-10-09T17:21:16.117352264Z"
          },
          {
            "registry": "pypi",
            "name": "speechmatics-batch",
            "version": "1.1.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-09T17:21:17.489135668Z"
          },
          {
            "registry": "pypi",
            "name": "speechmatics-rt",
            "version": "1.2.1",
            "released": "2026-10-05",
            "seenAt": "2026-10-09T17:21:17.680413296Z"
          }
        ],
        "githubStars": 20,
        "npmWeekly": 21097,
        "pypiWeekly": 16435,
        "securityTxt": {
          "url": "https://speechmatics.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-09T15:39:54.389957384Z"
        },
        "llmsTxt": {
          "url": "https://docs.speechmatics.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:02:49.844693739Z"
        },
        "domain": {
          "domain": "speechmatics.com",
          "registered": "2006-05-10",
          "source": "https://rdap.verisign.com/com/v1/domain/speechmatics.com",
          "checkedAt": "2026-10-04T13:05:04.524690122Z"
        },
        "pages": [
          {
            "url": "https://speechmatics.featurebase.app/en/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:45:44.23207324Z",
            "changedAt": "2026-10-09T18:45:44.23207324Z",
            "fingerprint": "acc13ccb4f33"
          },
          {
            "url": "https://www.speechmatics.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:54:33.566370746Z",
            "changedAt": "2026-10-06T16:17:05.219546965Z",
            "fingerprint": "01ed4b71fd7e"
          },
          {
            "url": "https://www.speechmatics.com/legal/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:54:29.268881414Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "3af3999aaeb0"
          },
          {
            "url": "https://www.speechmatics.com/legal/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:54:31.42363639Z",
            "changedAt": "2026-10-02T15:28:24.437639084Z",
            "fingerprint": "33b19ec6fdcb"
          }
        ],
        "updatedAt": "2026-10-10T02:55:11.457847081Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "OpenAI",
        "b": "Speechmatics",
        "name": "Vendor"
      },
      {
        "a": "https://api.openai.com/v1",
        "b": "https://eu1.asr.api.speechmatics.com/v2",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, websocket",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "not published",
        "b": "$0.0027 per minute of audio",
        "name": "Price for speech-to-text"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
        "b": "MIT (SDKs)",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-08-26",
        "b": "2026-09-22",
        "name": "Last release"
      },
      {
        "a": "couldn't be read",
        "b": "no date given",
        "name": "Terms last updated"
      },
      {
        "a": "couldn't be read",
        "b": "2026-05-27",
        "name": "Privacy policy last updated"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "couldn't be read",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "32k stars",
        "b": "20 stars, 58k npm/wk, 48k PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3.5/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on payments \u0026 pricing and transparency \u0026 trust.",
        "question": "Which is better for AI agents, OpenAI Speech to Text or Speechmatics Speech-to-Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do OpenAI Speech to Text and Speechmatics Speech-to-Text need an API key?"
      },
      {
        "answer": "Yes. OpenAI Speech to Text has a hosted endpoint at https://api.openai.com/v1 and Speechmatics Speech-to-Text at https://eu1.asr.api.speechmatics.com/v2.",
        "question": "Can an agent call OpenAI Speech to Text and Speechmatics Speech-to-Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 80 against 70",
          "Schema \u0026 documentation, 88 against 65",
          "Security \u0026 auth, 86 against 65"
        ],
        "also": [
          "Agent-ready, a grade of BB or better"
        ],
        "goodFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "slug": "openai-speech-to-text",
        "watchFor": "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027"
      },
      {
        "aheadOn": [
          "Payments \u0026 pricing, 40 against 20",
          "Transparency \u0026 trust, 75 against 61"
        ],
        "also": [
          "Free to start without a card"
        ],
        "goodFor": "Regulated or privacy-sensitive audio, multilingual batch with Melia 1, and voice agents that want speaker-attributed turns.",
        "slug": "speechmatics-stt",
        "watchFor": "Enhanced costs $0.40 to $0.43 an hour, above most rivals"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.json",
        "title": "Amazon Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-speechmatics-stt.json",
        "title": "Amazon Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.json",
        "title": "Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-speechmatics-stt.json",
        "title": "Cartesia Ink vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-speechmatics-stt.json",
        "title": "Gladia Speech-to-Text API + MCP vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.json",
        "title": "Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Groq Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.json",
        "title": "Groq Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json",
        "title": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.json",
        "title": "Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.json",
        "title": "OpenAI Speech to Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json",
        "title": "OpenAI Speech to Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-speechmatics-stt.json",
        "title": "Rev AI Speech-to-Text API vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.json",
        "title": "Soniox Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 10,
        "edge": "openai-speech-to-text",
        "key": "reliability",
        "name": "Reliability",
        "openai-speech-to-text": 80,
        "speechmatics-stt": 70,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 23,
        "edge": "openai-speech-to-text",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "openai-speech-to-text": 88,
        "speechmatics-stt": 65,
        "weight": 13
      },
      {
        "by": 2,
        "edge": "speechmatics-stt",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "openai-speech-to-text": 78,
        "speechmatics-stt": 80,
        "weight": 13
      },
      {
        "by": 21,
        "edge": "openai-speech-to-text",
        "key": "security",
        "name": "Security \u0026 auth",
        "openai-speech-to-text": 86,
        "speechmatics-stt": 65,
        "weight": 14
      },
      {
        "by": 20,
        "edge": "speechmatics-stt",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "openai-speech-to-text": 20,
        "speechmatics-stt": 40,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 0,
        "edge": "",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "openai-speech-to-text": 75,
        "speechmatics-stt": 75,
        "weight": 7
      },
      {
        "by": 14,
        "edge": "speechmatics-stt",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "openai-speech-to-text": 61,
        "speechmatics-stt": 75,
        "weight": 7
      }
    ],
    "summary": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on payments \u0026 pricing and transparency \u0026 trust. Both do speech-to-text.",
    "verdicts": {
      "openai-speech-to-text": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
      "speechmatics-stt": "Training is opt-in and real-time audio is not stored. Enhanced transcription costs $0.40 to $0.43 an hour."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt",
    "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.md",
    "slim": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.min.md"
  },
  "markdown": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on payments \u0026 pricing and transparency \u0026 trust. Both do speech-to-text.\n\n- OpenAI Speech to Text: grade BB, 72.4/100, rank #106 of 950. Markdown https://www.anchorterminal.com/tools/openai-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n- Speechmatics Speech-to-Text: grade B, 67.1/100, rank #262 of 950. Markdown https://www.anchorterminal.com/tools/speechmatics-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/speechmatics-stt.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### OpenAI Speech to Text (BB)\n\nGood for: Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.\n\nAhead on:\n- Reliability, 80 against 70\n- Schema \u0026 documentation, 88 against 65\n- Security \u0026 auth, 86 against 65\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n\nWatch for: `whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027\n\n### Speechmatics Speech-to-Text (B)\n\nGood for: Regulated or privacy-sensitive audio, multilingual batch with Melia 1, and voice agents that want speaker-attributed turns.\n\nAhead on:\n- Payments \u0026 pricing, 40 against 20\n- Transparency \u0026 trust, 75 against 61\n\nAlso in its favour:\n- Free to start without a card\n\nWatch for: Enhanced costs $0.40 to $0.43 an hour, above most rivals\n\n\n## Score by category\n\n| Category | Weight | OpenAI Speech to Text | Speechmatics Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 80 | 70 | OpenAI Speech to Text +10 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 88 | 65 | OpenAI Speech to Text +23 |\n| Agent ergonomics | 13% (16.2 this run) | 78 | 80 | Speechmatics Speech-to-Text +2 |\n| Security \u0026 auth | 14% (17.5 this run) | 86 | 65 | OpenAI Speech to Text +21 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 40 | Speechmatics Speech-to-Text +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 75 | even |\n| Transparency \u0026 trust | 7% (8.8 this run) | 61 | 75 | Speechmatics Speech-to-Text +14 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **72.4 · BB** | **67.1 · B** | |\n\n## Facts side by side\n\n| Fact | OpenAI Speech to Text | Speechmatics Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | OpenAI | Speechmatics |\n| Hosted endpoint | `https://api.openai.com/v1` | `https://eu1.asr.api.speechmatics.com/v2` |\n| Transports | HTTP, websocket | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| Price for speech-to-text | not published | $0.0027 per minute of audio |\n| x402 | no | no |\n| Licence | Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT | MIT (SDKs) |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-08-26 | 2026-09-22 |\n| Terms last updated | couldn't be read | no date given |\n| Privacy policy last updated | couldn't be read | 2026-05-27 |\n| Customer content may train models | couldn't be read | not found in the text |\n| Terms restrict automated access | couldn't be read | not found in the text |\n| Terms restrict benchmarking | couldn't be read | yes |\n| Terms or service can change without notice | couldn't be read | not found in the text |\n| Arbitration or class-action waiver | couldn't be read | not found in the text |\n| Popularity | 32k stars | 20 stars, 58k npm/wk, 48k PyPI/wk |\n| Agent reviews | none | 3.5/5 (2) |\n\n## Verdicts\n\n**OpenAI Speech to Text.** `gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.\n\n**Speechmatics Speech-to-Text.** Training is opt-in and real-time audio is not stored. Enhanced transcription costs $0.40 to $0.43 an hour.\n\n## Before you call either\n\n### OpenAI Speech to Text\n\n1. Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.\n2. Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.\n3. For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.\n4. Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.\n5. On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors.\n\n### Speechmatics Speech-to-Text\n\n1. Set `\"model\": \"enhanced\"` explicitly. The default is `standard`\n2. Use notifications instead of polling. Polling waits 5 seconds by default since the 23 September 2026 change, and `wait=0` turns that off\n3. Fetch batch transcripts within 7 days. After that the API returns 404 `expired`\n4. Pass a `fetch_data` URL for files over 1 GB\n5. Use `/v2/agent` with `linden-1` for live agents instead of the plain realtime path\n\n## Questions\n\n### Which is better for AI agents, OpenAI Speech to Text or Speechmatics Speech-to-Text?\n\nOpenAI Speech to Text scores 72.4 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on payments \u0026 pricing and transparency \u0026 trust.\n\n### Do OpenAI Speech to Text and Speechmatics Speech-to-Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call OpenAI Speech to Text and Speechmatics Speech-to-Text without installing anything?\n\nYes. OpenAI Speech to Text has a hosted endpoint at https://api.openai.com/v1 and Speechmatics Speech-to-Text at https://eu1.asr.api.speechmatics.com/v2.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json, and with the fewest tokens: https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"openai-speech-to-text\", \"b\": \"speechmatics-stt\"}`. From a terminal: `anchor compare openai-speech-to-text speechmatics-stt`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json and https://www.anchorterminal.com/api/v1/tools/speechmatics-stt.json\n\n## Other comparisons with OpenAI Speech to Text or Speechmatics Speech-to-Text\n\n- [Amazon Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.md)\n- [Amazon Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-speechmatics-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt.md)\n- [Azure AI Speech speech-to-text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-speechmatics-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt.md)\n- [Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-speechmatics-stt.md)\n- [Google Cloud Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.md)\n- [Groq Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.md)\n- [Groq Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.md)\n- [Mistral Voxtral Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.md)\n- [Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.md)\n- [OpenAI Speech to Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.md)\n- [OpenAI Speech to Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.md)\n- [Rev AI Speech-to-Text API vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/rev-ai-stt-vs-speechmatics-stt.md)\n- [Soniox Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
        "url": ""
      }
    ],
    "description": "OpenAI Speech to Text scores 72.4 (BB) to Speechmatics Speech-to-Text's 67.1 (B) for speech-to-text. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "OpenAI Speech to Text BB 72.4",
      "Speechmatics Speech-to-Text B 67.1",
      "scores"
    ],
    "h1": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-openai-speech-to-text-vs-speechmatics-stt.png",
    "path": "/compare/openai-speech-to-text-vs-speechmatics-stt",
    "published": "2026-10-01",
    "section": "tools",
    "title": "OpenAI Speech to Text vs Speechmatics Speech-to-Text (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt"
  },
  "tokens": {
    "markdown": 2850,
    "slim": 780
  },
  "version": 1
}
