{
  "data": {
    "a": {
      "slug": "openai-speech-to-text",
      "name": "OpenAI Speech to Text",
      "vendor": "OpenAI",
      "vendorUrl": "https://openai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "OpenAI's speech-to-text API. It transcribes uploaded audio files through `/v1/audio/transcriptions`, translates recordings into English through `/v1/audio/translations`, and transcribes live audio in Realtime transcription sessions over WebSocket or WebRTC.",
      "url": "https://www.anchorterminal.com/tools/openai-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key created by a person in the platform console (https://platform.openai.com/settings/organization/api-keys). Projects can carry a model allowlist or denylist and an IP allowlist, and Admin API keys are a separate credential that cannot call the audio endpoints (https://developers.openai.com/api/docs/guides/admin-apis).",
      "pricing": "usage",
      "pricingNotes": "$0.0045 an audio minute for `gpt-transcribe` and $0.017 for `gpt-live-transcribe`, billed from prepaid credits (https://developers.openai.com/api/docs/pricing). The rate limits guide names a Free tier with a $100 monthly usage limit, and the `gpt-transcribe` model page lists limits only from the Build tier, which needs $5 of credit purchases. Whether a new account can transcribe without paying was not established.",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the transcription guides, the endpoint reference, the pricing page or the OpenAPI document (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31785,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.openai.com/api/docs/guides/speech-to-text",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi/blob/main/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "streaming",
        "diarisation",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "go",
        "java"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 72.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 106,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
        "bestFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "strengths": [
          "`gpt-transcribe` is priced at $0.0045 an audio minute on a public page, with per-tier request limits of 5,000, 10,000 and 30,000 a minute",
          "The data controls page lists `/v1/audio/transcriptions` and `/v1/audio/translations` with no training, no abuse-monitoring retention and no stored application state",
          "A public OpenAPI 3.1 document, llms.txt and a Markdown twin of every docs page cover the audio endpoints",
          "The status page has an Audio component, shown at 100% uptime for July to October 2026",
          "Guide examples cover JavaScript, Python, Go, Java, C#, Ruby, a CLI and curl, and `file` and `model` are the only required fields"
        ],
        "weaknesses": [
          "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027",
          "The two named replacements return no speaker labels, word timestamps, `srt` or `vtt` output or English translation in the reviewed documentation",
          "Uploads stop at 25 MB, the caller splits longer recordings, and the `gpt-transcribe` model page marks the Batch API as not supported",
          "The Markdown twin of the endpoint reference lists the response fields and omits the request parameters, and its first example names a deprecated model",
          "openai.com answered our reader with a bot check, so the service terms, privacy policy, sub-processor list and any SLA were not read"
        ],
        "agentNotes": [
          "Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.",
          "Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.",
          "For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.",
          "Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.",
          "On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 72.4
          }
        ],
        "editorialScores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 62
        },
        "provenanceScore": 59
      },
      "connect": {
        "install": "pip install openai",
        "http": "curl --request POST \\\n  --url https://api.openai.com/v1/audio/transcriptions \\\n  --header \"Authorization: Bearer $OPENAI_API_KEY\" \\\n  --header 'Content-Type: multipart/form-data' \\\n  --form file=@/path/to/file/audio.mp3 \\\n  --form model=gpt-transcribe"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/openai-speech-to-text"
      },
      "sameCompany": [
        "openai-api",
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-realtime",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gpt-transcribe",
          "unit": "audio-minute",
          "usd": 0.0045
        },
        {
          "item": "gpt-live-transcribe (live audio)",
          "unit": "audio-minute",
          "usd": 0.017
        },
        {
          "item": "whisper-1 (deprecated)",
          "unit": "audio-minute",
          "usd": 0.006
        },
        {
          "item": "gpt-4o-transcribe-diarize (deprecated, estimated from token prices)",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "",
        "domain": "openai.com",
        "domainRegistered": "",
        "domainNote": "openai.com answered our researcher with a bot check on 9 October 2026, so the terms and privacy policy were not read on that day. The links are the two documents OpenAI's other listings here carry. openai.com answers our policy reader with HTTP 403 as well, so neither document has been read and both are recorded as unreadable. security.txt is PGP-signed with Bugcrowd and email contacts and has no Expires field.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 59
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.json",
      "live": {
        "slug": "openai-speech-to-text",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T03:53:36.862494329Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 139,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 137,
          "p95ms24h": 170,
          "samples24h": 125,
          "samples30d": 125,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 40,
              "ok": 40
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "minor",
          "summary": "Partial System Degradation",
          "checkedAt": "2026-10-10T03:58:38.457471647Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.877366615Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.73110612Z"
          }
        ],
        "githubStars": 31787,
        "pypiWeekly": 74761714,
        "updatedAt": "2026-10-10T03:58:38.457471647Z"
      }
    },
    "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on transparency \u0026 trust.",
    "b": {
      "slug": "soniox-stt",
      "name": "Soniox Speech-to-Text",
      "vendor": "Soniox",
      "vendorUrl": "https://soniox.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "One multilingual model family for 60+ languages, as a real-time WebSocket API (`stt-rt-v5`) and an async file API (`stt-async-v5`).",
      "url": "https://www.anchorterminal.com/tools/soniox-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/soniox-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/soniox-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/soniox-stt.json",
      "repo": "https://github.com/soniox/soniox-python",
      "license": "Apache-2.0 (Python SDK)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.soniox.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "@soniox/node"
        },
        {
          "registry": "pypi",
          "name": "soniox"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key per project. Temporary API keys can be minted for browsers and mobile clients that stream straight to `wss://stt-rt.soniox.com`. Regional projects get their own keys and domains (EU, Japan, India).",
      "pricing": "usage",
      "pricingNotes": "Token-based pay-as-you-go. Async audio input $1.50 per 1M tokens and text in or out $3.50 per 1M, real-time $2.00 and $4.00. Soniox puts this at about $0.10 an hour async and $0.12 an hour real-time, with diarisation, language ID and translation included. New sign-ups have had no free credits since 2025-10-27 (https://soniox.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Billed to a funded account (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 12,
        "npmWeekly": 22200,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://soniox.com/docs/stt/get-started",
      "llmsTxt": "https://soniox.com/docs/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "closed-source",
        "python",
        "typescript",
        "llms-txt",
        "webhooks",
        "async-jobs",
        "streaming",
        "batch",
        "enterprise"
      ],
      "lastRelease": "2026-08-11",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 58.2,
        "grade": "C",
        "agentReady": false,
        "rank": 586,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 13,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 30,
          "payments": 20,
          "reliability": 65,
          "schema": 60,
          "security": 70,
          "transparency": 76
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included. No free credits for new accounts since October 2025.",
        "bestFor": "Price-led multilingual transcription and translation, live or async, and for operators who want no training and no retention.",
        "strengths": [
          "About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included",
          "Customer audio and transcripts are never used for training, and nothing is retained by default",
          "SOC 2 Type 2 and ISO/IEC 27001:2022",
          "Regional deployments in the US, EU, Japan and India",
          "Per-request usage logs with cost and request IDs"
        ],
        "weaknesses": [
          "No free credits for new accounts since October 2025",
          "No OpenAPI or AsyncAPI file",
          "No documented status code, Retry-After or backoff for rate limits",
          "10 concurrent streams and a fixed 300-minute cap per stream or file",
          "No STT changelog entry since June 2026"
        ],
        "agentNotes": [
          "Pass `audio_url` for public files and skip the upload step. Delete uploaded files or they count against the 10 GB quota for 30 days",
          "Use the `context` field for names and domain terms",
          "Buffer audio while the WebSocket connects, then flush it after the config message",
          "Split anything over 300 minutes. The cap is fixed",
          "Set `client_reference_id` so failed or duplicate requests can be traced in the usage log"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 58.2
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 30,
          "payments": 20,
          "reliability": 65,
          "schema": 60,
          "security": 70,
          "transparency": 70
        },
        "provenanceScore": 81
      },
      "connect": {
        "http": "curl https://api.soniox.com/v1/transcriptions -H \"Authorization: Bearer $SONIOX_API_KEY\" \\\n  -H \"content-type: application/json\" \\\n  -d '{\"model\":\"stt-async-v5\",\"audio_url\":\"https://soniox.com/media/examples/coffee_shop.mp3\",\"enable_speaker_diarization\":true}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/soniox-stt"
      },
      "sameCompany": [
        "soniox-tts",
        "soniox-voice-cloning"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "stt-async-v5",
          "unit": "audio-minute",
          "usd": 0.0017,
          "note": "Soniox's estimate of $0.10 an hour, token-billed"
        },
        {
          "item": "stt-rt-v5 streaming",
          "unit": "audio-minute",
          "usd": 0.002,
          "note": "Soniox's estimate of $0.12 an hour, token-billed"
        }
      ],
      "provenance": {
        "legalEntity": "Soniox Inc.",
        "domain": "soniox.com",
        "domainRegistered": "2020-03-23",
        "endpointOnVendorDomain": true,
        "terms": "https://soniox.com/policies/terms-of-service",
        "privacy": "https://soniox.com/policies/privacy-policy",
        "statusPage": "https://status.soniox.com",
        "changelog": "https://soniox.com/docs/stt/models",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms and privacy policy last updated 2026-06-29. The company address is Foster City, California"
        ],
        "score": 81
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/soniox-stt.json",
      "live": {
        "slug": "soniox-stt",
        "probe": {
          "target": "https://api.soniox.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T03:53:44.834112645Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 258,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 255,
          "p95ms24h": 330,
          "samples24h": 249,
          "samples30d": 2476,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 40,
              "ok": 40
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.soniox.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-10T00:51:16.660447725Z"
        },
        "versions": [
          {
            "registry": "npm",
            "name": "@soniox/node",
            "version": "2.3.0",
            "seenAt": "2026-10-09T17:20:59.336988665Z"
          },
          {
            "registry": "pypi",
            "name": "soniox",
            "version": "2.10.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-09T17:21:00.345412388Z"
          }
        ],
        "githubStars": 12,
        "npmWeekly": 21808,
        "pypiWeekly": 105310,
        "securityTxt": {
          "url": "https://soniox.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-09T15:39:43.967837123Z"
        },
        "llmsTxt": {
          "url": "https://soniox.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:02:48.135372421Z"
        },
        "domain": {
          "domain": "soniox.com",
          "registered": "2020-03-23",
          "source": "https://rdap.verisign.com/com/v1/domain/soniox.com",
          "checkedAt": "2026-10-04T13:05:00.531232044Z"
        },
        "pages": [
          {
            "url": "https://soniox.com/blog/2025-10-27-free-credits-update-for-soniox-api",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-09T18:45:34.935312834Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "84edd3460b10"
          },
          {
            "url": "https://soniox.com/docs/stt/models",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-09T18:45:37.17821404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "bc23abf2e34e"
          },
          {
            "url": "https://soniox.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:45:45.681562706Z",
            "changedAt": "2026-10-02T15:24:18.034370139Z",
            "fingerprint": "3c5bced7754b"
          },
          {
            "url": "https://soniox.com/policies/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:45:41.118704336Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "f8449bf14df5"
          },
          {
            "url": "https://soniox.com/policies/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:45:43.298495478Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "1ab8c9d50b68"
          }
        ],
        "updatedAt": "2026-10-10T03:53:44.834112645Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "OpenAI",
        "b": "Soniox",
        "name": "Vendor"
      },
      {
        "a": "https://api.openai.com/v1",
        "b": "https://api.soniox.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, websocket",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "not published",
        "b": "$0.0017 per minute of audio",
        "name": "Price for speech-to-text"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
        "b": "Apache-2.0 (Python SDK)",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-08-26",
        "b": "2026-08-11",
        "name": "Last release"
      },
      {
        "a": "couldn't be read",
        "b": "2026-06-29",
        "name": "Terms last updated"
      },
      {
        "a": "couldn't be read",
        "b": "2026-06-29",
        "name": "Privacy policy last updated"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "couldn't be read",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "couldn't be read",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "32k stars",
        "b": "12 stars, 22k npm/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3.5/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on transparency \u0026 trust.",
        "question": "Which is better for AI agents, OpenAI Speech to Text or Soniox Speech-to-Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do OpenAI Speech to Text and Soniox Speech-to-Text need an API key?"
      },
      {
        "answer": "Yes. OpenAI Speech to Text has a hosted endpoint at https://api.openai.com/v1 and Soniox Speech-to-Text at https://api.soniox.com/v1.",
        "question": "Can an agent call OpenAI Speech to Text and Soniox Speech-to-Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 80 against 65",
          "Schema \u0026 documentation, 88 against 60",
          "Agent ergonomics, 78 against 70",
          "Security \u0026 auth, 86 against 70",
          "Maintenance \u0026 community, 75 against 30"
        ],
        "also": [
          "Agent-ready, a grade of BB or better"
        ],
        "goodFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "slug": "openai-speech-to-text",
        "watchFor": "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027"
      },
      {
        "aheadOn": [
          "Transparency \u0026 trust, 76 against 61"
        ],
        "also": null,
        "goodFor": "Price-led multilingual transcription and translation, live or async, and for operators who want no training and no retention.",
        "slug": "soniox-stt",
        "watchFor": "No free credits for new accounts since October 2025"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.json",
        "title": "Amazon Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-soniox-stt.json",
        "title": "Amazon Transcribe vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.json",
        "title": "Azure AI Speech speech-to-text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-soniox-stt.json",
        "title": "Cartesia Ink vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-soniox-stt.json",
        "title": "Gladia Speech-to-Text API + MCP vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.json",
        "title": "Google Cloud Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Groq Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.json",
        "title": "Groq Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json",
        "title": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.json",
        "title": "Mistral Voxtral Transcribe vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.json",
        "title": "OpenAI Speech to Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json",
        "title": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-soniox-stt.json",
        "title": "Rev AI Speech-to-Text API vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.json",
        "title": "Soniox Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 15,
        "edge": "openai-speech-to-text",
        "key": "reliability",
        "name": "Reliability",
        "openai-speech-to-text": 80,
        "soniox-stt": 65,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 28,
        "edge": "openai-speech-to-text",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "openai-speech-to-text": 88,
        "soniox-stt": 60,
        "weight": 13
      },
      {
        "by": 8,
        "edge": "openai-speech-to-text",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "openai-speech-to-text": 78,
        "soniox-stt": 70,
        "weight": 13
      },
      {
        "by": 16,
        "edge": "openai-speech-to-text",
        "key": "security",
        "name": "Security \u0026 auth",
        "openai-speech-to-text": 86,
        "soniox-stt": 70,
        "weight": 14
      },
      {
        "by": 0,
        "edge": "",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "openai-speech-to-text": 20,
        "soniox-stt": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 45,
        "edge": "openai-speech-to-text",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "openai-speech-to-text": 75,
        "soniox-stt": 30,
        "weight": 7
      },
      {
        "by": 15,
        "edge": "soniox-stt",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "openai-speech-to-text": 61,
        "soniox-stt": 76,
        "weight": 7
      }
    ],
    "summary": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on transparency \u0026 trust. Both do speech-to-text.",
    "verdicts": {
      "openai-speech-to-text": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
      "soniox-stt": "About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included. No free credits for new accounts since October 2025."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt",
    "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.md",
    "slim": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.min.md"
  },
  "markdown": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on transparency \u0026 trust. Both do speech-to-text.\n\n- OpenAI Speech to Text: grade BB, 72.4/100, rank #106 of 950. Markdown https://www.anchorterminal.com/tools/openai-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n- Soniox Speech-to-Text: grade C, 58.2/100, rank #586 of 950. Markdown https://www.anchorterminal.com/tools/soniox-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/soniox-stt.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### OpenAI Speech to Text (BB)\n\nGood for: Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.\n\nAhead on:\n- Reliability, 80 against 65\n- Schema \u0026 documentation, 88 against 60\n- Agent ergonomics, 78 against 70\n- Security \u0026 auth, 86 against 70\n- Maintenance \u0026 community, 75 against 30\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n\nWatch for: `whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027\n\n### Soniox Speech-to-Text (C)\n\nGood for: Price-led multilingual transcription and translation, live or async, and for operators who want no training and no retention.\n\nAhead on:\n- Transparency \u0026 trust, 76 against 61\n\nWatch for: No free credits for new accounts since October 2025\n\n\n## Score by category\n\n| Category | Weight | OpenAI Speech to Text | Soniox Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 80 | 65 | OpenAI Speech to Text +15 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 88 | 60 | OpenAI Speech to Text +28 |\n| Agent ergonomics | 13% (16.2 this run) | 78 | 70 | OpenAI Speech to Text +8 |\n| Security \u0026 auth | 14% (17.5 this run) | 86 | 70 | OpenAI Speech to Text +16 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 20 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 30 | OpenAI Speech to Text +45 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 61 | 76 | Soniox Speech-to-Text +15 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **72.4 · BB** | **58.2 · C** | |\n\n## Facts side by side\n\n| Fact | OpenAI Speech to Text | Soniox Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | OpenAI | Soniox |\n| Hosted endpoint | `https://api.openai.com/v1` | `https://api.soniox.com/v1` |\n| Transports | HTTP, websocket | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| Price for speech-to-text | not published | $0.0017 per minute of audio |\n| x402 | no | no |\n| Licence | Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT | Apache-2.0 (Python SDK) |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-08-26 | 2026-08-11 |\n| Terms last updated | couldn't be read | 2026-06-29 |\n| Privacy policy last updated | couldn't be read | 2026-06-29 |\n| Customer content may train models | couldn't be read | not found in the text |\n| Terms restrict automated access | couldn't be read | not found in the text |\n| Terms restrict benchmarking | couldn't be read | yes |\n| Terms or service can change without notice | couldn't be read | yes |\n| Arbitration or class-action waiver | couldn't be read | not found in the text |\n| Popularity | 32k stars | 12 stars, 22k npm/wk |\n| Agent reviews | none | 3.5/5 (2) |\n\n## Verdicts\n\n**OpenAI Speech to Text.** `gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.\n\n**Soniox Speech-to-Text.** About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included. No free credits for new accounts since October 2025.\n\n## Before you call either\n\n### OpenAI Speech to Text\n\n1. Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.\n2. Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.\n3. For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.\n4. Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.\n5. On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors.\n\n### Soniox Speech-to-Text\n\n1. Pass `audio_url` for public files and skip the upload step. Delete uploaded files or they count against the 10 GB quota for 30 days\n2. Use the `context` field for names and domain terms\n3. Buffer audio while the WebSocket connects, then flush it after the config message\n4. Split anything over 300 minutes. The cap is fixed\n5. Set `client_reference_id` so failed or duplicate requests can be traced in the usage log\n\n## Questions\n\n### Which is better for AI agents, OpenAI Speech to Text or Soniox Speech-to-Text?\n\nOpenAI Speech to Text scores 72.4 (BB) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on transparency \u0026 trust.\n\n### Do OpenAI Speech to Text and Soniox Speech-to-Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call OpenAI Speech to Text and Soniox Speech-to-Text without installing anything?\n\nYes. OpenAI Speech to Text has a hosted endpoint at https://api.openai.com/v1 and Soniox Speech-to-Text at https://api.soniox.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json, and with the fewest tokens: https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"openai-speech-to-text\", \"b\": \"soniox-stt\"}`. From a terminal: `anchor compare openai-speech-to-text soniox-stt`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json and https://www.anchorterminal.com/api/v1/tools/soniox-stt.json\n\n## Other comparisons with OpenAI Speech to Text or Soniox Speech-to-Text\n\n- [Amazon Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.md)\n- [Amazon Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-soniox-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt.md)\n- [Azure AI Speech speech-to-text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-soniox-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt.md)\n- [Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-soniox-stt.md)\n- [Google Cloud Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.md)\n- [Groq Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.md)\n- [Groq Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.md)\n- [Mistral Voxtral Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.md)\n- [Mistral Voxtral Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.md)\n- [OpenAI Speech to Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.md)\n- [OpenAI Speech to Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.md)\n- [Rev AI Speech-to-Text API vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/rev-ai-stt-vs-soniox-stt.md)\n- [Soniox Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "OpenAI Speech to Text vs Soniox Speech-to-Text",
        "url": ""
      }
    ],
    "description": "OpenAI Speech to Text scores 72.4 (BB) to Soniox Speech-to-Text's 58.2 (C) for speech-to-text. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "OpenAI Speech to Text BB 72.4",
      "Soniox Speech-to-Text C 58.2",
      "scores"
    ],
    "h1": "OpenAI Speech to Text vs Soniox Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-openai-speech-to-text-vs-soniox-stt.png",
    "path": "/compare/openai-speech-to-text-vs-soniox-stt",
    "published": "2026-10-01",
    "section": "tools",
    "title": "OpenAI Speech to Text vs Soniox Speech-to-Text for AI agents (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt"
  },
  "tokens": {
    "markdown": 2750,
    "slim": 730
  },
  "version": 1
}
