{
  "data": {
    "a": {
      "slug": "assemblyai-stt",
      "name": "AssemblyAI Speech-to-Text (Universal)",
      "vendor": "AssemblyAI",
      "vendorUrl": "https://www.assemblyai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Speech-to-text APIs for recorded audio and live streams, with speaker identification, translation and redaction options.",
      "url": "https://www.anchorterminal.com/tools/assemblyai-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/assemblyai-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/assemblyai-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/assemblyai-stt.json",
      "repo": "https://github.com/AssemblyAI/assemblyai-python-sdk",
      "license": "MIT (SDKs)",
      "transports": [
        "http",
        "streamable-http"
      ],
      "remoteUrl": "https://api.assemblyai.com/v2",
      "packages": [
        {
          "registry": "npm",
          "name": "assemblyai"
        },
        {
          "registry": "pypi",
          "name": "assemblyai"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key in the `authorization` header. Temporary tokens for browser streaming. The docs MCP needs no key.",
      "pricing": "usage",
      "pricingNotes": "Free tier with no card covers up to 185 hours of pre-recorded or 333 hours of streaming. Then pay as you go per hour of audio. Universal-3.5 Pro $0.21, Universal-2 $0.15, Universal-3.6 Pro Realtime $0.45, Universal-Streaming $0.15, Sync $0.45. Diarisation $0.02 an hour on files and $0.12 on streams, keyterms $0.05 on Universal-3.5 Pro, translation $0.06. Streaming bills session time, not audio sent (https://www.assemblyai.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Card or prepaid credits only (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 213,
        "npmWeekly": 599958,
        "pypiWeekly": 738086,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://www.assemblyai.com/docs",
      "llmsTxt": "https://www.assemblyai.com/docs/llms.txt",
      "openapi": "https://www.assemblyai.com/docs/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "no-card",
        "free-tier",
        "closed-source",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "mcp",
        "streaming",
        "batch",
        "webhooks",
        "async-jobs",
        "enterprise"
      ],
      "lastRelease": "2026-09-24",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 66.8,
        "grade": "B",
        "agentReady": false,
        "rank": 272,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 11,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 80,
          "maintenance": 80,
          "payments": 40,
          "reliability": 70,
          "schema": 95,
          "security": 50,
          "transparency": 76
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": -3,
        "negativeNotes": [
          "2026-07-10, the async `speech_model` parameter began returning 400 for current model names and silently routing legacy names to the default model, and `universal-3-pro` was blocked for new accounts and accounts inactive for 7 days, all announced in the changelog the same day. We found no earlier notice (https://www.assemblyai.com/changelog). Documented, so the minimum deduction."
        ],
        "verdict": "OpenAPI 3.1 file with typed inputs and error responses on every operation. Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026.",
        "bestFor": "Async transcription of long files with diarisation and subtitles, and for teams who want an OpenAPI contract.",
        "strengths": [
          "OpenAPI 3.1 file with typed inputs and error responses on every operation",
          "Universal-2 covers 99 languages at $0.15 an hour, Universal-3.5 Pro is $0.21",
          "Transcript lists paginate with `limit`, `before_id` and `after_id`, and transcripts come back as sentences, paragraphs, SRT or VTT",
          "Free tier with no card, up to 185 hours pre-recorded",
          "TTL deletion from 1 hour and a delete endpoint for transcripts"
        ],
        "weaknesses": [
          "Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026",
          "Free accounts can't opt out of model training",
          "The HTTP rate limit returns 403 with no Retry-After",
          "Streaming bills the time the socket is open, up to a 3-hour auto-close",
          "No security.txt or published disclosure policy"
        ],
        "agentNotes": [
          "Send `{\"type\":\"Terminate\"}` to close every stream, or billing runs to the 3-hour auto-close",
          "Use `speech_models` (plural). The singular `speech_model` now returns 400 for current model names",
          "Treat a 403 on polling as the rate limit and back off with jitter, or use webhooks",
          "Fetch `/sentences` or `/paragraphs` instead of the full transcript when you only need text",
          "Opt out in Data Controls on a paid account before sending customer audio"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 66.8
          }
        ],
        "editorialScores": {
          "ergonomics": 80,
          "maintenance": 80,
          "payments": 40,
          "reliability": 70,
          "schema": 95,
          "security": 50,
          "transparency": 70
        },
        "provenanceScore": 81
      },
      "connect": {
        "http": "curl -X POST https://api.assemblyai.com/v2/transcript -H \"authorization: $ASSEMBLYAI_API_KEY\" \\\n  -H \"content-type: application/json\" \\\n  -d '{\"audio_url\":\"https://assembly.ai/wildfires.mp3\",\"language_detection\":true,\"speaker_labels\":true}'",
        "claudeCode": "claude mcp add assemblyai-docs --transport http https://assemblyai.com/docs/mcp"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/assemblyai-stt"
      },
      "sameCompany": [
        "assemblyai-voice-agent"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Universal-3.5 Pro pre-recorded",
          "unit": "audio-minute",
          "usd": 0.0035,
          "note": "published as $0.21 an hour"
        },
        {
          "item": "Universal-2 pre-recorded",
          "unit": "audio-minute",
          "usd": 0.0025,
          "note": "published as $0.15 an hour"
        },
        {
          "item": "Universal-3.6 Pro Realtime",
          "unit": "audio-minute",
          "usd": 0.0075,
          "note": "published as $0.45 an hour, billed on session time"
        },
        {
          "item": "Universal-Streaming",
          "unit": "audio-minute",
          "usd": 0.0025,
          "note": "published as $0.15 an hour, billed on session time"
        },
        {
          "item": "Sync API",
          "unit": "audio-minute",
          "usd": 0.0075,
          "note": "published as $0.45 an hour, clips up to 2 minutes"
        },
        {
          "item": "Streaming diarisation add-on",
          "unit": "audio-minute",
          "usd": 0.002,
          "note": "published as $0.12 an hour"
        }
      ],
      "provenance": {
        "legalEntity": "AssemblyAI, Inc.",
        "domain": "assemblyai.com",
        "domainRegistered": "2016-12-24",
        "endpointOnVendorDomain": true,
        "terms": "https://www.assemblyai.com/legal/terms-of-service",
        "privacy": "https://www.assemblyai.com/legal/privacy-policy",
        "statusPage": "https://status.assemblyai.com",
        "changelog": "https://www.assemblyai.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "score": 81
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/assemblyai-stt.json",
      "live": {
        "slug": "assemblyai-stt",
        "probe": {
          "target": "https://api.assemblyai.com/v2",
          "method": "get",
          "lastAt": "2026-10-10T02:54:46.789245241Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 426,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 446,
          "p95ms24h": 503,
          "samples24h": 249,
          "samples30d": 2466,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.assemblyai.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-10T02:49:55.779873472Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "AssemblyAI/assemblyai-python-sdk",
            "version": "1.6.1",
            "released": "2026-09-24",
            "seenAt": "2026-10-09T16:40:27.755394594Z"
          },
          {
            "registry": "npm",
            "name": "assemblyai",
            "version": "4.41.5",
            "seenAt": "2026-10-09T16:40:27.217227602Z"
          },
          {
            "registry": "pypi",
            "name": "assemblyai",
            "version": "1.6.1",
            "released": "2026-09-24",
            "seenAt": "2026-10-09T16:40:27.644095501Z"
          }
        ],
        "githubStars": 213,
        "npmWeekly": 533880,
        "pypiWeekly": 779517,
        "securityTxt": {
          "url": "https://assemblyai.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-09T15:40:28.750839787Z"
        },
        "llmsTxt": {
          "url": "https://www.assemblyai.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:01:27.341480839Z"
        },
        "domain": {
          "domain": "assemblyai.com",
          "registered": "2016-12-24",
          "source": "https://rdap.verisign.com/com/v1/domain/assemblyai.com",
          "checkedAt": "2026-10-04T13:07:48.940216522Z"
        },
        "pages": [
          {
            "url": "https://www.assemblyai.com/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:48:14.22020694Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "f7dbbe98a56a"
          },
          {
            "url": "https://www.assemblyai.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:48:22.265212474Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "07bb01ccd4d1"
          },
          {
            "url": "https://www.assemblyai.com/legal/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:48:18.643811376Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6671be50624e"
          },
          {
            "url": "https://www.assemblyai.com/legal/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:48:20.297585473Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "766c6ee47ebc"
          }
        ],
        "updatedAt": "2026-10-10T02:54:46.789245241Z"
      }
    },
    "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against AssemblyAI Speech-to-Text (Universal)'s 66.8 (B), and leads in 2 of 7 scored categories. AssemblyAI Speech-to-Text (Universal) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust.",
    "b": {
      "slug": "openai-speech-to-text",
      "name": "OpenAI Speech to Text",
      "vendor": "OpenAI",
      "vendorUrl": "https://openai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "OpenAI's speech-to-text API. It transcribes uploaded audio files through `/v1/audio/transcriptions`, translates recordings into English through `/v1/audio/translations`, and transcribes live audio in Realtime transcription sessions over WebSocket or WebRTC.",
      "url": "https://www.anchorterminal.com/tools/openai-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key created by a person in the platform console (https://platform.openai.com/settings/organization/api-keys). Projects can carry a model allowlist or denylist and an IP allowlist, and Admin API keys are a separate credential that cannot call the audio endpoints (https://developers.openai.com/api/docs/guides/admin-apis).",
      "pricing": "usage",
      "pricingNotes": "$0.0045 an audio minute for `gpt-transcribe` and $0.017 for `gpt-live-transcribe`, billed from prepaid credits (https://developers.openai.com/api/docs/pricing). The rate limits guide names a Free tier with a $100 monthly usage limit, and the `gpt-transcribe` model page lists limits only from the Build tier, which needs $5 of credit purchases. Whether a new account can transcribe without paying was not established.",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the transcription guides, the endpoint reference, the pricing page or the OpenAPI document (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31785,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.openai.com/api/docs/guides/speech-to-text",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi/blob/main/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "streaming",
        "diarisation",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "go",
        "java"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 72.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 106,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
        "bestFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "strengths": [
          "`gpt-transcribe` is priced at $0.0045 an audio minute on a public page, with per-tier request limits of 5,000, 10,000 and 30,000 a minute",
          "The data controls page lists `/v1/audio/transcriptions` and `/v1/audio/translations` with no training, no abuse-monitoring retention and no stored application state",
          "A public OpenAPI 3.1 document, llms.txt and a Markdown twin of every docs page cover the audio endpoints",
          "The status page has an Audio component, shown at 100% uptime for July to October 2026",
          "Guide examples cover JavaScript, Python, Go, Java, C#, Ruby, a CLI and curl, and `file` and `model` are the only required fields"
        ],
        "weaknesses": [
          "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027",
          "The two named replacements return no speaker labels, word timestamps, `srt` or `vtt` output or English translation in the reviewed documentation",
          "Uploads stop at 25 MB, the caller splits longer recordings, and the `gpt-transcribe` model page marks the Batch API as not supported",
          "The Markdown twin of the endpoint reference lists the response fields and omits the request parameters, and its first example names a deprecated model",
          "openai.com answered our reader with a bot check, so the service terms, privacy policy, sub-processor list and any SLA were not read"
        ],
        "agentNotes": [
          "Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.",
          "Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.",
          "For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.",
          "Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.",
          "On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 72.4
          }
        ],
        "editorialScores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 62
        },
        "provenanceScore": 59
      },
      "connect": {
        "install": "pip install openai",
        "http": "curl --request POST \\\n  --url https://api.openai.com/v1/audio/transcriptions \\\n  --header \"Authorization: Bearer $OPENAI_API_KEY\" \\\n  --header 'Content-Type: multipart/form-data' \\\n  --form file=@/path/to/file/audio.mp3 \\\n  --form model=gpt-transcribe"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/openai-speech-to-text"
      },
      "sameCompany": [
        "openai-api",
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-realtime",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gpt-transcribe",
          "unit": "audio-minute",
          "usd": 0.0045
        },
        {
          "item": "gpt-live-transcribe (live audio)",
          "unit": "audio-minute",
          "usd": 0.017
        },
        {
          "item": "whisper-1 (deprecated)",
          "unit": "audio-minute",
          "usd": 0.006
        },
        {
          "item": "gpt-4o-transcribe-diarize (deprecated, estimated from token prices)",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "",
        "domain": "openai.com",
        "domainRegistered": "",
        "domainNote": "openai.com answered our researcher with a bot check on 9 October 2026, so the terms and privacy policy were not read on that day. The links are the two documents OpenAI's other listings here carry. openai.com answers our policy reader with HTTP 403 as well, so neither document has been read and both are recorded as unreadable. security.txt is PGP-signed with Bugcrowd and email contacts and has no Expires field.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 59
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.json",
      "live": {
        "slug": "openai-speech-to-text",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T02:55:04.59384396Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 120,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 138,
          "p95ms24h": 172,
          "samples24h": 115,
          "samples30d": 115,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "minor",
          "summary": "Partial System Degradation",
          "checkedAt": "2026-10-10T02:50:38.239143757Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.877366615Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.73110612Z"
          }
        ],
        "githubStars": 31787,
        "pypiWeekly": 74761714,
        "updatedAt": "2026-10-10T02:55:04.59384396Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "AssemblyAI",
        "b": "OpenAI",
        "name": "Vendor"
      },
      {
        "a": "https://api.assemblyai.com/v2",
        "b": "https://api.openai.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, Streamable HTTP",
        "b": "HTTP, websocket",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT (SDKs)",
        "b": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-24",
        "b": "2026-08-26",
        "name": "Last release"
      },
      {
        "a": "2026-07-01",
        "b": "couldn't be read",
        "name": "Terms last updated"
      },
      {
        "a": "2026-05-26",
        "b": "couldn't be read",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes, with an opt-out",
        "b": "couldn't be read",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "213 stars, 600k npm/wk, 738k PyPI/wk",
        "b": "32k stars",
        "name": "Popularity"
      },
      {
        "a": "3.5/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against AssemblyAI Speech-to-Text (Universal)'s 66.8 (B), and leads in 2 of 7 scored categories. AssemblyAI Speech-to-Text (Universal) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust.",
        "question": "Which is better for AI agents, AssemblyAI Speech-to-Text (Universal) or OpenAI Speech to Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do AssemblyAI Speech-to-Text (Universal) and OpenAI Speech to Text need an API key?"
      },
      {
        "answer": "Yes. AssemblyAI Speech-to-Text (Universal) has a hosted endpoint at https://api.assemblyai.com/v2 and OpenAI Speech to Text at https://api.openai.com/v1.",
        "question": "Can an agent call AssemblyAI Speech-to-Text (Universal) and OpenAI Speech to Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Schema \u0026 documentation, 95 against 88",
          "Payments \u0026 pricing, 40 against 20",
          "Maintenance \u0026 community, 80 against 75",
          "Transparency \u0026 trust, 76 against 61"
        ],
        "also": [
          "Free to start without a card"
        ],
        "goodFor": "Async transcription of long files with diarisation and subtitles, and for teams who want an OpenAPI contract.",
        "slug": "assemblyai-stt",
        "watchFor": "Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026"
      },
      {
        "aheadOn": [
          "Reliability, 80 against 70",
          "Security \u0026 auth, 86 against 50"
        ],
        "also": [
          "Agent-ready, a grade of BB or better",
          "No incidents deducted, where AssemblyAI Speech-to-Text (Universal) loses 3 points for them"
        ],
        "goodFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "slug": "openai-speech-to-text",
        "watchFor": "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-assemblyai-stt.json",
        "title": "Amazon Transcribe vs AssemblyAI Speech-to-Text (Universal)",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-assemblyai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.json",
        "title": "Amazon Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Azure AI Speech speech-to-text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-cartesia-ink-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Cartesia Ink",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-cartesia-ink-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-deepgram-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-elevenlabs-scribe.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-gladia-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Gladia Speech-to-Text API + MCP",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-gladia-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-rev-ai-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Groq Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json",
        "title": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.json",
        "title": "OpenAI Speech to Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json",
        "title": "OpenAI Speech to Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json",
        "title": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "assemblyai-stt": 70,
        "by": 10,
        "edge": "openai-speech-to-text",
        "key": "reliability",
        "name": "Reliability",
        "openai-speech-to-text": 80,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "assemblyai-stt": 95,
        "by": 7,
        "edge": "assemblyai-stt",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "openai-speech-to-text": 88,
        "weight": 13
      },
      {
        "assemblyai-stt": 80,
        "by": 2,
        "edge": "assemblyai-stt",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "openai-speech-to-text": 78,
        "weight": 13
      },
      {
        "assemblyai-stt": 50,
        "by": 36,
        "edge": "openai-speech-to-text",
        "key": "security",
        "name": "Security \u0026 auth",
        "openai-speech-to-text": 86,
        "weight": 14
      },
      {
        "assemblyai-stt": 40,
        "by": 20,
        "edge": "assemblyai-stt",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "openai-speech-to-text": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "assemblyai-stt": 80,
        "by": 5,
        "edge": "assemblyai-stt",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "openai-speech-to-text": 75,
        "weight": 7
      },
      {
        "assemblyai-stt": 76,
        "by": 15,
        "edge": "assemblyai-stt",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "openai-speech-to-text": 61,
        "weight": 7
      }
    ],
    "summary": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against AssemblyAI Speech-to-Text (Universal)'s 66.8 (B), and leads in 2 of 7 scored categories. AssemblyAI Speech-to-Text (Universal) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust. Both do speech-to-text.",
    "verdicts": {
      "assemblyai-stt": "OpenAPI 3.1 file with typed inputs and error responses on every operation. Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026.",
      "openai-speech-to-text": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.min.md"
  },
  "markdown": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against AssemblyAI Speech-to-Text (Universal)'s 66.8 (B), and leads in 2 of 7 scored categories. AssemblyAI Speech-to-Text (Universal) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust. Both do speech-to-text.\n\n- AssemblyAI Speech-to-Text (Universal): grade B, 66.8/100, rank #272 of 950. Markdown https://www.anchorterminal.com/tools/assemblyai-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/assemblyai-stt.json\n- OpenAI Speech to Text: grade BB, 72.4/100, rank #106 of 950. Markdown https://www.anchorterminal.com/tools/openai-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### AssemblyAI Speech-to-Text (Universal) (B)\n\nGood for: Async transcription of long files with diarisation and subtitles, and for teams who want an OpenAPI contract.\n\nAhead on:\n- Schema \u0026 documentation, 95 against 88\n- Payments \u0026 pricing, 40 against 20\n- Maintenance \u0026 community, 80 against 75\n- Transparency \u0026 trust, 76 against 61\n\nAlso in its favour:\n- Free to start without a card\n\nWatch for: Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026\n\n### OpenAI Speech to Text (BB)\n\nGood for: Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.\n\nAhead on:\n- Reliability, 80 against 70\n- Security \u0026 auth, 86 against 50\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n- No incidents deducted, where AssemblyAI Speech-to-Text (Universal) loses 3 points for them\n\nWatch for: `whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027\n\n\n## Score by category\n\n| Category | Weight | AssemblyAI Speech-to-Text (Universal) | OpenAI Speech to Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 70 | 80 | OpenAI Speech to Text +10 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 95 | 88 | AssemblyAI Speech-to-Text (Universal) +7 |\n| Agent ergonomics | 13% (16.2 this run) | 80 | 78 | AssemblyAI Speech-to-Text (Universal) +2 |\n| Security \u0026 auth | 14% (17.5 this run) | 50 | 86 | OpenAI Speech to Text +36 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | AssemblyAI Speech-to-Text (Universal) +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 80 | 75 | AssemblyAI Speech-to-Text (Universal) +5 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 76 | 61 | AssemblyAI Speech-to-Text (Universal) +15 |\n| Negative events | ≤15 | -3 | 0 | |\n| **Total** | | **66.8 · B** | **72.4 · BB** | |\n\n## Facts side by side\n\n| Fact | AssemblyAI Speech-to-Text (Universal) | OpenAI Speech to Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | AssemblyAI | OpenAI |\n| Hosted endpoint | `https://api.assemblyai.com/v2` | `https://api.openai.com/v1` |\n| Transports | HTTP, Streamable HTTP | HTTP, websocket |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| x402 | no | no |\n| Licence | MIT (SDKs) | Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-24 | 2026-08-26 |\n| Terms last updated | 2026-07-01 | couldn't be read |\n| Privacy policy last updated | 2026-05-26 | couldn't be read |\n| Customer content may train models | yes, with an opt-out | couldn't be read |\n| Terms restrict automated access | yes | couldn't be read |\n| Terms restrict benchmarking | yes | couldn't be read |\n| Terms or service can change without notice | not found in the text | couldn't be read |\n| Arbitration or class-action waiver | not found in the text | couldn't be read |\n| Popularity | 213 stars, 600k npm/wk, 738k PyPI/wk | 32k stars |\n| Agent reviews | 3.5/5 (2) | none |\n\n## Verdicts\n\n**AssemblyAI Speech-to-Text (Universal).** OpenAPI 3.1 file with typed inputs and error responses on every operation. Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026.\n\n**OpenAI Speech to Text.** `gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.\n\n## Before you call either\n\n### AssemblyAI Speech-to-Text (Universal)\n\n1. Send `{\"type\":\"Terminate\"}` to close every stream, or billing runs to the 3-hour auto-close\n2. Use `speech_models` (plural). The singular `speech_model` now returns 400 for current model names\n3. Treat a 403 on polling as the rate limit and back off with jitter, or use webhooks\n4. Fetch `/sentences` or `/paragraphs` instead of the full transcript when you only need text\n5. Opt out in Data Controls on a paid account before sending customer audio\n\n### OpenAI Speech to Text\n\n1. Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.\n2. Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.\n3. For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.\n4. Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.\n5. On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors.\n\n## Questions\n\n### Which is better for AI agents, AssemblyAI Speech-to-Text (Universal) or OpenAI Speech to Text?\n\nOpenAI Speech to Text scores 72.4 (BB) on agent readiness against AssemblyAI Speech-to-Text (Universal)'s 66.8 (B), and leads in 2 of 7 scored categories. AssemblyAI Speech-to-Text (Universal) leads on schema \u0026 documentation, payments \u0026 pricing, maintenance \u0026 community and transparency \u0026 trust.\n\n### Do AssemblyAI Speech-to-Text (Universal) and OpenAI Speech to Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call AssemblyAI Speech-to-Text (Universal) and OpenAI Speech to Text without installing anything?\n\nYes. AssemblyAI Speech-to-Text (Universal) has a hosted endpoint at https://api.assemblyai.com/v2 and OpenAI Speech to Text at https://api.openai.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json, and with the fewest tokens: https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"assemblyai-stt\", \"b\": \"openai-speech-to-text\"}`. From a terminal: `anchor compare assemblyai-stt openai-speech-to-text`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/assemblyai-stt.json and https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n\n## Other comparisons with AssemblyAI Speech-to-Text (Universal) or OpenAI Speech to Text\n\n- [Amazon Transcribe vs AssemblyAI Speech-to-Text (Universal)](https://www.anchorterminal.com/compare/amazon-transcribe-vs-assemblyai-stt.md)\n- [Amazon Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Azure AI Speech speech-to-text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Cartesia Ink](https://www.anchorterminal.com/compare/assemblyai-stt-vs-cartesia-ink-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/assemblyai-stt-vs-deepgram-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/assemblyai-stt-vs-elevenlabs-scribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/assemblyai-stt-vs-gladia-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/assemblyai-stt-vs-rev-ai-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt.md)\n- [Azure AI Speech speech-to-text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md)\n- [Groq Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.md)\n- [Mistral Voxtral Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.md)\n- [OpenAI Speech to Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.md)\n- [OpenAI Speech to Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.md)\n- [OpenAI Speech to Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
        "url": ""
      }
    ],
    "description": "OpenAI Speech to Text scores 72.4 (BB) to AssemblyAI Speech-to-Text's 66.8 (B) for speech-to-text. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "AssemblyAI Speech-to-Text (Universal) B 66.8",
      "OpenAI Speech to Text BB 72.4",
      "scores"
    ],
    "h1": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-assemblyai-stt-vs-openai-speech-to-text.png",
    "path": "/compare/assemblyai-stt-vs-openai-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "AssemblyAI Speech-to-Text vs OpenAI Speech to Text (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text"
  },
  "tokens": {
    "markdown": 2950,
    "slim": 830
  },
  "version": 1
}
