{
  "data": {
    "a": {
      "slug": "mistral-voxtral-transcribe",
      "name": "Mistral Voxtral Transcribe",
      "vendor": "Mistral AI",
      "vendorUrl": "https://mistral.ai",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Mistral AI's speech-to-text API. Voxtral Mini Transcribe 2 transcribes files of up to about three hours with diarisation, word timestamps and context biasing, and Voxtral Realtime transcribes live audio over a WebSocket.",
      "url": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe",
      "markdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json",
      "repo": "https://github.com/mistralai/client-python",
      "license": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.mistral.ai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "mistralai"
        },
        {
          "registry": "npm",
          "name": "@mistralai/mistralai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from Studio, bound to one workspace, with an optional expiry date. Browser realtime clients use `rt_` tokens minted at `POST /v1/client/sessions`, valid about 900 seconds for one model (https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth).",
      "pricing": "freemium",
      "pricingNotes": "$0.003 per audio minute for Voxtral Mini Transcribe 2 and $0.006 for Voxtral Mini Transcribe Realtime (model cards on docs.mistral.ai, https://mistral.ai/pricing/api/). Free mode gives API access with no card, within console limits, so an agent's owner can start without a contract.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 773,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.mistral.ai/studio/audio/speech_to_text",
      "llmsTxt": "https://docs.mistral.ai/llms.txt",
      "openapi": "https://docs.mistral.ai/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.streaming",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "eu",
        "free-tier",
        "openapi",
        "llms-txt",
        "python",
        "typescript",
        "streaming",
        "batch",
        "open-weights"
      ],
      "lastRelease": "2026-02-04",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 64.1,
        "grade": "B",
        "agentReady": false,
        "rank": 365,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 12,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 82
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": -3,
        "negativeNotes": [
          "2026-05-12. Compromised `mistralai` 2.4.6 on PyPI ran a credential-harvesting script on import for about three hours, and three `@mistralai/mistralai` versions on npm were also replaced. Mistral published advisory MAI-2026-002, removed the packages and closed its investigation on 14 May. Fixed and documented, so 3 points. https://docs.mistral.ai/resources/security-advisories/MAI-2026-002"
        ],
        "verdict": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.",
        "bestFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "strengths": [
          "Published prices of $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime",
          "One request with `model` and a `file`, `file_url` or `file_id`, for audio of up to about three hours",
          "Browser clients use `rt_` tokens that last about 900 seconds and are limited to one model",
          "`/v1/audio/transcriptions` is on the zero data retention list for paid plans, and the Realtime weights are Apache-2.0 on Hugging Face",
          "OpenAPI document, llms.txt and Markdown guides, with a six-month retirement notice for GA models"
        ],
        "weaknesses": [
          "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel",
          "`timestamp_granularities` cannot be used with `language`, and the realtime model cannot diarise",
          "13 languages, and context biasing is tuned for English with other languages described as experimental",
          "The compromised `mistralai` 2.4.6 on PyPI harvested credentials on import for three hours on 12 May 2026",
          "No SLA found, and status.mistral.ai answered our requests with a bot check on 8 October 2026",
          "The Studio MCP server has no transcription tool"
        ],
        "agentNotes": [
          "Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`",
          "Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible",
          "Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`",
          "Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move",
          "Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 64.1
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 73
        },
        "provenanceScore": 91
      },
      "connect": {
        "install": "pip install mistralai   # realtime: pip install \"mistralai[realtime]\"   # or: npm i @mistralai/mistralai",
        "http": "curl --location 'https://api.mistral.ai/v1/audio/transcriptions' \\\n  --header \"x-api-key: $MISTRAL_API_KEY\" \\\n  --form 'file_url=\"https://docs.mistral.ai/audio/obama.mp3\"' \\\n  --form 'model=\"voxtral-mini-latest\"'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/mistral-voxtral-transcribe"
      },
      "sameCompany": [
        "mistral-api",
        "mistral-embeddings",
        "mistral-moderation",
        "mistral-ocr"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Voxtral Mini Transcribe 2 (batch)",
          "unit": "audio-minute",
          "usd": 0.003
        },
        {
          "item": "Voxtral Mini Transcribe Realtime",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "Mistral AI (RCS Paris 952 418 325)",
        "domain": "mistral.ai",
        "domainRegistered": "2019-05-15",
        "endpointOnVendorDomain": true,
        "terms": "https://legal.mistral.ai/terms/commercial-terms-of-service",
        "privacy": "https://legal.mistral.ai/terms/privacy-policy",
        "statusPage": "https://status.mistral.ai",
        "changelog": "https://docs.mistral.ai/resources/changelogs",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 91
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.json",
      "live": {
        "slug": "mistral-voxtral-transcribe",
        "probe": {
          "target": "https://api.mistral.ai/v1",
          "method": "get",
          "lastAt": "2026-10-10T03:07:13.477589933Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 45,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 48,
          "p95ms24h": 84,
          "samples24h": 201,
          "samples30d": 201,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 169,
              "ok": 169
            },
            {
              "date": "2026-10-10",
              "probes": 32,
              "ok": 32
            }
          ]
        },
        "versions": [
          {
            "registry": "github",
            "name": "mistralai/client-python",
            "version": "v3.2.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:06:39.739433733Z"
          },
          {
            "registry": "npm",
            "name": "@mistralai/mistralai",
            "version": "2.7.0",
            "seenAt": "2026-10-09T17:06:39.382083947Z"
          },
          {
            "registry": "pypi",
            "name": "mistralai",
            "version": "3.2.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:06:39.265218804Z"
          }
        ],
        "githubStars": 773,
        "npmWeekly": 8338907,
        "pypiWeekly": 3268403,
        "securityTxt": {
          "url": "https://mistral.ai/.well-known/security.txt",
          "state": "valid",
          "expires": "2027-05-05T23:59:59.000Z",
          "checkedAt": "2026-10-09T15:39:43.017059181Z"
        },
        "llmsTxt": {
          "url": "https://docs.mistral.ai/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:02:26.746229848Z"
        },
        "updatedAt": "2026-10-10T03:07:13.477589933Z"
      }
    },
    "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 5 of 7 scored categories. Mistral Voxtral Transcribe leads on payments \u0026 pricing and transparency \u0026 trust.",
    "b": {
      "slug": "openai-speech-to-text",
      "name": "OpenAI Speech to Text",
      "vendor": "OpenAI",
      "vendorUrl": "https://openai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "OpenAI's speech-to-text API. It transcribes uploaded audio files through `/v1/audio/transcriptions`, translates recordings into English through `/v1/audio/translations`, and transcribes live audio in Realtime transcription sessions over WebSocket or WebRTC.",
      "url": "https://www.anchorterminal.com/tools/openai-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key created by a person in the platform console (https://platform.openai.com/settings/organization/api-keys). Projects can carry a model allowlist or denylist and an IP allowlist, and Admin API keys are a separate credential that cannot call the audio endpoints (https://developers.openai.com/api/docs/guides/admin-apis).",
      "pricing": "usage",
      "pricingNotes": "$0.0045 an audio minute for `gpt-transcribe` and $0.017 for `gpt-live-transcribe`, billed from prepaid credits (https://developers.openai.com/api/docs/pricing). The rate limits guide names a Free tier with a $100 monthly usage limit, and the `gpt-transcribe` model page lists limits only from the Build tier, which needs $5 of credit purchases. Whether a new account can transcribe without paying was not established.",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the transcription guides, the endpoint reference, the pricing page or the OpenAPI document (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31785,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.openai.com/api/docs/guides/speech-to-text",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi/blob/main/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "streaming",
        "diarisation",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "go",
        "java"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 72.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 106,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
        "bestFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "strengths": [
          "`gpt-transcribe` is priced at $0.0045 an audio minute on a public page, with per-tier request limits of 5,000, 10,000 and 30,000 a minute",
          "The data controls page lists `/v1/audio/transcriptions` and `/v1/audio/translations` with no training, no abuse-monitoring retention and no stored application state",
          "A public OpenAPI 3.1 document, llms.txt and a Markdown twin of every docs page cover the audio endpoints",
          "The status page has an Audio component, shown at 100% uptime for July to October 2026",
          "Guide examples cover JavaScript, Python, Go, Java, C#, Ruby, a CLI and curl, and `file` and `model` are the only required fields"
        ],
        "weaknesses": [
          "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027",
          "The two named replacements return no speaker labels, word timestamps, `srt` or `vtt` output or English translation in the reviewed documentation",
          "Uploads stop at 25 MB, the caller splits longer recordings, and the `gpt-transcribe` model page marks the Batch API as not supported",
          "The Markdown twin of the endpoint reference lists the response fields and omits the request parameters, and its first example names a deprecated model",
          "openai.com answered our reader with a bot check, so the service terms, privacy policy, sub-processor list and any SLA were not read"
        ],
        "agentNotes": [
          "Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.",
          "Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.",
          "For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.",
          "Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.",
          "On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 72.4
          }
        ],
        "editorialScores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 62
        },
        "provenanceScore": 59
      },
      "connect": {
        "install": "pip install openai",
        "http": "curl --request POST \\\n  --url https://api.openai.com/v1/audio/transcriptions \\\n  --header \"Authorization: Bearer $OPENAI_API_KEY\" \\\n  --header 'Content-Type: multipart/form-data' \\\n  --form file=@/path/to/file/audio.mp3 \\\n  --form model=gpt-transcribe"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/openai-speech-to-text"
      },
      "sameCompany": [
        "openai-api",
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-realtime",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gpt-transcribe",
          "unit": "audio-minute",
          "usd": 0.0045
        },
        {
          "item": "gpt-live-transcribe (live audio)",
          "unit": "audio-minute",
          "usd": 0.017
        },
        {
          "item": "whisper-1 (deprecated)",
          "unit": "audio-minute",
          "usd": 0.006
        },
        {
          "item": "gpt-4o-transcribe-diarize (deprecated, estimated from token prices)",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "",
        "domain": "openai.com",
        "domainRegistered": "",
        "domainNote": "openai.com answered our researcher with a bot check on 9 October 2026, so the terms and privacy policy were not read on that day. The links are the two documents OpenAI's other listings here carry. openai.com answers our policy reader with HTTP 403 as well, so neither document has been read and both are recorded as unreadable. security.txt is PGP-signed with Bugcrowd and email contacts and has no Expires field.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 59
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.json",
      "live": {
        "slug": "openai-speech-to-text",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T03:07:15.369334284Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 133,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 137,
          "p95ms24h": 172,
          "samples24h": 117,
          "samples30d": 117,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 32,
              "ok": 32
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "minor",
          "summary": "Partial System Degradation",
          "checkedAt": "2026-10-10T03:02:29.613651511Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.877366615Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.73110612Z"
          }
        ],
        "githubStars": 31787,
        "pypiWeekly": 74761714,
        "updatedAt": "2026-10-10T03:07:15.369334284Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Mistral AI",
        "b": "OpenAI",
        "name": "Vendor"
      },
      {
        "a": "https://api.mistral.ai/v1",
        "b": "https://api.openai.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, websocket",
        "b": "HTTP, websocket",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
        "b": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-02-04",
        "b": "2026-08-26",
        "name": "Last release"
      },
      {
        "a": "2026-09-25",
        "b": "couldn't be read",
        "name": "Terms last updated"
      },
      {
        "a": "2026-09-03",
        "b": "couldn't be read",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes, with an opt-out",
        "b": "couldn't be read",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "773 stars",
        "b": "32k stars",
        "name": "Popularity"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 5 of 7 scored categories. Mistral Voxtral Transcribe leads on payments \u0026 pricing and transparency \u0026 trust.",
        "question": "Which is better for AI agents, Mistral Voxtral Transcribe or OpenAI Speech to Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Mistral Voxtral Transcribe and OpenAI Speech to Text need an API key?"
      },
      {
        "answer": "Yes. Mistral Voxtral Transcribe has a hosted endpoint at https://api.mistral.ai/v1 and OpenAI Speech to Text at https://api.openai.com/v1.",
        "question": "Can an agent call Mistral Voxtral Transcribe and OpenAI Speech to Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Payments \u0026 pricing, 40 against 20",
          "Transparency \u0026 trust, 82 against 61"
        ],
        "also": null,
        "goodFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "slug": "mistral-voxtral-transcribe",
        "watchFor": "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel"
      },
      {
        "aheadOn": [
          "Reliability, 80 against 53",
          "Security \u0026 auth, 86 against 70",
          "Maintenance \u0026 community, 75 against 66"
        ],
        "also": [
          "Agent-ready, a grade of BB or better",
          "No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them"
        ],
        "goodFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "slug": "openai-speech-to-text",
        "watchFor": "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.json",
        "title": "Amazon Transcribe vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.json",
        "title": "Amazon Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Cartesia Ink vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Groq Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Groq Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.json",
        "title": "Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.json",
        "title": "Mistral Voxtral Transcribe vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.json",
        "title": "Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.json",
        "title": "OpenAI Speech to Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json",
        "title": "OpenAI Speech to Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json",
        "title": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 27,
        "edge": "openai-speech-to-text",
        "key": "reliability",
        "mistral-voxtral-transcribe": 53,
        "name": "Reliability",
        "openai-speech-to-text": 80,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 3,
        "edge": "openai-speech-to-text",
        "key": "schema",
        "mistral-voxtral-transcribe": 85,
        "name": "Schema \u0026 documentation",
        "openai-speech-to-text": 88,
        "weight": 13
      },
      {
        "by": 1,
        "edge": "openai-speech-to-text",
        "key": "ergonomics",
        "mistral-voxtral-transcribe": 77,
        "name": "Agent ergonomics",
        "openai-speech-to-text": 78,
        "weight": 13
      },
      {
        "by": 16,
        "edge": "openai-speech-to-text",
        "key": "security",
        "mistral-voxtral-transcribe": 70,
        "name": "Security \u0026 auth",
        "openai-speech-to-text": 86,
        "weight": 14
      },
      {
        "by": 20,
        "edge": "mistral-voxtral-transcribe",
        "key": "payments",
        "mistral-voxtral-transcribe": 40,
        "name": "Payments \u0026 pricing",
        "openai-speech-to-text": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 9,
        "edge": "openai-speech-to-text",
        "key": "maintenance",
        "mistral-voxtral-transcribe": 66,
        "name": "Maintenance \u0026 community",
        "openai-speech-to-text": 75,
        "weight": 7
      },
      {
        "by": 21,
        "edge": "mistral-voxtral-transcribe",
        "key": "transparency",
        "mistral-voxtral-transcribe": 82,
        "name": "Transparency \u0026 trust",
        "openai-speech-to-text": 61,
        "weight": 7
      }
    ],
    "summary": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 5 of 7 scored categories. Mistral Voxtral Transcribe leads on payments \u0026 pricing and transparency \u0026 trust. Both do speech-to-text.",
    "verdicts": {
      "mistral-voxtral-transcribe": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.",
      "openai-speech-to-text": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.min.md"
  },
  "markdown": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 5 of 7 scored categories. Mistral Voxtral Transcribe leads on payments \u0026 pricing and transparency \u0026 trust. Both do speech-to-text.\n\n- Mistral Voxtral Transcribe: grade B, 64.1/100, rank #365 of 950. Markdown https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md · JSON https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json\n- OpenAI Speech to Text: grade BB, 72.4/100, rank #106 of 950. Markdown https://www.anchorterminal.com/tools/openai-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### Mistral Voxtral Transcribe (B)\n\nGood for: Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.\n\nAhead on:\n- Payments \u0026 pricing, 40 against 20\n- Transparency \u0026 trust, 82 against 61\n\nWatch for: Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel\n\n### OpenAI Speech to Text (BB)\n\nGood for: Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.\n\nAhead on:\n- Reliability, 80 against 53\n- Security \u0026 auth, 86 against 70\n- Maintenance \u0026 community, 75 against 66\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n- No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them\n\nWatch for: `whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027\n\n\n## Score by category\n\n| Category | Weight | Mistral Voxtral Transcribe | OpenAI Speech to Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 53 | 80 | OpenAI Speech to Text +27 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 85 | 88 | OpenAI Speech to Text +3 |\n| Agent ergonomics | 13% (16.2 this run) | 77 | 78 | OpenAI Speech to Text +1 |\n| Security \u0026 auth | 14% (17.5 this run) | 70 | 86 | OpenAI Speech to Text +16 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | Mistral Voxtral Transcribe +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 66 | 75 | OpenAI Speech to Text +9 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 82 | 61 | Mistral Voxtral Transcribe +21 |\n| Negative events | ≤15 | -3 | 0 | |\n| **Total** | | **64.1 · B** | **72.4 · BB** | |\n\n## Facts side by side\n\n| Fact | Mistral Voxtral Transcribe | OpenAI Speech to Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Mistral AI | OpenAI |\n| Hosted endpoint | `https://api.mistral.ai/v1` | `https://api.openai.com/v1` |\n| Transports | HTTP, websocket | HTTP, websocket |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0 | Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-02-04 | 2026-08-26 |\n| Terms last updated | 2026-09-25 | couldn't be read |\n| Privacy policy last updated | 2026-09-03 | couldn't be read |\n| Customer content may train models | yes, with an opt-out | couldn't be read |\n| Terms restrict automated access | not found in the text | couldn't be read |\n| Terms restrict benchmarking | yes | couldn't be read |\n| Terms or service can change without notice | yes | couldn't be read |\n| Arbitration or class-action waiver | not found in the text | couldn't be read |\n| Popularity | 773 stars | 32k stars |\n\n## Verdicts\n\n**Mistral Voxtral Transcribe.** Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.\n\n**OpenAI Speech to Text.** `gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.\n\n## Before you call either\n\n### Mistral Voxtral Transcribe\n\n1. Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`\n2. Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible\n3. Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`\n4. Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move\n5. Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key\n\n### OpenAI Speech to Text\n\n1. Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.\n2. Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.\n3. For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.\n4. Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.\n5. On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors.\n\n## Questions\n\n### Which is better for AI agents, Mistral Voxtral Transcribe or OpenAI Speech to Text?\n\nOpenAI Speech to Text scores 72.4 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 5 of 7 scored categories. Mistral Voxtral Transcribe leads on payments \u0026 pricing and transparency \u0026 trust.\n\n### Do Mistral Voxtral Transcribe and OpenAI Speech to Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call Mistral Voxtral Transcribe and OpenAI Speech to Text without installing anything?\n\nYes. Mistral Voxtral Transcribe has a hosted endpoint at https://api.mistral.ai/v1 and OpenAI Speech to Text at https://api.openai.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json, and with the fewest tokens: https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"mistral-voxtral-transcribe\", \"b\": \"openai-speech-to-text\"}`. From a terminal: `anchor compare mistral-voxtral-transcribe openai-speech-to-text`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json and https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n\n## Other comparisons with Mistral Voxtral Transcribe or OpenAI Speech to Text\n\n- [Amazon Transcribe vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.md)\n- [Amazon Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Azure AI Speech speech-to-text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-mistral-voxtral-transcribe.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.md)\n- [ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.md)\n- [Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Google Cloud Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md)\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Groq Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.md)\n- [Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.md)\n- [Mistral Voxtral Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.md)\n- [Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.md)\n- [OpenAI Speech to Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.md)\n- [OpenAI Speech to Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.md)\n- [OpenAI Speech to Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
        "url": ""
      }
    ],
    "description": "OpenAI Speech to Text scores 72.4 (BB) to Mistral Voxtral Transcribe's 64.1 (B) for speech-to-text. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "Mistral Voxtral Transcribe B 64.1",
      "OpenAI Speech to Text BB 72.4",
      "scores"
    ],
    "h1": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-mistral-voxtral-transcribe-vs-openai-speech-to-text.png",
    "path": "/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Mistral Voxtral Transcribe vs OpenAI Speech to Text (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text"
  },
  "tokens": {
    "markdown": 2950,
    "slim": 780
  },
  "version": 1
}
