{
  "data": {
    "similar": [
      {
        "grade": "BB",
        "json": "https://www.anchorterminal.com/tools/amazon-transcribe.json",
        "name": "Amazon Transcribe",
        "score": 73.4,
        "shared": [
          "speech.stt",
          "speech.streaming",
          "speech.batch",
          "speech.diarisation",
          "speech.languages"
        ],
        "slug": "amazon-transcribe"
      },
      {
        "grade": "BB",
        "json": "https://www.anchorterminal.com/tools/azure-speech-to-text.json",
        "name": "Azure AI Speech speech-to-text",
        "score": 73,
        "shared": [
          "speech.stt",
          "speech.streaming",
          "speech.batch",
          "speech.diarisation",
          "speech.languages"
        ],
        "slug": "azure-speech-to-text"
      },
      {
        "grade": "BB",
        "json": "https://www.anchorterminal.com/tools/deepgram-stt.json",
        "name": "Deepgram Speech-to-Text (Nova-3, Flux)",
        "score": 70.3,
        "shared": [
          "speech.stt",
          "speech.streaming",
          "speech.batch",
          "speech.diarisation",
          "speech.languages"
        ],
        "slug": "deepgram-stt"
      },
      {
        "grade": "BB",
        "json": "https://www.anchorterminal.com/tools/google-speech-to-text.json",
        "name": "Google Cloud Speech-to-Text",
        "score": 70.2,
        "shared": [
          "speech.stt",
          "speech.streaming",
          "speech.batch",
          "speech.diarisation",
          "speech.languages"
        ],
        "slug": "google-speech-to-text"
      },
      {
        "grade": "B",
        "json": "https://www.anchorterminal.com/tools/gladia-stt.json",
        "name": "Gladia Speech-to-Text API + MCP",
        "score": 69.5,
        "shared": [
          "speech.stt",
          "speech.streaming",
          "speech.batch",
          "speech.diarisation",
          "speech.languages"
        ],
        "slug": "gladia-stt"
      },
      {
        "grade": "B",
        "json": "https://www.anchorterminal.com/tools/elevenlabs-scribe.json",
        "name": "ElevenLabs Scribe Speech to Text API",
        "score": 68.9,
        "shared": [
          "speech.stt",
          "speech.streaming",
          "speech.batch",
          "speech.diarisation",
          "speech.languages"
        ],
        "slug": "elevenlabs-scribe"
      }
    ],
    "tool": {
      "slug": "mistral-voxtral-transcribe",
      "name": "Mistral Voxtral Transcribe",
      "vendor": "Mistral AI",
      "vendorUrl": "https://mistral.ai",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Mistral AI's speech-to-text API. Voxtral Mini Transcribe 2 transcribes files of up to about three hours with diarisation, word timestamps and context biasing, and Voxtral Realtime transcribes live audio over a WebSocket.",
      "url": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe",
      "markdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json",
      "repo": "https://github.com/mistralai/client-python",
      "license": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.mistral.ai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "mistralai"
        },
        {
          "registry": "npm",
          "name": "@mistralai/mistralai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from Studio, bound to one workspace, with an optional expiry date. Browser realtime clients use `rt_` tokens minted at `POST /v1/client/sessions`, valid about 900 seconds for one model (https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth).",
      "pricing": "freemium",
      "pricingNotes": "$0.003 per audio minute for Voxtral Mini Transcribe 2 and $0.006 for Voxtral Mini Transcribe Realtime (model cards on docs.mistral.ai, https://mistral.ai/pricing/api/). Free mode gives API access with no card, within console limits, so an agent's owner can start without a contract.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 773,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.mistral.ai/studio/audio/speech_to_text",
      "llmsTxt": "https://docs.mistral.ai/llms.txt",
      "openapi": "https://docs.mistral.ai/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.streaming",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "eu",
        "free-tier",
        "openapi",
        "llms-txt",
        "python",
        "typescript",
        "streaming",
        "batch",
        "open-weights"
      ],
      "lastRelease": "2026-02-04",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 64.1,
        "grade": "B",
        "agentReady": false,
        "rank": 329,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 10,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 82
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "breakdown": [
          {
            "key": "reliability",
            "name": "Reliability",
            "weight": 16,
            "effectiveWeight": 20,
            "score": 53,
            "points": 10.6,
            "reason": "Hosted reading. The docs point to status.mistral.ai, which our sibling Mistral listings read on 1 and 2 October 2026 as a Rootly page with component uptime bars (20). On 8 October the page answered curl and WebFetch with a Cloudflare check and HTTP 403, so the 90-day record for audio was not read (5, unread and not a finding against the vendor). The limits page names audio seconds a minute and audio seconds a month as the audio limits, with the numbers shown only in the Admin Panel (3 of 15). The error glossary tells callers to back off exponentially on 429, 500, 502, 503 and 504 and to read the `Retry-After` header, and the official SDKs retry with backoff. Transcription is a stateless call (15). No SLA found (0). The model cards mark Voxtral Mini Transcribe 2 and Voxtral Mini Transcribe Realtime as GA (10)."
          },
          {
            "key": "performance",
            "name": "Performance",
            "weight": 10,
            "effectiveWeight": 0,
            "pending": true,
            "points": 0,
            "reason": "Pending. Latency is measured per call by our probes, which haven't run yet, so this run doesn't score it. Its weight is shared across the assessed categories until the first probe window closes."
          },
          {
            "key": "schema",
            "name": "Schema \u0026 documentation",
            "weight": 13,
            "effectiveWeight": 16.25,
            "score": 85,
            "points": 13.81,
            "reason": "Model reading. The public OpenAPI document covers `/v1/audio/transcriptions`, its SSE streaming form and `/v1/client/sessions`. The realtime WebSocket messages are described by SDK types and guide examples, not by the document (22 of 25). llms.txt with Markdown twins of the speech-to-text, offline, realtime and client authentication pages (10). The guides say which model each endpoint runs, that `timestamp_granularities` and `language` cannot be combined and that realtime cannot diarise, and carry an OpenAI SDK compatibility table. In the OpenAPI document `diarize`, `context_bias` and `temperature` have no description (14 of 20). `timestamp_granularities` is an enum of `segment` and `word`, `language` is a two-letter pattern and `context_bias` items have a pattern, while `model` is a free string (12 of 15). curl, Python and TypeScript examples on each guide and a shared error glossary. The OpenAPI response example still names `voxtral-mini-2507` (12 of 15). Dated model identifiers and a dated changelog (15)."
          },
          {
            "key": "ergonomics",
            "name": "Agent ergonomics",
            "weight": 13,
            "effectiveWeight": 16.25,
            "score": 77,
            "points": 12.51,
            "reason": "API reading, as for the other speech-to-text listings. Segments, word timestamps and speaker labels are returned only when asked for, and the default response is the text, the language and usage. No subtitle or plain-text output format was found (18 of 25). There is no transcript store to page through. Volume goes through the batch API, which accepts `/v1/audio/transcriptions` (12 of 20). Errors return `type`, `param` and `code` with a fix per status in the glossary (17 of 20). The call is stateless and safe to retry, the SDKs retry with backoff, and no idempotency key exists or is needed (15 of 20). `model` is the only required field beside the audio source, and there are official Python and TypeScript SDKs plus OpenAI SDK compatibility for the batch endpoint (15)."
          },
          {
            "key": "security",
            "name": "Security \u0026 auth",
            "weight": 14,
            "effectiveWeight": 17.5,
            "score": 70,
            "points": 12.25,
            "reason": "Model reading, scored like the other Mistral listings. Keys are bound to one workspace, can carry an expiry date and a connector scope, and organisations can enforce a maximum key lifetime. Browser clients get `rt_` tokens from `POST /v1/client/sessions` that last about 900 seconds, work for one model and travel in `Sec-WebSocket-Protocol`. A key cannot be limited to transcription (26 of 30). The commercial terms effective 25 September 2026 say Mistral will not train on customer data unless the customer opted in, or did not opt out where a product defaults to opt-in, without naming the API's default. The privacy and data controls page the docs link to returned 404 (10 of 20). Input and output are kept 30 rolling days for abuse monitoring, and `/v1/audio/transcriptions` is on the zero data retention list for paid plans, on request and subject to approval (12 of 15). Audit logs cover API key actions on Enterprise plans only (10 of 15). security.txt is valid until 5 May 2027 and names a HackerOne submission form, and the docs publish three security advisories for 2026. Certifications sit in a trust centre drawn by script that we could not read (12 of 20)."
          },
          {
            "key": "payments",
            "name": "Payments \u0026 pricing",
            "weight": 10,
            "effectiveWeight": 12.5,
            "score": 40,
            "points": 5,
            "reason": "No x402, MPP or L402 (0). $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime on public pages (20). Free mode gives API access with no card, within limits shown in the console (20). A person signs up in a browser (0)."
          },
          {
            "key": "tasks",
            "name": "Task success",
            "weight": 10,
            "effectiveWeight": 0,
            "pending": true,
            "points": 0,
            "reason": "Pending. Task success needs the category task suites run through each tool, which haven't run yet, so this run doesn't score it. Its weight is shared across the assessed categories until then. A data provider's data-quality score is published on its listing now and becomes half of this category when it's scored."
          },
          {
            "key": "maintenance",
            "name": "Maintenance \u0026 community",
            "weight": 7,
            "effectiveWeight": 8.75,
            "score": 66,
            "points": 5.78,
            "reason": "Model reading. The last dated change to transcription is the 4 February 2026 release of both models with diarisation and context biasing, eight months ago. The SDKs that carry the realtime client were released on 6 October 2026 (Python 3.1.0), so we give 10 of 30, a departure from the strict reading of 0. The lifecycle policy promises six months' notice for GA models (12 of 12). No transcription model was retired in the last 90 days of the changelog (8 of 8). Five dated platform changelog entries between 16 July and 6 October, none about audio (8 of 15). The eight newest open issues on client-python have zero or one comment (5 of 10). Official `mistralai` 3.1.0 on PyPI and `@mistralai/mistralai` 2.7.0 on npm (15). Both are generated from the OpenAPI document, with test, lint and example workflows in the repository (8 of 10)."
          },
          {
            "key": "transparency",
            "name": "Transparency \u0026 trust",
            "weight": 7,
            "effectiveWeight": 8.75,
            "score": 82,
            "points": 7.18,
            "note": "editorial 73, provenance 91",
            "reason": "The hosted service is closed under commercial terms. The Realtime weights are published on Hugging Face under Apache-2.0 and the SDKs are Apache-2.0, while Voxtral Mini Transcribe 2 has no public weights (20 of 30). The privacy policy effective 3 September 2026 (30 rolling days), the zero data retention page and the terms agree on retention, and the terms leave the API's training default unnamed (20 of 30). A model lifecycle page with notice periods per stage (20). Opt-in EU and US regional endpoints are documented, the global endpoint commits to no location, and the provider list is in a trust centre we could not read (13 of 20)."
          }
        ],
        "assessment": {
          "date": "2026-10-08",
          "basis": "public evidence",
          "confidence": "medium",
          "notes": {
            "ergonomics": "API reading, as for the other speech-to-text listings. Segments, word timestamps and speaker labels are returned only when asked for, and the default response is the text, the language and usage. No subtitle or plain-text output format was found (18 of 25). There is no transcript store to page through. Volume goes through the batch API, which accepts `/v1/audio/transcriptions` (12 of 20). Errors return `type`, `param` and `code` with a fix per status in the glossary (17 of 20). The call is stateless and safe to retry, the SDKs retry with backoff, and no idempotency key exists or is needed (15 of 20). `model` is the only required field beside the audio source, and there are official Python and TypeScript SDKs plus OpenAI SDK compatibility for the batch endpoint (15).",
            "maintenance": "Model reading. The last dated change to transcription is the 4 February 2026 release of both models with diarisation and context biasing, eight months ago. The SDKs that carry the realtime client were released on 6 October 2026 (Python 3.1.0), so we give 10 of 30, a departure from the strict reading of 0. The lifecycle policy promises six months' notice for GA models (12 of 12). No transcription model was retired in the last 90 days of the changelog (8 of 8). Five dated platform changelog entries between 16 July and 6 October, none about audio (8 of 15). The eight newest open issues on client-python have zero or one comment (5 of 10). Official `mistralai` 3.1.0 on PyPI and `@mistralai/mistralai` 2.7.0 on npm (15). Both are generated from the OpenAPI document, with test, lint and example workflows in the repository (8 of 10).",
            "payments": "No x402, MPP or L402 (0). $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime on public pages (20). Free mode gives API access with no card, within limits shown in the console (20). A person signs up in a browser (0).",
            "reliability": "Hosted reading. The docs point to status.mistral.ai, which our sibling Mistral listings read on 1 and 2 October 2026 as a Rootly page with component uptime bars (20). On 8 October the page answered curl and WebFetch with a Cloudflare check and HTTP 403, so the 90-day record for audio was not read (5, unread and not a finding against the vendor). The limits page names audio seconds a minute and audio seconds a month as the audio limits, with the numbers shown only in the Admin Panel (3 of 15). The error glossary tells callers to back off exponentially on 429, 500, 502, 503 and 504 and to read the `Retry-After` header, and the official SDKs retry with backoff. Transcription is a stateless call (15). No SLA found (0). The model cards mark Voxtral Mini Transcribe 2 and Voxtral Mini Transcribe Realtime as GA (10).",
            "schema": "Model reading. The public OpenAPI document covers `/v1/audio/transcriptions`, its SSE streaming form and `/v1/client/sessions`. The realtime WebSocket messages are described by SDK types and guide examples, not by the document (22 of 25). llms.txt with Markdown twins of the speech-to-text, offline, realtime and client authentication pages (10). The guides say which model each endpoint runs, that `timestamp_granularities` and `language` cannot be combined and that realtime cannot diarise, and carry an OpenAI SDK compatibility table. In the OpenAPI document `diarize`, `context_bias` and `temperature` have no description (14 of 20). `timestamp_granularities` is an enum of `segment` and `word`, `language` is a two-letter pattern and `context_bias` items have a pattern, while `model` is a free string (12 of 15). curl, Python and TypeScript examples on each guide and a shared error glossary. The OpenAPI response example still names `voxtral-mini-2507` (12 of 15). Dated model identifiers and a dated changelog (15).",
            "security": "Model reading, scored like the other Mistral listings. Keys are bound to one workspace, can carry an expiry date and a connector scope, and organisations can enforce a maximum key lifetime. Browser clients get `rt_` tokens from `POST /v1/client/sessions` that last about 900 seconds, work for one model and travel in `Sec-WebSocket-Protocol`. A key cannot be limited to transcription (26 of 30). The commercial terms effective 25 September 2026 say Mistral will not train on customer data unless the customer opted in, or did not opt out where a product defaults to opt-in, without naming the API's default. The privacy and data controls page the docs link to returned 404 (10 of 20). Input and output are kept 30 rolling days for abuse monitoring, and `/v1/audio/transcriptions` is on the zero data retention list for paid plans, on request and subject to approval (12 of 15). Audit logs cover API key actions on Enterprise plans only (10 of 15). security.txt is valid until 5 May 2027 and names a HackerOne submission form, and the docs publish three security advisories for 2026. Certifications sit in a trust centre drawn by script that we could not read (12 of 20).",
            "transparency": "The hosted service is closed under commercial terms. The Realtime weights are published on Hugging Face under Apache-2.0 and the SDKs are Apache-2.0, while Voxtral Mini Transcribe 2 has no public weights (20 of 30). The privacy policy effective 3 September 2026 (30 rolling days), the zero data retention page and the terms agree on retention, and the terms leave the API's training default unnamed (20 of 30). A model lifecycle page with notice periods per stage (20). Opt-in EU and US regional endpoints are documented, the global endpoint commits to no location, and the provider list is in a trust centre we could not read (13 of 20)."
          },
          "sources": [
            {
              "what": "speech-to-text overview",
              "url": "https://docs.mistral.ai/studio/audio/speech_to_text",
              "seen": "2026-10-08"
            },
            {
              "what": "offline transcription guide",
              "url": "https://docs.mistral.ai/studio/audio/speech_to_text/offline_transcription",
              "seen": "2026-10-08"
            },
            {
              "what": "realtime transcription guide",
              "url": "https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription",
              "seen": "2026-10-08"
            },
            {
              "what": "realtime client authentication",
              "url": "https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth",
              "seen": "2026-10-08"
            },
            {
              "what": "OpenAPI document",
              "url": "https://docs.mistral.ai/openapi.yaml",
              "seen": "2026-10-08"
            },
            {
              "what": "Voxtral Mini Transcribe 2 model card and price",
              "url": "https://docs.mistral.ai/models/voxtral-mini-transcribe-26-02",
              "seen": "2026-10-08"
            },
            {
              "what": "Voxtral Mini Transcribe Realtime model card and price",
              "url": "https://docs.mistral.ai/models/voxtral-mini-transcribe-realtime-26-02",
              "seen": "2026-10-08"
            },
            {
              "what": "API pricing",
              "url": "https://mistral.ai/pricing/api/",
              "seen": "2026-10-08"
            },
            {
              "what": "changelog",
              "url": "https://docs.mistral.ai/resources/changelogs",
              "seen": "2026-10-08"
            },
            {
              "what": "model lifecycle policy",
              "url": "https://docs.mistral.ai/inference/model-lifecycle",
              "seen": "2026-10-08"
            },
            {
              "what": "error glossary",
              "url": "https://docs.mistral.ai/resources/error-glossary",
              "seen": "2026-10-08"
            },
            {
              "what": "usage and limits",
              "url": "https://docs.mistral.ai/admin/billing-usage/usage-limits",
              "seen": "2026-10-08"
            },
            {
              "what": "API keys",
              "url": "https://docs.mistral.ai/admin/identity-access/api-keys",
              "seen": "2026-10-08"
            },
            {
              "what": "zero data retention",
              "url": "https://docs.mistral.ai/admin/monitor-comply/zero-data-retention",
              "seen": "2026-10-08"
            },
            {
              "what": "audit logs",
              "url": "https://docs.mistral.ai/admin/monitor-comply/audit-logs/overview",
              "seen": "2026-10-08"
            },
            {
              "what": "security advisory MAI-2026-002",
              "url": "https://docs.mistral.ai/resources/security-advisories/MAI-2026-002",
              "seen": "2026-10-08"
            },
            {
              "what": "Studio MCP server",
              "url": "https://docs.mistral.ai/resources/mcp",
              "seen": "2026-10-08"
            },
            {
              "what": "commercial terms of service",
              "url": "https://legal.mistral.ai/terms/commercial-terms-of-service",
              "seen": "2026-10-08"
            },
            {
              "what": "privacy policy",
              "url": "https://legal.mistral.ai/terms/privacy-policy",
              "seen": "2026-10-08"
            },
            {
              "what": "security.txt",
              "url": "https://mistral.ai/.well-known/security.txt",
              "seen": "2026-10-08"
            },
            {
              "what": "Python SDK releases",
              "url": "https://pypi.org/pypi/mistralai/json",
              "seen": "2026-10-08"
            },
            {
              "what": "Python SDK repository",
              "url": "https://github.com/mistralai/client-python",
              "seen": "2026-10-08"
            },
            {
              "what": "Realtime open weights",
              "url": "https://huggingface.co/mistralai/Voxtral-Mini-4B-Realtime-2602",
              "seen": "2026-10-08"
            }
          ],
          "openQuestions": [
            "unchecked: status.mistral.ai answered every request with a Cloudflare check (HTTP 403) on 8 October 2026, so the audio component's uptime and the 90-day incident record were not read",
            "unchecked: trust.mistral.ai is drawn by script, so certifications and the provider list were not read",
            "unchecked: the privacy and data controls page linked from the zero data retention page returned 404, so the API's training default was not established",
            "unchecked: whether Free mode includes transcription minutes and whether sign-up needs a phone number. The quickstart says no card is needed",
            "The pricing table shows $0.0003 a minute in a second column for Voxtral Mini Transcribe 2 without a label we could tie to batch or cached use, so it is not listed as a price",
            "Whether zero data retention covers the realtime WebSocket at `/v1/audio/transcriptions/realtime`. The list names `/v1/audio/transcriptions` only",
            "The transcription curl examples send the key in an `x-api-key` header, while the error glossary gives `Authorization: Bearer`",
            "The other Mistral listings took no deduction for advisory MAI-2026-002. This one takes 3 because the realtime client is only reachable through the SDK"
          ]
        },
        "negative": -3,
        "negativeNotes": [
          "2026-05-12. Compromised `mistralai` 2.4.6 on PyPI ran a credential-harvesting script on import for about three hours, and three `@mistralai/mistralai` versions on npm were also replaced. Mistral published advisory MAI-2026-002, removed the packages and closed its investigation on 14 May. Fixed and documented, so 3 points. https://docs.mistral.ai/resources/security-advisories/MAI-2026-002"
        ],
        "verdict": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.",
        "bestFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "strengths": [
          "Published prices of $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime",
          "One request with `model` and a `file`, `file_url` or `file_id`, for audio of up to about three hours",
          "Browser clients use `rt_` tokens that last about 900 seconds and are limited to one model",
          "`/v1/audio/transcriptions` is on the zero data retention list for paid plans, and the Realtime weights are Apache-2.0 on Hugging Face",
          "OpenAPI document, llms.txt and Markdown guides, with a six-month retirement notice for GA models"
        ],
        "weaknesses": [
          "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel",
          "`timestamp_granularities` cannot be used with `language`, and the realtime model cannot diarise",
          "13 languages, and context biasing is tuned for English with other languages described as experimental",
          "The compromised `mistralai` 2.4.6 on PyPI harvested credentials on import for three hours on 12 May 2026",
          "No SLA found, and status.mistral.ai answered our requests with a bot check on 8 October 2026",
          "The Studio MCP server has no transcription tool"
        ],
        "agentNotes": [
          "Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`",
          "Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible",
          "Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`",
          "Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move",
          "Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 64.1
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 73
        },
        "provenanceScore": 91
      },
      "connect": {
        "install": "pip install mistralai   # realtime: pip install \"mistralai[realtime]\"   # or: npm i @mistralai/mistralai",
        "http": "curl --location 'https://api.mistral.ai/v1/audio/transcriptions' \\\n  --header \"x-api-key: $MISTRAL_API_KEY\" \\\n  --form 'file_url=\"https://docs.mistral.ai/audio/obama.mp3\"' \\\n  --form 'model=\"voxtral-mini-latest\"'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/mistral-voxtral-transcribe"
      },
      "sameCompany": [
        "mistral-api",
        "mistral-embeddings",
        "mistral-moderation",
        "mistral-ocr"
      ],
      "notable": [
        "Voxtral Mini Transcribe 2 (`voxtral-mini-2602`, alias `voxtral-mini-latest`) and Voxtral Mini Transcribe Realtime (`voxtral-mini-transcribe-realtime-2602`) were released on 2026-02-04 and are marked GA (https://docs.mistral.ai/resources/changelogs)",
        "Batch transcription accepts about three hours of audio in one request, as `file`, `file_url` or `file_id` (https://docs.mistral.ai/studio/audio/speech_to_text/offline_transcription)",
        "`timestamp_granularities` is not compatible with `language`, and realtime is not compatible with `diarize` (https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription)",
        "Realtime runs over `wss://api.mistral.ai/v1/audio/transcriptions/realtime` with a `target_streaming_delay_ms` setting. Mistral says latency can be configured below 200 ms",
        "The endpoint follows the OpenAI audio transcriptions request shape, with `diarize` and `context_bias` passed through `extra_body` and no `file_url` support in that mode",
        "The Studio MCP server at `https://api.mistral.ai/mcp` has no transcription tool (https://docs.mistral.ai/resources/mcp)"
      ],
      "area": "voice",
      "details": [
        {
          "label": "Models",
          "value": "`voxtral-mini-latest` (Voxtral Mini Transcribe 2, `voxtral-mini-2602`) for files, `voxtral-mini-transcribe-realtime-2602` for live audio"
        },
        {
          "label": "Languages",
          "value": "13: English, Chinese, Hindi, Spanish, Arabic, French, Portuguese, Russian, German, Japanese, Korean, Italian and Dutch, with language detection"
        },
        {
          "label": "Max audio",
          "value": "About 3 hours a request on Voxtral Mini Transcribe 2"
        },
        {
          "label": "Diarisation",
          "value": "`diarize` on the batch endpoint only"
        },
        {
          "label": "Extras",
          "value": "Segment or word timestamps, `context_bias` of up to 100 terms (tuned for English), SSE streaming of a file transcription"
        },
        {
          "label": "Streaming latency",
          "value": "Vendor says configurable below 200 ms through `target_streaming_delay_ms`"
        },
        {
          "label": "Realtime input",
          "value": "PCM audio such as `pcm_s16le` at 16 kHz over a WebSocket"
        },
        {
          "label": "Rate limits",
          "value": "Audio seconds a minute and a month per workspace, numbers shown in the Admin Panel"
        },
        {
          "label": "Free tier",
          "value": "Free mode, no card, console limits"
        },
        {
          "label": "Data retention",
          "value": "30 rolling days for abuse monitoring. Zero data retention on request for paid plans covers `/v1/audio/transcriptions`"
        },
        {
          "label": "Regions",
          "value": "Global endpoint, with opt-in EU and US endpoints at 1.1 times list price"
        },
        {
          "label": "Open weights",
          "value": "Voxtral Mini 4B Realtime 2602 on Hugging Face, Apache-2.0"
        },
        {
          "label": "MCP server",
          "value": "None for transcription"
        }
      ],
      "unitPrices": [
        {
          "item": "Voxtral Mini Transcribe 2 (batch)",
          "unit": "audio-minute",
          "usd": 0.003
        },
        {
          "item": "Voxtral Mini Transcribe Realtime",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "Mistral AI (RCS Paris 952 418 325)",
        "domain": "mistral.ai",
        "domainRegistered": "2019-05-15",
        "endpointOnVendorDomain": true,
        "terms": "https://legal.mistral.ai/terms/commercial-terms-of-service",
        "privacy": "https://legal.mistral.ai/terms/privacy-policy",
        "statusPage": "https://status.mistral.ai",
        "changelog": "https://docs.mistral.ai/resources/changelogs",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 91,
        "checks": [
          {
            "check": "Legal entity named",
            "value": "Mistral AI (RCS Paris 952 418 325)",
            "points": 20,
            "max": 20,
            "state": "ok"
          },
          {
            "check": "Domain age",
            "value": "mistral.ai, registered 2019-05-15 (7 years)",
            "points": 11,
            "max": 15,
            "state": "part"
          },
          {
            "check": "Endpoint on the vendor's domain",
            "value": "api.mistral.ai",
            "points": 15,
            "max": 15,
            "state": "ok"
          },
          {
            "check": "Terms of service",
            "value": "read, states 6 of the 7 things a reader expects, and has 2 clauses that cost points",
            "points": 5.1,
            "max": 10,
            "state": "part"
          },
          {
            "check": "Privacy policy",
            "value": "read, states 8 of the 8 things a reader expects",
            "points": 10,
            "max": 10,
            "state": "ok"
          },
          {
            "check": "Status page",
            "value": "status.mistral.ai",
            "points": 10,
            "max": 10,
            "state": "ok"
          },
          {
            "check": "Changelog",
            "value": "published",
            "points": 10,
            "max": 10,
            "state": "ok"
          },
          {
            "check": "security.txt",
            "value": "valid",
            "points": 10,
            "max": 10,
            "state": "ok"
          }
        ],
        "policies": [
          {
            "kind": "terms",
            "url": "https://legal.mistral.ai/terms/commercial-terms-of-service",
            "state": "read",
            "readAt": "2026-10-08",
            "statedDate": "2026-09-25",
            "words": 6763,
            "points": 5.1,
            "max": 10,
            "expected": [
              {
                "key": "terms.date",
                "label": "Gives the date it was last updated",
                "found": true,
                "quote": "Effective: September 25, 2026",
                "says": "Last updated 2026-09-25"
              },
              {
                "key": "terms.law",
                "label": "Names the governing law or courts",
                "found": true,
                "quote": "…with it or its subject matter or formation will be governed by and construed in accordance with the laws of the Republic of Singapore.",
                "says": "The law of the Republic of Singapore"
              },
              {
                "key": "terms.liability",
                "label": "States a limit on its liability",
                "found": true,
                "quote": "…TOTAL AGGREGATE LIABILITY ARISING OUT OF OR RELATING TO THESE TERMS OR THE MISTRAL AI PRODUCTS WILL NOT EXCEED THE TOTAL AMOUNTS PAID BY CUSTOMER TO MISTRAL AI IN THE TWELVE (12) MONTHS PRECEDING THE EVENT(S) GIVING RISE TO THE CLAIM.",
                "says": "Capped at the fees paid in the 12 months before the claim"
              },
              {
                "key": "terms.termination",
                "label": "Says how the agreement or account can be ended",
                "found": true,
                "quote": "Mistral AI may terminate or suspend access to the Beta Products at any time in Mistral AI’s sole discretion without liability."
              },
              {
                "key": "terms.changes",
                "label": "Says how changes to the terms are announced",
                "found": true,
                "quote": "Material updates to the Terms become effective thirty (30) days after notice is provided to Customer.",
                "says": "Gives thirty days of notice before a change"
              },
              {
                "key": "terms.use",
                "label": "Lists what users may not do",
                "found": true,
                "quote": "Customer will not, and will not permit any other person (including any End User) to:"
              },
              {
                "key": "terms.sla",
                "label": "Refers to a service level or uptime commitment",
                "found": false
              }
            ],
            "toKnow": [
              {
                "key": "training.optout",
                "label": "Says it may use customer content to train or improve models, and gives an opt-out",
                "found": true,
                "quote": "…royalty-free, fully-paid license (with the right to sublicense to our service providers) to use Customer Data and Outputs solely as provided in the preceding sentence to train Mistral AI’s artificial intelligence models."
              },
              {
                "key": "terms.benchmark",
                "label": "Restricts benchmarking or competitive use",
                "found": true,
                "quote": "To the extent permitted by applicable law, Customer may not use image Outputs to develop or train any image generation product that competes with a Mistral AI Product.",
                "costsPoints": true
              },
              {
                "key": "terms.nonotice",
                "label": "Says the terms or the service can change without notice",
                "found": true,
                "quote": "Any updates (i) that are not material or (ii) made for compliance with applicable law or to address a material security risk become effective immediately upon being posted at https://legal.mistral.ai/terms.",
                "costsPoints": true
              }
            ],
            "notes": [
              {
                "date": "2026-10-08",
                "text": "Customer data and outputs from Labs or Preview Models may be used for training, and opt-out choices made for other products, including zero data retention, do not apply.",
                "quote": "(i) Mistral AI may use Customer Data and Outputs generated from Labs or Preview Models to train its artificial intelligence models and (ii) the opt-out preferences you selected for other Mistral AI Products (including through zero data retention) does not apply to Labs or Preview Models."
              },
              {
                "date": "2026-10-08",
                "text": "When the customer gives feedback, Mistral AI may use it together with the associated customer data and output in any manner without restriction.",
                "quote": "Mistral AI may use, copy, disclose, license, distribute, and exploit such Feedback, along with the Customer Data and Output associated with such Feedback, in any manner without any obligation, royalty, or restriction based on intellectual property rights or otherwise."
              },
              {
                "date": "2026-10-08",
                "text": "The customer may not state or imply that output was generated by a human when the Mistral AI Products generated it.",
                "quote": "Customer may not represent or imply that the Output was generated by a human when it was generated by the Mistral AI Products."
              }
            ]
          },
          {
            "kind": "privacy",
            "url": "https://legal.mistral.ai/terms/privacy-policy",
            "state": "read",
            "readAt": "2026-10-08",
            "statedDate": "2026-09-03",
            "words": 3844,
            "points": 10,
            "max": 10,
            "expected": [
              {
                "key": "privacy.date",
                "label": "Gives the date it was last updated",
                "found": true,
                "quote": "Effective: September 3, 2026",
                "says": "Last updated 2026-09-03"
              },
              {
                "key": "privacy.collected",
                "label": "Says what personal data is collected",
                "found": true,
                "quote": "This Privacy Policy explains, in a clear and simple way, how we collect, use, and protect your personal data when you use our generative AI Mistral AI Products such as Vibe or Mistral AI Studio (the \"Mistral AI Products\")."
              },
              {
                "key": "privacy.retention",
                "label": "Says how long data is kept",
                "found": true,
                "quote": "Contracts, commercial licenses to use our Models and contact details relating to both: we keep your data for the duration of the contract and for 5 years following the termination of the contract.",
                "says": "Names a period of 5 years"
              },
              {
                "key": "privacy.processors",
                "label": "Says who else receives the data",
                "found": true,
                "quote": "Additionally, we may share all or part of your personal data with our service providers."
              },
              {
                "key": "privacy.sale",
                "label": "Says whether personal data is sold or shared for advertising",
                "found": true,
                "quote": "Mistral AI has not “sold,” “shared,” or engaged in “targeted advertising” with (as those terms are defined under U.S."
              },
              {
                "key": "privacy.rights",
                "label": "Says what rights people have over their data",
                "found": true,
                "quote": "You have the right to know if we process your personal data."
              },
              {
                "key": "privacy.contact",
                "label": "Gives a privacy contact",
                "found": true,
                "quote": "By sending us a letter at Mistral AI, Attn: DPO, Mistral AI, 15 rue des Halles, 75001 Paris, France.",
                "says": "Names a data protection officer"
              },
              {
                "key": "privacy.transfers",
                "label": "Says where data is transferred or stored",
                "found": true,
                "quote": "Additionally, we attach the most recent version of the European Commission’s Standard Contractual Clauses to all such contracts.",
                "says": "Relies on standard contractual clauses"
              }
            ],
            "toKnow": [
              {
                "key": "training.optout",
                "label": "Says it may use customer content to train or improve models, and gives an opt-out",
                "found": true,
                "quote": "We’ve introduced a user control which allows you to object to the use of your input and output data for model training directly from your account."
              }
            ],
            "notes": [
              {
                "date": "2026-10-08",
                "text": "The policy does not apply where a business uses the products to process personal data, in which case Mistral AI is the processor.",
                "quote": "This Privacy Policy does not apply if you use our Mistral AI Products to process personal data in the context of your business activities."
              },
              {
                "date": "2026-10-08",
                "text": "Mistral AI says that rights requests concerning model training have technical limitations and may involve a complex technical process.",
                "quote": "However, when your request concerns the training of our models, it’s important to note that your rights have technical limitations and fulfilling your requests might involve a complex technical process."
              }
            ]
          }
        ]
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.json",
      "live": {
        "slug": "mistral-voxtral-transcribe",
        "probe": {
          "target": "https://api.mistral.ai/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:46:34.670206805Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 38,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 47,
          "p95ms24h": 77,
          "samples24h": 44,
          "samples30d": 44,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 44,
              "ok": 44
            }
          ]
        },
        "updatedAt": "2026-10-09T11:46:34.670206805Z"
      }
    },
    "verify": {
      "accepts": "a page on mistral.ai or one of its subdomains, or the README of github.com/mistralai/client-python",
      "badgeUrl": "https://www.anchorterminal.com/badges/mistral-voxtral-transcribe.svg",
      "body": {
        "slug": "mistral-voxtral-transcribe",
        "url": "the page with the badge or the link"
      },
      "docs": "https://www.anchorterminal.com/builders/#verify",
      "effect": "none, it never changes a grade, rank or review",
      "endpoint": "https://www.anchorterminal.com/api/v1/verify",
      "listingUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe",
      "mcpTool": "verify_listing",
      "recheck": "weekly; two failed checks in a row and it lapses, a later pass restores it",
      "snippets": {
        "html": "\u003ca href=\"https://www.anchorterminal.com/tools/mistral-voxtral-transcribe\"\u003e\u003cimg src=\"https://www.anchorterminal.com/badges/mistral-voxtral-transcribe.svg\" alt=\"Mistral Voxtral Transcribe on Anchor Terminal\" height=\"20\"\u003e\u003c/a\u003e",
        "markdown": "[![Mistral Voxtral Transcribe on Anchor Terminal](https://www.anchorterminal.com/badges/mistral-voxtral-transcribe.svg)](https://www.anchorterminal.com/tools/mistral-voxtral-transcribe)",
        "link": "\u003ca href=\"https://www.anchorterminal.com/tools/mistral-voxtral-transcribe\"\u003eMistral Voxtral Transcribe on Anchor Terminal\u003c/a\u003e"
      }
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe",
    "json": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md",
    "slim": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.min.md"
  },
  "markdown": "## Overview\n\n**Grade B · 64.1/100 · rank #329 of 842 · #10 in Speech-to-text · not agent-ready · confidence medium**\n\n\nMore from Mistral AI, listed separately because each is its own product: [Mistral AI API](https://www.anchorterminal.com/tools/mistral-api.md) (Model APIs \u0026 inference), [Mistral Embed and Codestral Embed](https://www.anchorterminal.com/tools/mistral-embeddings.md) (Embeddings \u0026 rerankers), [Mistral Moderation API](https://www.anchorterminal.com/tools/mistral-moderation.md) (Guardrails \u0026 safety filters), [Mistral OCR API](https://www.anchorterminal.com/tools/mistral-ocr.md) (Document parsing \u0026 extraction).\n\n## Assessment\n\nBatch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.\n\n## Facts\n\n| Field | Value |\n| --- | --- |\n| Vendor | Mistral AI (https://mistral.ai) |\n| Kind | Model API |\n| Category | Speech-to-text (https://www.anchorterminal.com/categories/speech-to-text) |\n| Transport | HTTP, websocket |\n| Endpoint | `https://api.mistral.ai/v1` |\n| Auth | API key · Self-serve bearer key from Studio, bound to one workspace, with an optional expiry date. Browser realtime clients use `rt_` tokens minted at `POST /v1/client/sessions`, valid about 900 seconds for one model (https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth). |\n| Pricing | Freemium (Freemium) · $0.003 per audio minute for Voxtral Mini Transcribe 2 and $0.006 for Voxtral Mini Transcribe Realtime (model cards on docs.mistral.ai, https://mistral.ai/pricing/api/). Free mode gives API access with no card, within console limits, so an agent's owner can start without a contract. |\n| x402 | No · No x402, MPP or L402 in the docs, the OpenAPI document or the pricing page (checked 2026-10-08). |\n| Licence | Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0 |\n| Packages | pypi: `mistralai`; npm: `@mistralai/mistralai` |\n| Source | https://github.com/mistralai/client-python |\n| Docs | https://docs.mistral.ai/studio/audio/speech_to_text |\n| llms.txt | https://docs.mistral.ai/llms.txt |\n| Last release | 2026-02-04 |\n| GitHub stars | 773 (as of 2026-10-08) |\n| Models | `voxtral-mini-latest` (Voxtral Mini Transcribe 2, `voxtral-mini-2602`) for files, `voxtral-mini-transcribe-realtime-2602` for live audio |\n| Languages | 13: English, Chinese, Hindi, Spanish, Arabic, French, Portuguese, Russian, German, Japanese, Korean, Italian and Dutch, with language detection |\n| Max audio | About 3 hours a request on Voxtral Mini Transcribe 2 |\n| Diarisation | `diarize` on the batch endpoint only |\n| Extras | Segment or word timestamps, `context_bias` of up to 100 terms (tuned for English), SSE streaming of a file transcription |\n| Streaming latency | Vendor says configurable below 200 ms through `target_streaming_delay_ms` |\n| Realtime input | PCM audio such as `pcm_s16le` at 16 kHz over a WebSocket |\n| Rate limits | Audio seconds a minute and a month per workspace, numbers shown in the Admin Panel |\n| Free tier | Free mode, no card, console limits |\n| Data retention | 30 rolling days for abuse monitoring. Zero data retention on request for paid plans covers `/v1/audio/transcriptions` |\n| Regions | Global endpoint, with opt-in EU and US endpoints at 1.1 times list price |\n| Open weights | Voxtral Mini 4B Realtime 2602 on Hugging Face, Apache-2.0 |\n| MCP server | None for transcription |\n| Capabilities | speech.stt, speech.batch, speech.streaming, speech.diarisation, speech.languages |\n| Tags | official, hosted, model, eu, free-tier, openapi, llms-txt, python, typescript, streaming, batch, open-weights |\n| JSON | https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json |\n\n## Score breakdown (methodology v0.4, October 2026 research run)\n\nAssessed 2026-10-08 from public evidence against the published checklist (https://www.anchorterminal.com/benchmark/#checklist). Confidence: medium. Performance and Task success pending (no score, not in the total); the total is Σ(score × weight) ÷ 80 over the 7 assessed categories. \"This run\" is each category's share of the 100 points.\n\n| Category | Weight | This run | Score (0–100) | Points |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% | 20 | 53 | 10.6 |\n| Performance | 10% | pending | pending | n/a |\n| Schema \u0026 documentation | 13% | 16.2 | 85 | 13.8 |\n| Agent ergonomics | 13% | 16.2 | 77 | 12.5 |\n| Security \u0026 auth | 14% | 17.5 | 70 | 12.2 |\n| Payments \u0026 pricing | 10% | 12.5 | 40 | 5.0 |\n| Task success | 10% | pending | pending | n/a |\n| Maintenance \u0026 community | 7% | 8.8 | 66 | 5.8 |\n| Transparency \u0026 trust (editorial 73, provenance 91) | 7% | 8.8 | 82 | 7.2 |\n| Negative events | up to −15 | up to −15 | 2026-05-12. Compromised `mistralai` 2.4.6 on PyPI ran a credential-harvesting script on import for about three hours, and three `@mistralai/mistralai` versions on npm were also replaced. Mistral published advisory MAI-2026-002, removed the packages and closed its investigation on 14 May. Fixed and documented, so 3 points. https://docs.mistral.ai/resources/security-advisories/MAI-2026-002  | -3 |\n| **Total** | | | | **64.1 → B** |\n\n### Why each score\n\n- Reliability 53: Hosted reading. The docs point to status.mistral.ai, which our sibling Mistral listings read on 1 and 2 October 2026 as a Rootly page with component uptime bars (20). On 8 October the page answered curl and WebFetch with a Cloudflare check and HTTP 403, so the 90-day record for audio was not read (5, unread and not a finding against the vendor). The limits page names audio seconds a minute and audio seconds a month as the audio limits, with the numbers shown only in the Admin Panel (3 of 15). The error glossary tells callers to back off exponentially on 429, 500, 502, 503 and 504 and to read the `Retry-After` header, and the official SDKs retry with backoff. Transcription is a stateless call (15). No SLA found (0). The model cards mark Voxtral Mini Transcribe 2 and Voxtral Mini Transcribe Realtime as GA (10).\n- Performance: Pending. Latency is measured per call by our probes, which haven't run yet, so this run doesn't score it. Its weight is shared across the assessed categories until the first probe window closes.\n- Schema \u0026 documentation 85: Model reading. The public OpenAPI document covers `/v1/audio/transcriptions`, its SSE streaming form and `/v1/client/sessions`. The realtime WebSocket messages are described by SDK types and guide examples, not by the document (22 of 25). llms.txt with Markdown twins of the speech-to-text, offline, realtime and client authentication pages (10). The guides say which model each endpoint runs, that `timestamp_granularities` and `language` cannot be combined and that realtime cannot diarise, and carry an OpenAI SDK compatibility table. In the OpenAPI document `diarize`, `context_bias` and `temperature` have no description (14 of 20). `timestamp_granularities` is an enum of `segment` and `word`, `language` is a two-letter pattern and `context_bias` items have a pattern, while `model` is a free string (12 of 15). curl, Python and TypeScript examples on each guide and a shared error glossary. The OpenAPI response example still names `voxtral-mini-2507` (12 of 15). Dated model identifiers and a dated changelog (15).\n- Agent ergonomics 77: API reading, as for the other speech-to-text listings. Segments, word timestamps and speaker labels are returned only when asked for, and the default response is the text, the language and usage. No subtitle or plain-text output format was found (18 of 25). There is no transcript store to page through. Volume goes through the batch API, which accepts `/v1/audio/transcriptions` (12 of 20). Errors return `type`, `param` and `code` with a fix per status in the glossary (17 of 20). The call is stateless and safe to retry, the SDKs retry with backoff, and no idempotency key exists or is needed (15 of 20). `model` is the only required field beside the audio source, and there are official Python and TypeScript SDKs plus OpenAI SDK compatibility for the batch endpoint (15).\n- Security \u0026 auth 70: Model reading, scored like the other Mistral listings. Keys are bound to one workspace, can carry an expiry date and a connector scope, and organisations can enforce a maximum key lifetime. Browser clients get `rt_` tokens from `POST /v1/client/sessions` that last about 900 seconds, work for one model and travel in `Sec-WebSocket-Protocol`. A key cannot be limited to transcription (26 of 30). The commercial terms effective 25 September 2026 say Mistral will not train on customer data unless the customer opted in, or did not opt out where a product defaults to opt-in, without naming the API's default. The privacy and data controls page the docs link to returned 404 (10 of 20). Input and output are kept 30 rolling days for abuse monitoring, and `/v1/audio/transcriptions` is on the zero data retention list for paid plans, on request and subject to approval (12 of 15). Audit logs cover API key actions on Enterprise plans only (10 of 15). security.txt is valid until 5 May 2027 and names a HackerOne submission form, and the docs publish three security advisories for 2026. Certifications sit in a trust centre drawn by script that we could not read (12 of 20).\n- Payments \u0026 pricing 40: No x402, MPP or L402 (0). $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime on public pages (20). Free mode gives API access with no card, within limits shown in the console (20). A person signs up in a browser (0).\n- Task success: Pending. Task success needs the category task suites run through each tool, which haven't run yet, so this run doesn't score it. Its weight is shared across the assessed categories until then. A data provider's data-quality score is published on its listing now and becomes half of this category when it's scored.\n- Maintenance \u0026 community 66: Model reading. The last dated change to transcription is the 4 February 2026 release of both models with diarisation and context biasing, eight months ago. The SDKs that carry the realtime client were released on 6 October 2026 (Python 3.1.0), so we give 10 of 30, a departure from the strict reading of 0. The lifecycle policy promises six months' notice for GA models (12 of 12). No transcription model was retired in the last 90 days of the changelog (8 of 8). Five dated platform changelog entries between 16 July and 6 October, none about audio (8 of 15). The eight newest open issues on client-python have zero or one comment (5 of 10). Official `mistralai` 3.1.0 on PyPI and `@mistralai/mistralai` 2.7.0 on npm (15). Both are generated from the OpenAPI document, with test, lint and example workflows in the repository (8 of 10).\n- Transparency \u0026 trust 82: The hosted service is closed under commercial terms. The Realtime weights are published on Hugging Face under Apache-2.0 and the SDKs are Apache-2.0, while Voxtral Mini Transcribe 2 has no public weights (20 of 30). The privacy policy effective 3 September 2026 (30 rolling days), the zero data retention page and the terms agree on retention, and the terms leave the API's training default unnamed (20 of 30). A model lifecycle page with notice periods per stage (20). Opt-in EU and US regional endpoints are documented, the global endpoint commits to no location, and the provider list is in a trust centre we could not read (13 of 20).\n\nFix list for a coding agent, everything this grade says the listing lacks, the biggest gain first (18 items): https://www.anchorterminal.com/fixes/mistral-voxtral-transcribe.md (JSON https://www.anchorterminal.com/fixes/mistral-voxtral-transcribe.json)\n\n### What we couldn't check\n\n- unchecked: status.mistral.ai answered every request with a Cloudflare check (HTTP 403) on 8 October 2026, so the audio component's uptime and the 90-day incident record were not read\n- unchecked: trust.mistral.ai is drawn by script, so certifications and the provider list were not read\n- unchecked: the privacy and data controls page linked from the zero data retention page returned 404, so the API's training default was not established\n- unchecked: whether Free mode includes transcription minutes and whether sign-up needs a phone number. The quickstart says no card is needed\n- The pricing table shows $0.0003 a minute in a second column for Voxtral Mini Transcribe 2 without a label we could tie to batch or cached use, so it is not listed as a price\n- Whether zero data retention covers the realtime WebSocket at `/v1/audio/transcriptions/realtime`. The list names `/v1/audio/transcriptions` only\n- The transcription curl examples send the key in an `x-api-key` header, while the error glossary gives `Authorization: Bearer`\n- The other Mistral listings took no deduction for advisory MAI-2026-002. This one takes 3 because the realtime client is only reachable through the SDK\n\n### Sources\n\n- speech-to-text overview: \u003chttps://docs.mistral.ai/studio/audio/speech_to_text\u003e (seen 2026-10-08)\n- offline transcription guide: \u003chttps://docs.mistral.ai/studio/audio/speech_to_text/offline_transcription\u003e (seen 2026-10-08)\n- realtime transcription guide: \u003chttps://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription\u003e (seen 2026-10-08)\n- realtime client authentication: \u003chttps://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth\u003e (seen 2026-10-08)\n- OpenAPI document: \u003chttps://docs.mistral.ai/openapi.yaml\u003e (seen 2026-10-08)\n- Voxtral Mini Transcribe 2 model card and price: \u003chttps://docs.mistral.ai/models/voxtral-mini-transcribe-26-02\u003e (seen 2026-10-08)\n- Voxtral Mini Transcribe Realtime model card and price: \u003chttps://docs.mistral.ai/models/voxtral-mini-transcribe-realtime-26-02\u003e (seen 2026-10-08)\n- API pricing: \u003chttps://mistral.ai/pricing/api/\u003e (seen 2026-10-08)\n- changelog: \u003chttps://docs.mistral.ai/resources/changelogs\u003e (seen 2026-10-08)\n- model lifecycle policy: \u003chttps://docs.mistral.ai/inference/model-lifecycle\u003e (seen 2026-10-08)\n- error glossary: \u003chttps://docs.mistral.ai/resources/error-glossary\u003e (seen 2026-10-08)\n- usage and limits: \u003chttps://docs.mistral.ai/admin/billing-usage/usage-limits\u003e (seen 2026-10-08)\n- API keys: \u003chttps://docs.mistral.ai/admin/identity-access/api-keys\u003e (seen 2026-10-08)\n- zero data retention: \u003chttps://docs.mistral.ai/admin/monitor-comply/zero-data-retention\u003e (seen 2026-10-08)\n- audit logs: \u003chttps://docs.mistral.ai/admin/monitor-comply/audit-logs/overview\u003e (seen 2026-10-08)\n- security advisory MAI-2026-002: \u003chttps://docs.mistral.ai/resources/security-advisories/MAI-2026-002\u003e (seen 2026-10-08)\n- Studio MCP server: \u003chttps://docs.mistral.ai/resources/mcp\u003e (seen 2026-10-08)\n- commercial terms of service: \u003chttps://legal.mistral.ai/terms/commercial-terms-of-service\u003e (seen 2026-10-08)\n- privacy policy: \u003chttps://legal.mistral.ai/terms/privacy-policy\u003e (seen 2026-10-08)\n- security.txt: \u003chttps://mistral.ai/.well-known/security.txt\u003e (seen 2026-10-08)\n- Python SDK releases: \u003chttps://pypi.org/pypi/mistralai/json\u003e (seen 2026-10-08)\n- Python SDK repository: \u003chttps://github.com/mistralai/client-python\u003e (seen 2026-10-08)\n- Realtime open weights: \u003chttps://huggingface.co/mistralai/Voxtral-Mini-4B-Realtime-2602\u003e (seen 2026-10-08)\n\n## Who's behind it (provenance 91/100, checked 2026-10-08)\n\n| Check | Finding | Points |\n| --- | --- | --- |\n| Legal entity named | Mistral AI (RCS Paris 952 418 325) | 20/20 |\n| Domain age | mistral.ai, registered 2019-05-15 (7 years) | 11/15 |\n| Endpoint on the vendor's domain | api.mistral.ai | 15/15 |\n| Terms of service | read, states 6 of the 7 things a reader expects, and has 2 clauses that cost points | 5.1/10 |\n| Privacy policy | read, states 8 of the 8 things a reader expects | 10/10 |\n| Status page | status.mistral.ai | 10/10 |\n| Changelog | published | 10/10 |\n| security.txt | valid | 10/10 |\n\n### Terms and privacy, as read\n\nA reading by a fixed set of rules, each answered with the vendor's own sentence. Not legal advice.\n\n**Terms of service** (https://legal.mistral.ai/terms/commercial-terms-of-service), read 2026-10-08, dated 2026-09-25, states 6 of the 7 things a reader expects.\n\n- To know. Says it may use customer content to train or improve models, and gives an opt-out. \"…royalty-free, fully-paid license (with the right to sublicense to our service providers) to use Customer Data and Outputs solely as provided in the preceding sentence to train Mistral AI’s artificial intelligence models.\"\n- To know. Restricts benchmarking or competitive use (costs points). \"To the extent permitted by applicable law, Customer may not use image Outputs to develop or train any image generation product that competes with a Mistral AI Product.\"\n- To know. Says the terms or the service can change without notice (costs points). \"Any updates (i) that are not material or (ii) made for compliance with applicable law or to address a material security risk become effective immediately upon being posted at https://legal.mistral.ai/terms.\"\n- Gives the date it was last updated. Last updated 2026-09-25.\n- Names the governing law or courts. The law of the Republic of Singapore.\n- States a limit on its liability. Capped at the fees paid in the 12 months before the claim.\n- Says how changes to the terms are announced. Gives thirty days of notice before a change.\n- Not found in the text. Refers to a service level or uptime commitment.\n- Also in the text (2026-10-08). Customer data and outputs from Labs or Preview Models may be used for training, and opt-out choices made for other products, including zero data retention, do not apply. \"(i) Mistral AI may use Customer Data and Outputs generated from Labs or Preview Models to train its artificial intelligence models and (ii) the opt-out preferences you selected for other Mistral AI Products (including through zero data retention) does not apply to Labs or Preview Models.\"\n- Also in the text (2026-10-08). When the customer gives feedback, Mistral AI may use it together with the associated customer data and output in any manner without restriction. \"Mistral AI may use, copy, disclose, license, distribute, and exploit such Feedback, along with the Customer Data and Output associated with such Feedback, in any manner without any obligation, royalty, or restriction based on intellectual property rights or otherwise.\"\n- Also in the text (2026-10-08). The customer may not state or imply that output was generated by a human when the Mistral AI Products generated it. \"Customer may not represent or imply that the Output was generated by a human when it was generated by the Mistral AI Products.\"\n\n**Privacy policy** (https://legal.mistral.ai/terms/privacy-policy), read 2026-10-08, dated 2026-09-03, states 8 of the 8 things a reader expects.\n\n- To know. Says it may use customer content to train or improve models, and gives an opt-out. \"We’ve introduced a user control which allows you to object to the use of your input and output data for model training directly from your account.\"\n- Gives the date it was last updated. Last updated 2026-09-03.\n- Says how long data is kept. Names a period of 5 years.\n- Gives a privacy contact. Names a data protection officer.\n- Says where data is transferred or stored. Relies on standard contractual clauses.\n- Also in the text (2026-10-08). The policy does not apply where a business uses the products to process personal data, in which case Mistral AI is the processor. \"This Privacy Policy does not apply if you use our Mistral AI Products to process personal data in the context of your business activities.\"\n- Also in the text (2026-10-08). Mistral AI says that rights requests concerning model training have technical limitations and may involve a complex technical process. \"However, when your request concerns the training of our models, it’s important to note that your rights have technical limitations and fulfilling your requests might involve a complex technical process.\"\n\n## Live (updated 2026-10-09 11:46 UTC)\n\n- Right now: up, HTTP 404, 38 ms, checked 2026-10-09 11:46 UTC (get on `https://api.mistral.ai/v1`)\n- Uptime 24h 100.0% (44 probes) · 30 days 100.0% (44 probes) · p50 47 ms · p95 77 ms\n- Always current: https://www.anchorterminal.com/api/v1/live/mistral-voxtral-transcribe.json\n\n## Probe metrics\n\nNot measured yet. Our benchmark probes haven't run, so there's no availability, latency or error rate from a run and Performance is pending. Live uptime, where we poll the endpoint, is under Live and doesn't change the score.\n\n## Prices\n\n| Item | Price | Unit | Note |\n| --- | --- | --- | --- |\n| Voxtral Mini Transcribe 2 (batch) | $0.003 | per minute of audio |  |\n| Voxtral Mini Transcribe Realtime | $0.006 | per minute of audio |  |\n\nAcross all listings: https://www.anchorterminal.com/prices/index.md\n\n## Strengths\n\n- Published prices of $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime\n- One request with `model` and a `file`, `file_url` or `file_id`, for audio of up to about three hours\n- Browser clients use `rt_` tokens that last about 900 seconds and are limited to one model\n- `/v1/audio/transcriptions` is on the zero data retention list for paid plans, and the Realtime weights are Apache-2.0 on Hugging Face\n- OpenAPI document, llms.txt and Markdown guides, with a six-month retirement notice for GA models\n\n## Weaknesses\n\n- Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel\n- `timestamp_granularities` cannot be used with `language`, and the realtime model cannot diarise\n- 13 languages, and context biasing is tuned for English with other languages described as experimental\n- The compromised `mistralai` 2.4.6 on PyPI harvested credentials on import for three hours on 12 May 2026\n- No SLA found, and status.mistral.ai answered our requests with a bot check on 8 October 2026\n- The Studio MCP server has no transcription tool\n\n## Before you call it (notes for agents)\n\n1. Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`\n2. Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible\n3. Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`\n4. Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move\n5. Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key\n\n## Connect\n\nInstall:\n\n```bash\npip install mistralai   # realtime: pip install \"mistralai[realtime]\"   # or: npm i @mistralai/mistralai\n```\n\nFirst request:\n\n```bash\ncurl --location 'https://api.mistral.ai/v1/audio/transcriptions' \\\n  --header \"x-api-key: $MISTRAL_API_KEY\" \\\n  --form 'file_url=\"https://docs.mistral.ai/audio/obama.mp3\"' \\\n  --form 'model=\"voxtral-mini-latest\"'\n```\n\n## Similar tools\n\nRanked by shared capabilities, then score. Same-category tools with no shared capability key are listed last.\n\n| Tool | Grade | Score | Rank | Shared capabilities | x402 | Markdown |\n| --- | --- | --- | --- | --- | --- | --- |\n| Amazon Transcribe | BB | 73.4 | 86 | speech.stt, speech.streaming, speech.batch, speech.diarisation, speech.languages | no | https://www.anchorterminal.com/tools/amazon-transcribe.md |\n| Azure AI Speech speech-to-text | BB | 73 | 90 | speech.stt, speech.streaming, speech.batch, speech.diarisation, speech.languages | no | https://www.anchorterminal.com/tools/azure-speech-to-text.md |\n| Deepgram Speech-to-Text (Nova-3, Flux) | BB | 70.3 | 150 | speech.stt, speech.streaming, speech.batch, speech.diarisation, speech.languages | no | https://www.anchorterminal.com/tools/deepgram-stt.md |\n| Google Cloud Speech-to-Text | BB | 70.2 | 152 | speech.stt, speech.streaming, speech.batch, speech.diarisation, speech.languages | no | https://www.anchorterminal.com/tools/google-speech-to-text.md |\n| Gladia Speech-to-Text API + MCP | B | 69.5 | 171 | speech.stt, speech.streaming, speech.batch, speech.diarisation, speech.languages | no | https://www.anchorterminal.com/tools/gladia-stt.md |\n| ElevenLabs Scribe Speech to Text API | B | 68.9 | 189 | speech.stt, speech.streaming, speech.batch, speech.diarisation, speech.languages | no | https://www.anchorterminal.com/tools/elevenlabs-scribe.md |\n\n## Panel reviews (0)\n\nReviewed by the Anchor panel (https://www.anchorterminal.com/reviewers/index.md): .\n\nDesk reviews, written from public documentation, pricing, terms, source and status history on 1 October 2026. No calls made. For a desk review, the outcome says whether the reviewer's questions could be answered from public material: success, partial or failure. How reviews work: https://www.anchorterminal.com/reviews/how-it-works.md\n\n## Notable\n\n- Voxtral Mini Transcribe 2 (`voxtral-mini-2602`, alias `voxtral-mini-latest`) and Voxtral Mini Transcribe Realtime (`voxtral-mini-transcribe-realtime-2602`) were released on 2026-02-04 and are marked GA (source: \u003chttps://docs.mistral.ai/resources/changelogs\u003e)\n- Batch transcription accepts about three hours of audio in one request, as `file`, `file_url` or `file_id` (source: \u003chttps://docs.mistral.ai/studio/audio/speech_to_text/offline_transcription\u003e)\n- `timestamp_granularities` is not compatible with `language`, and realtime is not compatible with `diarize` (source: \u003chttps://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription\u003e)\n- Realtime runs over `wss://api.mistral.ai/v1/audio/transcriptions/realtime` with a `target_streaming_delay_ms` setting. Mistral says latency can be configured below 200 ms\n- The endpoint follows the OpenAI audio transcriptions request shape, with `diarize` and `context_bias` passed through `extra_body` and no `file_url` support in that mode\n- The Studio MCP server at `https://api.mistral.ai/mcp` has no transcription tool (source: \u003chttps://docs.mistral.ai/resources/mcp\u003e)\n\n## Compare\n\n- [Amazon Transcribe vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.md): BB 73.4 vs B 64.1\n- [AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.md): B 66.8 vs B 64.1\n- [Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.md): BB 73 vs B 64.1\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.md): BB 70.3 vs B 64.1\n- [ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.md): B 68.9 vs B 64.1\n- [Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.md): B 69.5 vs B 64.1\n- [Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.md): BB 70.2 vs B 64.1\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md): BB 71.8 vs B 64.1\n- [Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.md): B 64.1 vs C 57.8\n- [Mistral Voxtral Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.md): B 64.1 vs C 58.2\n- [Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.md): B 64.1 vs B 67.1\n\n## Verify this listing\n\nFor the vendor. The badge or a plain link to this page verifies the listing, from a page on mistral.ai or one of its subdomains, or the README of github.com/mistralai/client-python. It shows the listing is the vendor's and that the vendor knows it's here, and it never changes a grade, rank or review. The vendor sends the page's address to `POST https://www.anchorterminal.com/api/v1/verify` as `{\"slug\": \"mistral-voxtral-transcribe\", \"url\": \"…\"}`, or calls the `verify_listing` tool at https://www.anchorterminal.com/mcp. We fetch the page once, then again every week; two failed checks in a row and the verification lapses, and a later pass restores it. What we check: https://www.anchorterminal.com/builders/index.md#verify\n\nHTML badge:\n\n```html\n\u003ca href=\"https://www.anchorterminal.com/tools/mistral-voxtral-transcribe\"\u003e\u003cimg src=\"https://www.anchorterminal.com/badges/mistral-voxtral-transcribe.svg\" alt=\"Mistral Voxtral Transcribe on Anchor Terminal\" height=\"20\"\u003e\u003c/a\u003e\n```\n\nMarkdown badge, for a README:\n\n```markdown\n[![Mistral Voxtral Transcribe on Anchor Terminal](https://www.anchorterminal.com/badges/mistral-voxtral-transcribe.svg)](https://www.anchorterminal.com/tools/mistral-voxtral-transcribe)\n```\n\nPlain link:\n\n```html\n\u003ca href=\"https://www.anchorterminal.com/tools/mistral-voxtral-transcribe\"\u003eMistral Voxtral Transcribe on Anchor Terminal\u003c/a\u003e\n```\n\n## Share this listing\n\nFor the vendor. Sharing assets for social media, two PNGs of 1200 × 630 that say Mistral Voxtral Transcribe is listed on Anchor Terminal, with the vendor's logo and this page's address and no grade or score.\n\n- Dark: https://www.anchorterminal.com/assets/share/mistral-voxtral-transcribe-dark.png\n- Light: https://www.anchorterminal.com/assets/share/mistral-voxtral-transcribe-light.png\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Terminal",
        "url": "https://www.anchorterminal.com/tools/"
      },
      {
        "name": "Speech-to-text",
        "url": "https://www.anchorterminal.com/categories/speech-to-text"
      },
      {
        "name": "Mistral Voxtral Transcribe",
        "url": ""
      }
    ],
    "description": "Mistral AI's speech-to-text API. Voxtral Mini Transcribe 2 transcribes files of up to about three hours with diarisation, word timestamps and context biasing, and Voxtral Realtime transcribes live audio over a WebSocket.",
    "facts": [
      "rank #329 of 842",
      "API key auth",
      "0 desk reviews"
    ],
    "h1": "Mistral Voxtral Transcribe",
    "image": "https://www.anchorterminal.com/assets/og/tools-mistral-voxtral-transcribe.png",
    "path": "/tools/mistral-voxtral-transcribe",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Mistral Voxtral Transcribe review for AI agents, grade B (64.1/100)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe"
  },
  "tokens": {
    "markdown": 7650,
    "slim": 1530
  },
  "version": 1
}
