{
  "data": {
    "a": {
      "slug": "mistral-voxtral-transcribe",
      "name": "Mistral Voxtral Transcribe",
      "vendor": "Mistral AI",
      "vendorUrl": "https://mistral.ai",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Mistral AI's speech-to-text API. Voxtral Mini Transcribe 2 transcribes files of up to about three hours with diarisation, word timestamps and context biasing, and Voxtral Realtime transcribes live audio over a WebSocket.",
      "url": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe",
      "markdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json",
      "repo": "https://github.com/mistralai/client-python",
      "license": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.mistral.ai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "mistralai"
        },
        {
          "registry": "npm",
          "name": "@mistralai/mistralai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from Studio, bound to one workspace, with an optional expiry date. Browser realtime clients use `rt_` tokens minted at `POST /v1/client/sessions`, valid about 900 seconds for one model (https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth).",
      "pricing": "freemium",
      "pricingNotes": "$0.003 per audio minute for Voxtral Mini Transcribe 2 and $0.006 for Voxtral Mini Transcribe Realtime (model cards on docs.mistral.ai, https://mistral.ai/pricing/api/). Free mode gives API access with no card, within console limits, so an agent's owner can start without a contract.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 773,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.mistral.ai/studio/audio/speech_to_text",
      "llmsTxt": "https://docs.mistral.ai/llms.txt",
      "openapi": "https://docs.mistral.ai/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.streaming",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "eu",
        "free-tier",
        "openapi",
        "llms-txt",
        "python",
        "typescript",
        "streaming",
        "batch",
        "open-weights"
      ],
      "lastRelease": "2026-02-04",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 64.1,
        "grade": "B",
        "agentReady": false,
        "rank": 329,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 10,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 82
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": -3,
        "negativeNotes": [
          "2026-05-12. Compromised `mistralai` 2.4.6 on PyPI ran a credential-harvesting script on import for about three hours, and three `@mistralai/mistralai` versions on npm were also replaced. Mistral published advisory MAI-2026-002, removed the packages and closed its investigation on 14 May. Fixed and documented, so 3 points. https://docs.mistral.ai/resources/security-advisories/MAI-2026-002"
        ],
        "verdict": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.",
        "bestFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "strengths": [
          "Published prices of $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime",
          "One request with `model` and a `file`, `file_url` or `file_id`, for audio of up to about three hours",
          "Browser clients use `rt_` tokens that last about 900 seconds and are limited to one model",
          "`/v1/audio/transcriptions` is on the zero data retention list for paid plans, and the Realtime weights are Apache-2.0 on Hugging Face",
          "OpenAPI document, llms.txt and Markdown guides, with a six-month retirement notice for GA models"
        ],
        "weaknesses": [
          "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel",
          "`timestamp_granularities` cannot be used with `language`, and the realtime model cannot diarise",
          "13 languages, and context biasing is tuned for English with other languages described as experimental",
          "The compromised `mistralai` 2.4.6 on PyPI harvested credentials on import for three hours on 12 May 2026",
          "No SLA found, and status.mistral.ai answered our requests with a bot check on 8 October 2026",
          "The Studio MCP server has no transcription tool"
        ],
        "agentNotes": [
          "Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`",
          "Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible",
          "Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`",
          "Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move",
          "Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 64.1
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 73
        },
        "provenanceScore": 91
      },
      "connect": {
        "install": "pip install mistralai   # realtime: pip install \"mistralai[realtime]\"   # or: npm i @mistralai/mistralai",
        "http": "curl --location 'https://api.mistral.ai/v1/audio/transcriptions' \\\n  --header \"x-api-key: $MISTRAL_API_KEY\" \\\n  --form 'file_url=\"https://docs.mistral.ai/audio/obama.mp3\"' \\\n  --form 'model=\"voxtral-mini-latest\"'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/mistral-voxtral-transcribe"
      },
      "sameCompany": [
        "mistral-api",
        "mistral-embeddings",
        "mistral-moderation",
        "mistral-ocr"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Voxtral Mini Transcribe 2 (batch)",
          "unit": "audio-minute",
          "usd": 0.003
        },
        {
          "item": "Voxtral Mini Transcribe Realtime",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "Mistral AI (RCS Paris 952 418 325)",
        "domain": "mistral.ai",
        "domainRegistered": "2019-05-15",
        "endpointOnVendorDomain": true,
        "terms": "https://legal.mistral.ai/terms/commercial-terms-of-service",
        "privacy": "https://legal.mistral.ai/terms/privacy-policy",
        "statusPage": "https://status.mistral.ai",
        "changelog": "https://docs.mistral.ai/resources/changelogs",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 91
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.json",
      "live": {
        "slug": "mistral-voxtral-transcribe",
        "probe": {
          "target": "https://api.mistral.ai/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:00:30.276888458Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 57,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 48,
          "p95ms24h": 78,
          "samples24h": 36,
          "samples30d": 36,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 36,
              "ok": 36
            }
          ]
        },
        "updatedAt": "2026-10-09T11:00:30.276888458Z"
      }
    },
    "answer": "Mistral Voxtral Transcribe scores 64.1 (B) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on reliability.",
    "b": {
      "slug": "soniox-stt",
      "name": "Soniox Speech-to-Text",
      "vendor": "Soniox",
      "vendorUrl": "https://soniox.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "One multilingual model family for 60+ languages, as a real-time WebSocket API (`stt-rt-v5`) and an async file API (`stt-async-v5`).",
      "url": "https://www.anchorterminal.com/tools/soniox-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/soniox-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/soniox-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/soniox-stt.json",
      "repo": "https://github.com/soniox/soniox-python",
      "license": "Apache-2.0 (Python SDK)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.soniox.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "@soniox/node"
        },
        {
          "registry": "pypi",
          "name": "soniox"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key per project. Temporary API keys can be minted for browsers and mobile clients that stream straight to `wss://stt-rt.soniox.com`. Regional projects get their own keys and domains (EU, Japan, India).",
      "pricing": "usage",
      "pricingNotes": "Token-based pay-as-you-go. Async audio input $1.50 per 1M tokens and text in or out $3.50 per 1M, real-time $2.00 and $4.00. Soniox puts this at about $0.10 an hour async and $0.12 an hour real-time, with diarisation, language ID and translation included. New sign-ups have had no free credits since 2025-10-27 (https://soniox.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Billed to a funded account (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 12,
        "npmWeekly": 22200,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://soniox.com/docs/stt/get-started",
      "llmsTxt": "https://soniox.com/docs/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "closed-source",
        "python",
        "typescript",
        "llms-txt",
        "webhooks",
        "async-jobs",
        "streaming",
        "batch",
        "enterprise"
      ],
      "lastRelease": "2026-08-11",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 58.2,
        "grade": "C",
        "agentReady": false,
        "rank": 529,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 11,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 30,
          "payments": 20,
          "reliability": 65,
          "schema": 60,
          "security": 70,
          "transparency": 76
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included. No free credits for new accounts since October 2025.",
        "bestFor": "Price-led multilingual transcription and translation, live or async, and for operators who want no training and no retention.",
        "strengths": [
          "About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included",
          "Customer audio and transcripts are never used for training, and nothing is retained by default",
          "SOC 2 Type 2 and ISO/IEC 27001:2022",
          "Regional deployments in the US, EU, Japan and India",
          "Per-request usage logs with cost and request IDs"
        ],
        "weaknesses": [
          "No free credits for new accounts since October 2025",
          "No OpenAPI or AsyncAPI file",
          "No documented status code, Retry-After or backoff for rate limits",
          "10 concurrent streams and a fixed 300-minute cap per stream or file",
          "No STT changelog entry since June 2026"
        ],
        "agentNotes": [
          "Pass `audio_url` for public files and skip the upload step. Delete uploaded files or they count against the 10 GB quota for 30 days",
          "Use the `context` field for names and domain terms",
          "Buffer audio while the WebSocket connects, then flush it after the config message",
          "Split anything over 300 minutes. The cap is fixed",
          "Set `client_reference_id` so failed or duplicate requests can be traced in the usage log"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 58.2
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 30,
          "payments": 20,
          "reliability": 65,
          "schema": 60,
          "security": 70,
          "transparency": 70
        },
        "provenanceScore": 81
      },
      "connect": {
        "http": "curl https://api.soniox.com/v1/transcriptions -H \"Authorization: Bearer $SONIOX_API_KEY\" \\\n  -H \"content-type: application/json\" \\\n  -d '{\"model\":\"stt-async-v5\",\"audio_url\":\"https://soniox.com/media/examples/coffee_shop.mp3\",\"enable_speaker_diarization\":true}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/soniox-stt"
      },
      "sameCompany": [
        "soniox-tts",
        "soniox-voice-cloning"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "stt-async-v5",
          "unit": "audio-minute",
          "usd": 0.0017,
          "note": "Soniox's estimate of $0.10 an hour, token-billed"
        },
        {
          "item": "stt-rt-v5 streaming",
          "unit": "audio-minute",
          "usd": 0.002,
          "note": "Soniox's estimate of $0.12 an hour, token-billed"
        }
      ],
      "provenance": {
        "legalEntity": "Soniox Inc.",
        "domain": "soniox.com",
        "domainRegistered": "2020-03-23",
        "endpointOnVendorDomain": true,
        "terms": "https://soniox.com/policies/terms-of-service",
        "privacy": "https://soniox.com/policies/privacy-policy",
        "statusPage": "https://status.soniox.com",
        "changelog": "https://soniox.com/docs/stt/models",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms and privacy policy last updated 2026-06-29. The company address is Foster City, California"
        ],
        "score": 81
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/soniox-stt.json",
      "live": {
        "slug": "soniox-stt",
        "probe": {
          "target": "https://api.soniox.com/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:00:37.529212182Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 280,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 255,
          "p95ms24h": 289,
          "samples24h": 260,
          "samples30d": 2303,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 117,
              "ok": 117
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.soniox.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-09T07:58:32.968280938Z"
        },
        "versions": [
          {
            "registry": "npm",
            "name": "@soniox/node",
            "version": "2.3.0",
            "seenAt": "2026-10-08T16:29:50.412742768Z"
          },
          {
            "registry": "pypi",
            "name": "soniox",
            "version": "2.10.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-08T16:29:51.372607135Z"
          }
        ],
        "githubStars": 12,
        "npmWeekly": 26231,
        "pypiWeekly": 124290,
        "securityTxt": {
          "url": "https://soniox.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:38:50.110847376Z"
        },
        "llmsTxt": {
          "url": "https://soniox.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:52.768403714Z"
        },
        "domain": {
          "domain": "soniox.com",
          "registered": "2020-03-23",
          "source": "https://rdap.verisign.com/com/v1/domain/soniox.com",
          "checkedAt": "2026-10-04T13:05:00.531232044Z"
        },
        "pages": [
          {
            "url": "https://soniox.com/blog/2025-10-27-free-credits-update-for-soniox-api",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-08T18:24:29.452410806Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "84edd3460b10"
          },
          {
            "url": "https://soniox.com/docs/stt/models",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-08T18:24:31.648986358Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "bc23abf2e34e"
          },
          {
            "url": "https://soniox.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:24:39.904877842Z",
            "changedAt": "2026-10-02T15:24:18.034370139Z",
            "fingerprint": "3c5bced7754b"
          },
          {
            "url": "https://soniox.com/policies/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:24:35.536298045Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "f8449bf14df5"
          },
          {
            "url": "https://soniox.com/policies/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:24:37.831433566Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "1ab8c9d50b68"
          }
        ],
        "updatedAt": "2026-10-09T11:00:37.529212182Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Mistral AI",
        "b": "Soniox",
        "name": "Vendor"
      },
      {
        "a": "https://api.mistral.ai/v1",
        "b": "https://api.soniox.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, websocket",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "not published",
        "b": "$0.0017 per minute of audio",
        "name": "Price for speech stt"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
        "b": "Apache-2.0 (Python SDK)",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-02-04",
        "b": "2026-08-11",
        "name": "Last release"
      },
      {
        "a": "2026-09-25",
        "b": "2026-06-29",
        "name": "Terms last updated"
      },
      {
        "a": "2026-09-03",
        "b": "2026-06-29",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes, with an opt-out",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "773 stars",
        "b": "12 stars, 22k npm/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3.5/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Mistral Voxtral Transcribe scores 64.1 (B) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on reliability.",
        "question": "Which is better for AI agents, Mistral Voxtral Transcribe or Soniox Speech-to-Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Mistral Voxtral Transcribe and Soniox Speech-to-Text need an API key?"
      },
      {
        "answer": "Yes. Mistral Voxtral Transcribe has a hosted endpoint at https://api.mistral.ai/v1 and Soniox Speech-to-Text at https://api.soniox.com/v1.",
        "question": "Can an agent call Mistral Voxtral Transcribe and Soniox Speech-to-Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Schema \u0026 documentation, 85 against 60",
          "Agent ergonomics, 77 against 70",
          "Payments \u0026 pricing, 40 against 20",
          "Maintenance \u0026 community, 66 against 30",
          "Transparency \u0026 trust, 82 against 76"
        ],
        "also": null,
        "goodFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "slug": "mistral-voxtral-transcribe",
        "watchFor": "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel"
      },
      {
        "aheadOn": [
          "Reliability, 65 against 53"
        ],
        "also": [
          "No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them"
        ],
        "goodFor": "Price-led multilingual transcription and translation, live or async, and for operators who want no training and no retention.",
        "slug": "soniox-stt",
        "watchFor": "No free credits for new accounts since October 2025"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech stt"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.json",
        "title": "Amazon Transcribe vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-soniox-stt.json",
        "title": "Amazon Transcribe vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.json",
        "title": "Azure AI Speech speech-to-text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-soniox-stt.json",
        "title": "Gladia Speech-to-Text API + MCP vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.json",
        "title": "Google Cloud Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Groq Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.json",
        "title": "Groq Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.json",
        "title": "Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.json",
        "title": "Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-soniox-stt.json",
        "title": "Rev AI Speech-to-Text API vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.json",
        "title": "Soniox Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 12,
        "edge": "soniox-stt",
        "key": "reliability",
        "mistral-voxtral-transcribe": 53,
        "name": "Reliability",
        "soniox-stt": 65,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 25,
        "edge": "mistral-voxtral-transcribe",
        "key": "schema",
        "mistral-voxtral-transcribe": 85,
        "name": "Schema \u0026 documentation",
        "soniox-stt": 60,
        "weight": 13
      },
      {
        "by": 7,
        "edge": "mistral-voxtral-transcribe",
        "key": "ergonomics",
        "mistral-voxtral-transcribe": 77,
        "name": "Agent ergonomics",
        "soniox-stt": 70,
        "weight": 13
      },
      {
        "by": 0,
        "edge": "",
        "key": "security",
        "mistral-voxtral-transcribe": 70,
        "name": "Security \u0026 auth",
        "soniox-stt": 70,
        "weight": 14
      },
      {
        "by": 20,
        "edge": "mistral-voxtral-transcribe",
        "key": "payments",
        "mistral-voxtral-transcribe": 40,
        "name": "Payments \u0026 pricing",
        "soniox-stt": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 36,
        "edge": "mistral-voxtral-transcribe",
        "key": "maintenance",
        "mistral-voxtral-transcribe": 66,
        "name": "Maintenance \u0026 community",
        "soniox-stt": 30,
        "weight": 7
      },
      {
        "by": 6,
        "edge": "mistral-voxtral-transcribe",
        "key": "transparency",
        "mistral-voxtral-transcribe": 82,
        "name": "Transparency \u0026 trust",
        "soniox-stt": 76,
        "weight": 7
      }
    ],
    "summary": "Mistral Voxtral Transcribe scores 64.1 (B) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on reliability. Both do speech stt.",
    "verdicts": {
      "mistral-voxtral-transcribe": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.",
      "soniox-stt": "About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included. No free credits for new accounts since October 2025."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt",
    "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.md",
    "slim": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.min.md"
  },
  "markdown": "Mistral Voxtral Transcribe scores 64.1 (B) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on reliability. Both do speech stt.\n\n- Mistral Voxtral Transcribe: grade B, 64.1/100, rank #329 of 842. Markdown https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md · JSON https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json\n- Soniox Speech-to-Text: grade C, 58.2/100, rank #529 of 842. Markdown https://www.anchorterminal.com/tools/soniox-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/soniox-stt.json\n\n## Which one, for what\n\n### Mistral Voxtral Transcribe (B)\n\nGood for: Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.\n\nAhead on:\n- Schema \u0026 documentation, 85 against 60\n- Agent ergonomics, 77 against 70\n- Payments \u0026 pricing, 40 against 20\n- Maintenance \u0026 community, 66 against 30\n- Transparency \u0026 trust, 82 against 76\n\nWatch for: Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel\n\n### Soniox Speech-to-Text (C)\n\nGood for: Price-led multilingual transcription and translation, live or async, and for operators who want no training and no retention.\n\nAhead on:\n- Reliability, 65 against 53\n\nAlso in its favour:\n- No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them\n\nWatch for: No free credits for new accounts since October 2025\n\n\n## Score by category\n\n| Category | Weight | Mistral Voxtral Transcribe | Soniox Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 53 | 65 | Soniox Speech-to-Text +12 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 85 | 60 | Mistral Voxtral Transcribe +25 |\n| Agent ergonomics | 13% (16.2 this run) | 77 | 70 | Mistral Voxtral Transcribe +7 |\n| Security \u0026 auth | 14% (17.5 this run) | 70 | 70 | even |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | Mistral Voxtral Transcribe +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 66 | 30 | Mistral Voxtral Transcribe +36 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 82 | 76 | Mistral Voxtral Transcribe +6 |\n| Negative events | ≤15 | -3 | 0 | |\n| **Total** | | **64.1 · B** | **58.2 · C** | |\n\n## Facts side by side\n\n| Fact | Mistral Voxtral Transcribe | Soniox Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Mistral AI | Soniox |\n| Hosted endpoint | `https://api.mistral.ai/v1` | `https://api.soniox.com/v1` |\n| Transports | HTTP, websocket | HTTP |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| Price for speech stt | not published | $0.0017 per minute of audio |\n| x402 | no | no |\n| Licence | Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0 | Apache-2.0 (Python SDK) |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-02-04 | 2026-08-11 |\n| Terms last updated | 2026-09-25 | 2026-06-29 |\n| Privacy policy last updated | 2026-09-03 | 2026-06-29 |\n| Customer content may train models | yes, with an opt-out | not found in the text |\n| Terms restrict automated access | not found in the text | not found in the text |\n| Terms restrict benchmarking | yes | yes |\n| Terms or service can change without notice | yes | yes |\n| Arbitration or class-action waiver | not found in the text | not found in the text |\n| Popularity | 773 stars | 12 stars, 22k npm/wk |\n| Agent reviews | none | 3.5/5 (2) |\n\n## Verdicts\n\n**Mistral Voxtral Transcribe.** Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.\n\n**Soniox Speech-to-Text.** About $0.10 an hour async and $0.12 real time, with diarisation, language ID and translation included. No free credits for new accounts since October 2025.\n\n## Before you call either\n\n### Mistral Voxtral Transcribe\n\n1. Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`\n2. Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible\n3. Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`\n4. Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move\n5. Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key\n\n### Soniox Speech-to-Text\n\n1. Pass `audio_url` for public files and skip the upload step. Delete uploaded files or they count against the 10 GB quota for 30 days\n2. Use the `context` field for names and domain terms\n3. Buffer audio while the WebSocket connects, then flush it after the config message\n4. Split anything over 300 minutes. The cap is fixed\n5. Set `client_reference_id` so failed or duplicate requests can be traced in the usage log\n\n## Questions\n\n### Which is better for AI agents, Mistral Voxtral Transcribe or Soniox Speech-to-Text?\n\nMistral Voxtral Transcribe scores 64.1 (B) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on reliability.\n\n### Do Mistral Voxtral Transcribe and Soniox Speech-to-Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call Mistral Voxtral Transcribe and Soniox Speech-to-Text without installing anything?\n\nYes. Mistral Voxtral Transcribe has a hosted endpoint at https://api.mistral.ai/v1 and Soniox Speech-to-Text at https://api.soniox.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.json, and with the fewest tokens: https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"mistral-voxtral-transcribe\", \"b\": \"soniox-stt\"}`. From a terminal: `anchor compare mistral-voxtral-transcribe soniox-stt`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json and https://www.anchorterminal.com/api/v1/tools/soniox-stt.json\n\n## Other comparisons with Mistral Voxtral Transcribe or Soniox Speech-to-Text\n\n- [Amazon Transcribe vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.md)\n- [Amazon Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-soniox-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt.md)\n- [Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Azure AI Speech speech-to-text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-soniox-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.md)\n- [ElevenLabs Scribe Speech to Text API vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt.md)\n- [Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.md)\n- [Gladia Speech-to-Text API + MCP vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-soniox-stt.md)\n- [Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Google Cloud Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.md)\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Groq Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.md)\n- [Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.md)\n- [Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.md)\n- [Rev AI Speech-to-Text API vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/rev-ai-stt-vs-soniox-stt.md)\n- [Soniox Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Mistral Voxtral Transcribe vs Soniox Speech-to-Text",
        "url": ""
      }
    ],
    "description": "Mistral Voxtral Transcribe scores 64.1 (B) on agent readiness against Soniox Speech-to-Text's 58.2 (C), and leads in 5 of 7 scored categories. Soniox Speech-to-Text leads on reliability. Both do speech stt. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Mistral Voxtral Transcribe B 64.1",
      "Soniox Speech-to-Text C 58.2",
      "scores"
    ],
    "h1": "Mistral Voxtral Transcribe vs Soniox Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-mistral-voxtral-transcribe-vs-soniox-stt.png",
    "path": "/compare/mistral-voxtral-transcribe-vs-soniox-stt",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Mistral Voxtral Transcribe vs Soniox Speech-to-Text for AI agents",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt"
  },
  "tokens": {
    "markdown": 2550,
    "slim": 730
  },
  "version": 1
}
