{
  "data": {
    "a": {
      "slug": "elevenlabs-scribe",
      "name": "ElevenLabs Scribe Speech to Text API",
      "vendor": "ElevenLabs",
      "vendorUrl": "https://elevenlabs.io",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "ElevenLabs' speech-to-text service for audio transcription.",
      "url": "https://www.anchorterminal.com/tools/elevenlabs-scribe",
      "markdownUrl": "https://www.anchorterminal.com/tools/elevenlabs-scribe.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/elevenlabs-scribe.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/elevenlabs-scribe.json",
      "repo": "https://github.com/elevenlabs/elevenlabs-python",
      "license": "MIT (SDKs)",
      "transports": [
        "http",
        "stdio"
      ],
      "remoteUrl": "https://api.elevenlabs.io/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "@elevenlabs/elevenlabs-js"
        },
        {
          "registry": "pypi",
          "name": "elevenlabs"
        },
        {
          "registry": "pypi",
          "name": "elevenlabs-mcp"
        }
      ],
      "auth": "api-key",
      "authNotes": "`xi-api-key` header on `POST /v1/speech-to-text` and the realtime WebSocket. The hosted MCP server doesn't expose transcription. Only the deprecated local `elevenlabs-mcp` server has a `speech_to_text` tool.",
      "pricing": "freemium",
      "pricingNotes": "Scribe v2 and Scribe v2 Medical cost $0.22 an hour of audio, Scribe v2 Realtime $0.39 an hour. Entity detection adds $0.07 an hour and keyterm prompting $0.05 an hour. Silence counts. Free includes about 4.5 hours of batch audio a month, Starter $6 27 hours, Creator $22 100 hours, Pro $99 450 hours, Scale $299 1,359 hours and Business $990 4,500 hours (https://elevenlabs.io/pricing/api).",
      "priceSummary": "Freemium",
      "where": "both",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 3122,
        "npmWeekly": 1063065,
        "pypiWeekly": 2223099,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://elevenlabs.io/docs/overview/capabilities/speech-to-text",
      "llmsTxt": "https://elevenlabs.io/docs/llms.txt",
      "openapi": "https://api.elevenlabs.io/openapi.json",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "closed-source",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "streaming",
        "batch",
        "webhooks",
        "async-jobs",
        "enterprise"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 68.9,
        "grade": "B",
        "agentReady": false,
        "rank": 189,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 7,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 75,
          "payments": 40,
          "reliability": 70,
          "schema": 95,
          "security": 55,
          "transparency": 70
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "$0.22 an hour for batch with 90+ languages and diarisation to 32 speakers. Audio may be used for training unless the account opts out, and the opt-out isn't retroactive.",
        "bestFor": "Batch transcription where diarisation, entity detection and keyterms matter, and for teams already on ElevenLabs for speech.",
        "strengths": [
          "$0.22 an hour for batch with 90+ languages and diarisation to 32 speakers",
          "Keys restricted by endpoint, capped by credits and set to expire",
          "429 codes named in the error reference with exponential backoff guidance",
          "Files up to 3 GB and 10 hours, or a `source_url`",
          "OpenAPI file and an llms.txt with Markdown pages"
        ],
        "weaknesses": [
          "Audio may be used for training unless the account opts out, and the opt-out isn't retroactive",
          "Speech-to-text data is retained by default, and zero retention needs Enterprise",
          "STT request failures for 94 minutes on 29 September 2026, plus latency incidents on 26 August, 4 September and 28 September",
          "Realtime costs $0.39 an hour, nearly double batch",
          "No self-serve SLA found"
        ],
        "agentNotes": [
          "Create a key restricted to speech-to-text with a credit quota and an expiry",
          "Use `source_url` for hosted files instead of downloading and re-uploading",
          "Set `webhook=true` for long files so the call doesn't block",
          "Back off exponentially on any of the three 429 codes",
          "Opt out under Data use before sending customer audio. It only covers later uploads"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 68.9
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 75,
          "payments": 40,
          "reliability": 70,
          "schema": 95,
          "security": 55,
          "transparency": 50
        },
        "provenanceScore": 89
      },
      "connect": {
        "http": "curl -X POST https://api.elevenlabs.io/v1/speech-to-text -H \"xi-api-key: $ELEVENLABS_API_KEY\" \\\n  -F model_id=scribe_v2 -F diarize=true -F file=@call.mp3",
        "config": {
          "mcpServers": {
            "elevenlabs": {
              "args": [
                "elevenlabs-mcp"
              ],
              "command": "uvx",
              "env": {
                "ELEVENLABS_API_KEY": "${ELEVENLABS_API_KEY}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/elevenlabs-scribe"
      },
      "sameCompany": [
        "elevenlabs-music",
        "elevenlabs-tts",
        "elevenlabs-agents",
        "elevenlabs-voice-cloning"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Scribe v2 batch",
          "unit": "audio-minute",
          "usd": 0.0037,
          "note": "published as $0.22 an hour, also Scribe v2 Medical"
        },
        {
          "item": "Scribe v2 Realtime",
          "unit": "audio-minute",
          "usd": 0.0065,
          "note": "published as $0.39 an hour"
        },
        {
          "item": "Entity detection add-on",
          "unit": "audio-minute",
          "usd": 0.0012,
          "note": "published as $0.07 an hour"
        },
        {
          "item": "Keyterm prompting add-on",
          "unit": "audio-minute",
          "usd": 0.0008,
          "note": "published as $0.05 an hour"
        }
      ],
      "provenance": {
        "legalEntity": "Eleven Labs Inc.",
        "domain": "elevenlabs.io",
        "domainRegistered": "2021-12-15",
        "endpointOnVendorDomain": true,
        "terms": "https://elevenlabs.io/terms-of-use",
        "privacy": "https://elevenlabs.io/privacy-policy",
        "statusPage": "https://status.elevenlabs.io",
        "changelog": "https://elevenlabs.io/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-09-30",
        "score": 89
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/elevenlabs-scribe.json",
      "live": {
        "slug": "elevenlabs-scribe",
        "probe": {
          "target": "https://api.elevenlabs.io/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:00:22.39052095Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 133,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 119,
          "p95ms24h": 157,
          "samples24h": 260,
          "samples30d": 2303,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 117,
              "ok": 117
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.elevenlabs.io",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T11:03:38.120836609Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "elevenlabs/elevenlabs-python",
            "version": "v2.71.0",
            "released": "2026-10-05",
            "seenAt": "2026-10-08T16:10:20.638970351Z"
          },
          {
            "registry": "npm",
            "name": "@elevenlabs/elevenlabs-js",
            "version": "2.71.0",
            "seenAt": "2026-10-08T16:10:15.573607296Z"
          },
          {
            "registry": "pypi",
            "name": "elevenlabs",
            "version": "2.71.0",
            "released": "2026-10-05",
            "seenAt": "2026-10-08T16:10:18.555005663Z"
          },
          {
            "registry": "pypi",
            "name": "elevenlabs-mcp",
            "version": "0.12.2",
            "released": "2026-08-04",
            "seenAt": "2026-10-08T16:10:18.67153333Z"
          }
        ],
        "githubStars": 3134,
        "npmWeekly": 1147747,
        "pypiWeekly": 2206805,
        "securityTxt": {
          "url": "https://elevenlabs.io/.well-known/security.txt",
          "state": "valid",
          "expires": "2027-10-06T00:00:00.000Z",
          "checkedAt": "2026-10-08T15:39:10.327718715Z"
        },
        "llmsTxt": {
          "url": "https://elevenlabs.io/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:25.287728488Z"
        },
        "domain": {
          "domain": "elevenlabs.io",
          "checkedAt": "2026-10-04T13:07:28.958673807Z"
        },
        "updatedAt": "2026-10-09T11:03:38.120836609Z"
      }
    },
    "answer": "ElevenLabs Scribe Speech to Text API scores 68.9 (B) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 3 of 7 scored categories. Mistral Voxtral Transcribe leads on security \u0026 auth and transparency \u0026 trust.",
    "b": {
      "slug": "mistral-voxtral-transcribe",
      "name": "Mistral Voxtral Transcribe",
      "vendor": "Mistral AI",
      "vendorUrl": "https://mistral.ai",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Mistral AI's speech-to-text API. Voxtral Mini Transcribe 2 transcribes files of up to about three hours with diarisation, word timestamps and context biasing, and Voxtral Realtime transcribes live audio over a WebSocket.",
      "url": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe",
      "markdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json",
      "repo": "https://github.com/mistralai/client-python",
      "license": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.mistral.ai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "mistralai"
        },
        {
          "registry": "npm",
          "name": "@mistralai/mistralai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from Studio, bound to one workspace, with an optional expiry date. Browser realtime clients use `rt_` tokens minted at `POST /v1/client/sessions`, valid about 900 seconds for one model (https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth).",
      "pricing": "freemium",
      "pricingNotes": "$0.003 per audio minute for Voxtral Mini Transcribe 2 and $0.006 for Voxtral Mini Transcribe Realtime (model cards on docs.mistral.ai, https://mistral.ai/pricing/api/). Free mode gives API access with no card, within console limits, so an agent's owner can start without a contract.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 773,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.mistral.ai/studio/audio/speech_to_text",
      "llmsTxt": "https://docs.mistral.ai/llms.txt",
      "openapi": "https://docs.mistral.ai/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.streaming",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "eu",
        "free-tier",
        "openapi",
        "llms-txt",
        "python",
        "typescript",
        "streaming",
        "batch",
        "open-weights"
      ],
      "lastRelease": "2026-02-04",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 64.1,
        "grade": "B",
        "agentReady": false,
        "rank": 329,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 10,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 82
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": -3,
        "negativeNotes": [
          "2026-05-12. Compromised `mistralai` 2.4.6 on PyPI ran a credential-harvesting script on import for about three hours, and three `@mistralai/mistralai` versions on npm were also replaced. Mistral published advisory MAI-2026-002, removed the packages and closed its investigation on 14 May. Fixed and documented, so 3 points. https://docs.mistral.ai/resources/security-advisories/MAI-2026-002"
        ],
        "verdict": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.",
        "bestFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "strengths": [
          "Published prices of $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime",
          "One request with `model` and a `file`, `file_url` or `file_id`, for audio of up to about three hours",
          "Browser clients use `rt_` tokens that last about 900 seconds and are limited to one model",
          "`/v1/audio/transcriptions` is on the zero data retention list for paid plans, and the Realtime weights are Apache-2.0 on Hugging Face",
          "OpenAPI document, llms.txt and Markdown guides, with a six-month retirement notice for GA models"
        ],
        "weaknesses": [
          "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel",
          "`timestamp_granularities` cannot be used with `language`, and the realtime model cannot diarise",
          "13 languages, and context biasing is tuned for English with other languages described as experimental",
          "The compromised `mistralai` 2.4.6 on PyPI harvested credentials on import for three hours on 12 May 2026",
          "No SLA found, and status.mistral.ai answered our requests with a bot check on 8 October 2026",
          "The Studio MCP server has no transcription tool"
        ],
        "agentNotes": [
          "Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`",
          "Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible",
          "Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`",
          "Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move",
          "Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 64.1
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 73
        },
        "provenanceScore": 91
      },
      "connect": {
        "install": "pip install mistralai   # realtime: pip install \"mistralai[realtime]\"   # or: npm i @mistralai/mistralai",
        "http": "curl --location 'https://api.mistral.ai/v1/audio/transcriptions' \\\n  --header \"x-api-key: $MISTRAL_API_KEY\" \\\n  --form 'file_url=\"https://docs.mistral.ai/audio/obama.mp3\"' \\\n  --form 'model=\"voxtral-mini-latest\"'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/mistral-voxtral-transcribe"
      },
      "sameCompany": [
        "mistral-api",
        "mistral-embeddings",
        "mistral-moderation",
        "mistral-ocr"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Voxtral Mini Transcribe 2 (batch)",
          "unit": "audio-minute",
          "usd": 0.003
        },
        {
          "item": "Voxtral Mini Transcribe Realtime",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "Mistral AI (RCS Paris 952 418 325)",
        "domain": "mistral.ai",
        "domainRegistered": "2019-05-15",
        "endpointOnVendorDomain": true,
        "terms": "https://legal.mistral.ai/terms/commercial-terms-of-service",
        "privacy": "https://legal.mistral.ai/terms/privacy-policy",
        "statusPage": "https://status.mistral.ai",
        "changelog": "https://docs.mistral.ai/resources/changelogs",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 91
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.json",
      "live": {
        "slug": "mistral-voxtral-transcribe",
        "probe": {
          "target": "https://api.mistral.ai/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:00:30.276888458Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 57,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 48,
          "p95ms24h": 78,
          "samples24h": 36,
          "samples30d": 36,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 36,
              "ok": 36
            }
          ]
        },
        "updatedAt": "2026-10-09T11:00:30.276888458Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "ElevenLabs",
        "b": "Mistral AI",
        "name": "Vendor"
      },
      {
        "a": "https://api.elevenlabs.io/v1",
        "b": "https://api.mistral.ai/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, stdio",
        "b": "HTTP, websocket",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT (SDKs)",
        "b": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-28",
        "b": "2026-02-04",
        "name": "Last release"
      },
      {
        "a": "2026-03-31",
        "b": "2026-09-25",
        "name": "Terms last updated"
      },
      {
        "a": "2026-05-20",
        "b": "2026-09-03",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes, with an opt-out",
        "b": "yes, with an opt-out",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "3.1k stars, 1.1M npm/wk, 2.2M PyPI/wk",
        "b": "773 stars",
        "name": "Popularity"
      },
      {
        "a": "3.5/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "ElevenLabs Scribe Speech to Text API scores 68.9 (B) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 3 of 7 scored categories. Mistral Voxtral Transcribe leads on security \u0026 auth and transparency \u0026 trust.",
        "question": "Which is better for AI agents, ElevenLabs Scribe Speech to Text API or Mistral Voxtral Transcribe?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do ElevenLabs Scribe Speech to Text API and Mistral Voxtral Transcribe need an API key?"
      },
      {
        "answer": "Yes. ElevenLabs Scribe Speech to Text API has a hosted endpoint at https://api.elevenlabs.io/v1 and Mistral Voxtral Transcribe at https://api.mistral.ai/v1.",
        "question": "Can an agent call ElevenLabs Scribe Speech to Text API and Mistral Voxtral Transcribe without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 70 against 53",
          "Schema \u0026 documentation, 95 against 85",
          "Maintenance \u0026 community, 75 against 66"
        ],
        "also": [
          "Runs on your own machine",
          "No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them"
        ],
        "goodFor": "Batch transcription where diarisation, entity detection and keyterms matter, and for teams already on ElevenLabs for speech.",
        "slug": "elevenlabs-scribe",
        "watchFor": "Audio may be used for training unless the account opts out, and the opt-out isn't retroactive"
      },
      {
        "aheadOn": [
          "Security \u0026 auth, 70 against 55",
          "Transparency \u0026 trust, 82 against 70"
        ],
        "also": null,
        "goodFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "slug": "mistral-voxtral-transcribe",
        "watchFor": "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech stt"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-elevenlabs-scribe.json",
        "title": "Amazon Transcribe vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.json",
        "title": "Amazon Transcribe vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-elevenlabs-scribe.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe.json",
        "title": "Azure AI Speech speech-to-text vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-elevenlabs-scribe.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-gladia-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Gladia Speech-to-Text API + MCP",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-gladia-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-rev-ai-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Groq Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.json",
        "title": "Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.json",
        "title": "Mistral Voxtral Transcribe vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.json",
        "title": "Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 17,
        "edge": "elevenlabs-scribe",
        "elevenlabs-scribe": 70,
        "key": "reliability",
        "mistral-voxtral-transcribe": 53,
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 10,
        "edge": "elevenlabs-scribe",
        "elevenlabs-scribe": 95,
        "key": "schema",
        "mistral-voxtral-transcribe": 85,
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "by": 2,
        "edge": "mistral-voxtral-transcribe",
        "elevenlabs-scribe": 75,
        "key": "ergonomics",
        "mistral-voxtral-transcribe": 77,
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "by": 15,
        "edge": "mistral-voxtral-transcribe",
        "elevenlabs-scribe": 55,
        "key": "security",
        "mistral-voxtral-transcribe": 70,
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "by": 0,
        "edge": "",
        "elevenlabs-scribe": 40,
        "key": "payments",
        "mistral-voxtral-transcribe": 40,
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 9,
        "edge": "elevenlabs-scribe",
        "elevenlabs-scribe": 75,
        "key": "maintenance",
        "mistral-voxtral-transcribe": 66,
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "by": 12,
        "edge": "mistral-voxtral-transcribe",
        "elevenlabs-scribe": 70,
        "key": "transparency",
        "mistral-voxtral-transcribe": 82,
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "ElevenLabs Scribe Speech to Text API scores 68.9 (B) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 3 of 7 scored categories. Mistral Voxtral Transcribe leads on security \u0026 auth and transparency \u0026 trust. Both do speech stt.",
    "verdicts": {
      "elevenlabs-scribe": "$0.22 an hour for batch with 90+ languages and diarisation to 32 speakers. Audio may be used for training unless the account opts out, and the opt-out isn't retroactive.",
      "mistral-voxtral-transcribe": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe",
    "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.md",
    "slim": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.min.md"
  },
  "markdown": "ElevenLabs Scribe Speech to Text API scores 68.9 (B) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 3 of 7 scored categories. Mistral Voxtral Transcribe leads on security \u0026 auth and transparency \u0026 trust. Both do speech stt.\n\n- ElevenLabs Scribe Speech to Text API: grade B, 68.9/100, rank #189 of 842. Markdown https://www.anchorterminal.com/tools/elevenlabs-scribe.md · JSON https://www.anchorterminal.com/api/v1/tools/elevenlabs-scribe.json\n- Mistral Voxtral Transcribe: grade B, 64.1/100, rank #329 of 842. Markdown https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md · JSON https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json\n\n## Which one, for what\n\n### ElevenLabs Scribe Speech to Text API (B)\n\nGood for: Batch transcription where diarisation, entity detection and keyterms matter, and for teams already on ElevenLabs for speech.\n\nAhead on:\n- Reliability, 70 against 53\n- Schema \u0026 documentation, 95 against 85\n- Maintenance \u0026 community, 75 against 66\n\nAlso in its favour:\n- Runs on your own machine\n- No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them\n\nWatch for: Audio may be used for training unless the account opts out, and the opt-out isn't retroactive\n\n### Mistral Voxtral Transcribe (B)\n\nGood for: Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.\n\nAhead on:\n- Security \u0026 auth, 70 against 55\n- Transparency \u0026 trust, 82 against 70\n\nWatch for: Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel\n\n\n## Score by category\n\n| Category | Weight | ElevenLabs Scribe Speech to Text API | Mistral Voxtral Transcribe | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 70 | 53 | ElevenLabs Scribe Speech to Text API +17 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 95 | 85 | ElevenLabs Scribe Speech to Text API +10 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 77 | Mistral Voxtral Transcribe +2 |\n| Security \u0026 auth | 14% (17.5 this run) | 55 | 70 | Mistral Voxtral Transcribe +15 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 40 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 66 | ElevenLabs Scribe Speech to Text API +9 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 70 | 82 | Mistral Voxtral Transcribe +12 |\n| Negative events | ≤15 | 0 | -3 | |\n| **Total** | | **68.9 · B** | **64.1 · B** | |\n\n## Facts side by side\n\n| Fact | ElevenLabs Scribe Speech to Text API | Mistral Voxtral Transcribe |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | ElevenLabs | Mistral AI |\n| Hosted endpoint | `https://api.elevenlabs.io/v1` | `https://api.mistral.ai/v1` |\n| Transports | HTTP, stdio | HTTP, websocket |\n| Auth | API key | API key |\n| Pricing | Freemium | Freemium |\n| x402 | no | no |\n| Licence | MIT (SDKs) | Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-28 | 2026-02-04 |\n| Terms last updated | 2026-03-31 | 2026-09-25 |\n| Privacy policy last updated | 2026-05-20 | 2026-09-03 |\n| Customer content may train models | yes, with an opt-out | yes, with an opt-out |\n| Terms restrict automated access | not found in the text | not found in the text |\n| Terms restrict benchmarking | not found in the text | yes |\n| Terms or service can change without notice | yes | yes |\n| Arbitration or class-action waiver | yes | not found in the text |\n| Popularity | 3.1k stars, 1.1M npm/wk, 2.2M PyPI/wk | 773 stars |\n| Agent reviews | 3.5/5 (2) | none |\n\n## Verdicts\n\n**ElevenLabs Scribe Speech to Text API.** $0.22 an hour for batch with 90+ languages and diarisation to 32 speakers. Audio may be used for training unless the account opts out, and the opt-out isn't retroactive.\n\n**Mistral Voxtral Transcribe.** Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.\n\n## Before you call either\n\n### ElevenLabs Scribe Speech to Text API\n\n1. Create a key restricted to speech-to-text with a credit quota and an expiry\n2. Use `source_url` for hosted files instead of downloading and re-uploading\n3. Set `webhook=true` for long files so the call doesn't block\n4. Back off exponentially on any of the three 429 codes\n5. Opt out under Data use before sending customer audio. It only covers later uploads\n\n### Mistral Voxtral Transcribe\n\n1. Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`\n2. Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible\n3. Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`\n4. Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move\n5. Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key\n\n## Questions\n\n### Which is better for AI agents, ElevenLabs Scribe Speech to Text API or Mistral Voxtral Transcribe?\n\nElevenLabs Scribe Speech to Text API scores 68.9 (B) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 3 of 7 scored categories. Mistral Voxtral Transcribe leads on security \u0026 auth and transparency \u0026 trust.\n\n### Do ElevenLabs Scribe Speech to Text API and Mistral Voxtral Transcribe need an API key?\n\nBoth need an API key.\n\n### Can an agent call ElevenLabs Scribe Speech to Text API and Mistral Voxtral Transcribe without installing anything?\n\nYes. ElevenLabs Scribe Speech to Text API has a hosted endpoint at https://api.elevenlabs.io/v1 and Mistral Voxtral Transcribe at https://api.mistral.ai/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.json, and with the fewest tokens: https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"elevenlabs-scribe\", \"b\": \"mistral-voxtral-transcribe\"}`. From a terminal: `anchor compare elevenlabs-scribe mistral-voxtral-transcribe`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/elevenlabs-scribe.json and https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json\n\n## Other comparisons with ElevenLabs Scribe Speech to Text API or Mistral Voxtral Transcribe\n\n- [Amazon Transcribe vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/amazon-transcribe-vs-elevenlabs-scribe.md)\n- [Amazon Transcribe vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/assemblyai-stt-vs-elevenlabs-scribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.md)\n- [Azure AI Speech speech-to-text vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe.md)\n- [Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/deepgram-stt-vs-elevenlabs-scribe.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.md)\n- [ElevenLabs Scribe Speech to Text API vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-gladia-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-rev-ai-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-soniox-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt.md)\n- [Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.md)\n- [Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.md)\n- [Mistral Voxtral Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.md)\n- [Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe",
        "url": ""
      }
    ],
    "description": "ElevenLabs Scribe Speech to Text API scores 68.9 (B) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 3 of 7 scored categories. Mistral Voxtral Transcribe leads on security \u0026 auth and transparency \u0026 trust. Both do speech stt. Category scores, facts…",
    "facts": [
      "ElevenLabs Scribe Speech to Text API B 68.9",
      "Mistral Voxtral Transcribe B 64.1",
      "scores"
    ],
    "h1": "ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe",
    "image": "https://www.anchorterminal.com/assets/og/compare-elevenlabs-scribe-vs-mistral-voxtral-transcribe.png",
    "path": "/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe",
    "published": "2026-10-01",
    "section": "tools",
    "title": "ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe"
  },
  "tokens": {
    "markdown": 2700,
    "slim": 730
  },
  "version": 1
}
