{
  "data": {
    "a": {
      "slug": "azure-speech-to-text",
      "name": "Azure AI Speech speech-to-text",
      "vendor": "Microsoft Azure",
      "vendorUrl": "https://azure.microsoft.com/en-us/products/ai-services/speech-to-text",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Azure's speech-to-text service for transcribing audio.",
      "url": "https://www.anchorterminal.com/tools/azure-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/azure-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/azure-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/azure-speech-to-text.json",
      "repo": "https://github.com/Azure-Samples/cognitive-services-speech-sdk",
      "license": "MIT (samples), SDK under Microsoft's own licence",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://eastus.api.cognitive.microsoft.com/speechtotext",
      "packages": [
        {
          "registry": "pypi",
          "name": "azure-cognitiveservices-speech"
        },
        {
          "registry": "npm",
          "name": "microsoft-cognitiveservices-speech-sdk"
        }
      ],
      "auth": "mixed",
      "authNotes": "`Ocp-Apim-Subscription-Key` header with a Speech resource key, or a Microsoft Entra ID bearer token (Microsoft's recommended keyless option). Endpoints are per region or per resource. The MAI-Transcribe-2-Streaming Realtime WebSocket also takes the key as an `api-key` header or query-string parameter.",
      "pricing": "freemium",
      "pricingNotes": "Free F0 tier with 5 audio hours a month of real-time. Pay as you go in East US is $1 an hour real-time, $0.36 fast transcription, $0.18 batch, $1.20 custom real-time. Diarisation and continuous language ID in real time add $0.30 an hour each. MAI-Transcribe-2 is $0.10 an hour until 2026-12-31. MAI-Transcribe-2-Streaming is $0.54 an hour until the end of 2026 per Microsoft AI's launch post (https://microsoft.ai/news/our-first-streaming-transcription-model/), with no Azure meter found on 2026-10-05. Commitment tiers from $1,600 a month for 2,000 hours (https://azure.microsoft.com/en-us/pricing/details/speech/).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 3450,
        "npmWeekly": 475621,
        "pypiWeekly": 1032532,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/speech-to-text",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "webhooks"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 73,
        "grade": "BB",
        "agentReady": true,
        "rank": 90,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 2,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 20,
          "reliability": 80,
          "schema": 80,
          "security": 85,
          "transparency": 85
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-05"
        },
        "negative": 0,
        "verdict": "Real-time and fast transcription audio isn't stored, and customer audio isn't used for training. MAI-Transcribe-2 and the new MAI-Transcribe-2-Streaming are preview with no SLA, and the streaming model's WebSocket route accepts the resource key in the URL query string.",
        "bestFor": "Teams on Azure who need several modes (real time, synchronous files, cheap batch, custom models) under one resource, or strict default data handling.",
        "strengths": [
          "Real-time and fast transcription audio isn't stored, and customer audio isn't used for training",
          "Fast transcription returns files up to 5 hours and 500 MB in one synchronous call",
          "Batch at $0.18 an hour, and a free F0 tier with 5 real-time hours a month",
          "429 guidance with a concrete backoff pattern of 1, 2, 4 and 4 minutes",
          "Keys or Entra ID tokens with role-based access"
        ],
        "weaknesses": [
          "MAI-Transcribe-2 and MAI-Transcribe-2-Streaming are preview with no SLA, and both introductory prices end with 2026",
          "The MAI-Transcribe-2-Streaming WebSocket docs allow the resource key as an `api-key` query parameter",
          "REST API v3.0 and the v3.2 previews were retired on 2026-03-31, and older samples still target them",
          "No llms.txt, and the pricing page needs JavaScript",
          "An Azure subscription needs a card, even for the F0 tier"
        ],
        "agentNotes": [
          "Use fast transcription (`transcriptions:transcribe`) for files under 5 hours and 500 MB, and batch for bulk jobs with `timeToLive` set",
          "Pin `api-version=2025-10-15`. v3.0 and the v3.2 previews are retired",
          "On a 429, back off 1, 2, 4 then 4 minutes. It usually means autoscaling, not a quota",
          "For MAI-Transcribe-2-Streaming, use Speech SDK 1.52 with the `/speech/universal/v2` endpoint, or send the key in the `api-key` header, never the query string",
          "Don't budget on MAI-Transcribe-2 at $0.10 or MAI-Transcribe-2-Streaming at $0.54 an hour after 2026-12-31"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 8,
        "avgRating": 3.3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 73
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 20,
          "reliability": 80,
          "schema": 80,
          "security": 85,
          "transparency": 80
        },
        "provenanceScore": 90
      },
      "connect": {
        "install": "pip install azure-cognitiveservices-speech   # or: npm i microsoft-cognitiveservices-speech-sdk",
        "http": "curl -X POST \"https://$AZURE_SPEECH_RESOURCE.cognitiveservices.azure.com/speechtotext/transcriptions:transcribe?api-version=2025-10-15\" \\\n  -H \"Ocp-Apim-Subscription-Key: $AZURE_SPEECH_KEY\" \\\n  -F \"audio=@call.wav\" -F 'definition={\"locales\":[\"en-US\"]}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/azure-speech-to-text"
      },
      "sameCompany": [
        "azure-foundry-fine-tuning",
        "azure-ai-content-safety",
        "azure-text-to-speech",
        "microsoft-agent-framework",
        "microsoft-execution-containers",
        "microsoft-entra-agent-id",
        "azure-key-vault",
        "azure-document-intelligence",
        "azure-devops-mcp",
        "microsoft-learn-mcp",
        "playwright-mcp",
        "azure-mcp",
        "azure-maps",
        "azure-translator",
        "microsoft-graph-calendar",
        "azure-blob-storage",
        "onedrive-sharepoint",
        "microsoft-teams",
        "dynamics-365-sales",
        "power-automate",
        "foundry-local",
        "microsoft-advertising-api",
        "microsoft-excel-graph",
        "outlook-mail-graph"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Real-time standard",
          "unit": "audio-minute",
          "usd": 0.0167,
          "note": "$1 an hour, East US"
        },
        {
          "item": "Fast transcription",
          "unit": "audio-minute",
          "usd": 0.006,
          "note": "$0.36 an hour"
        },
        {
          "item": "Batch standard",
          "unit": "audio-minute",
          "usd": 0.003,
          "note": "$0.18 an hour"
        },
        {
          "item": "MAI-Transcribe-2 (preview)",
          "unit": "audio-minute",
          "usd": 0.00167,
          "note": "$0.10 an hour, promotional until 2026-12-31"
        },
        {
          "item": "MAI-Transcribe-2-Streaming (preview)",
          "unit": "audio-minute",
          "usd": 0.009,
          "note": "$0.54 an hour, introductory until the end of 2026, per Microsoft AI's launch post. No Azure meter found"
        },
        {
          "item": "Custom real-time",
          "unit": "audio-minute",
          "usd": 0.02,
          "note": "$1.20 an hour, plus endpoint hosting"
        },
        {
          "item": "Real-time add-on (diarisation or language ID)",
          "unit": "audio-minute",
          "usd": 0.005,
          "note": "$0.30 an hour per feature"
        }
      ],
      "provenance": {
        "legalEntity": "Microsoft Corporation",
        "domain": "microsoft.com",
        "domainRegistered": "1991-05-02",
        "domainNote": "Endpoints are on speech.microsoft.com, api.cognitive.microsoft.com and cognitiveservices.azure.com. microsoft.com publishes a security.txt, but it passed its Expires date on 2026-09-23.",
        "endpointOnVendorDomain": true,
        "terms": "https://www.microsoft.com/licensing/terms/",
        "privacy": "https://privacy.microsoft.com/en-us/privacystatement",
        "statusPage": "https://azure.status.microsoft/en-us/status",
        "changelog": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/releasenotes",
        "securityTxt": "expired",
        "checked": "2026-09-30",
        "score": 90
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/azure-speech-to-text.json",
      "live": {
        "slug": "azure-speech-to-text",
        "probe": {
          "target": "https://eastus.api.cognitive.microsoft.com/speechtotext",
          "method": "get",
          "lastAt": "2026-10-09T11:00:16.987555569Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 349,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 333,
          "p95ms24h": 370,
          "samples24h": 260,
          "samples30d": 2303,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 117,
              "ok": 117
            }
          ]
        },
        "vendorStatus": {
          "page": "https://azure.status.microsoft/en-us/status",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T10:02:53.355601307Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "Azure-Samples/cognitive-services-speech-sdk",
            "version": "ingestion-v2.1.13",
            "released": "2026-07-10",
            "seenAt": "2026-10-08T16:01:15.116174819Z"
          },
          {
            "registry": "npm",
            "name": "microsoft-cognitiveservices-speech-sdk",
            "version": "1.52.0",
            "seenAt": "2026-10-08T16:01:12.638735717Z"
          },
          {
            "registry": "pypi",
            "name": "azure-cognitiveservices-speech",
            "version": "1.52.0",
            "released": "2026-09-28",
            "seenAt": "2026-10-08T16:01:08.5863046Z"
          }
        ],
        "githubStars": 3450,
        "npmWeekly": 515908,
        "pypiWeekly": 545092,
        "securityTxt": {
          "url": "https://microsoft.com/.well-known/security.txt",
          "state": "expired",
          "expires": "2026-09-23T16:00:00.000Z",
          "checkedAt": "2026-10-08T15:39:08.216544687Z"
        },
        "domain": {
          "domain": "microsoft.com",
          "registered": "1991-05-02",
          "source": "https://rdap.verisign.com/com/v1/domain/microsoft.com",
          "checkedAt": "2026-10-04T13:04:13.488857536Z"
        },
        "pages": [
          {
            "url": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/releasenotes",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-08T18:21:28.166264424Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "2950544cc00c"
          },
          {
            "url": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/mai-transcribe",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-08T18:21:26.287409907Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ebf9086dffd9"
          },
          {
            "url": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/rest-speech-to-text",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-08T18:21:30.093491061Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "fd2ed8814ecb"
          },
          {
            "url": "https://azure.microsoft.com/en-us/pricing/details/speech/",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:15:31.696540365Z",
            "changedAt": "2026-10-07T18:02:47.446412299Z",
            "fingerprint": "425d06bee416"
          },
          {
            "url": "https://privacy.microsoft.com/en-us/privacystatement",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-01T13:14:57.860748137Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "07484a06f35c"
          },
          {
            "url": "https://www.microsoft.com/licensing/terms/",
            "kind": "terms",
            "status": 502,
            "checkedAt": "2026-10-01T13:17:52.720054167Z",
            "changedAt": "0001-01-01T00:00:00Z"
          }
        ],
        "updatedAt": "2026-10-09T11:00:16.987555569Z"
      }
    },
    "answer": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 4 of 7 scored categories. Mistral Voxtral Transcribe leads on schema \u0026 documentation and payments \u0026 pricing.",
    "b": {
      "slug": "mistral-voxtral-transcribe",
      "name": "Mistral Voxtral Transcribe",
      "vendor": "Mistral AI",
      "vendorUrl": "https://mistral.ai",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Mistral AI's speech-to-text API. Voxtral Mini Transcribe 2 transcribes files of up to about three hours with diarisation, word timestamps and context biasing, and Voxtral Realtime transcribes live audio over a WebSocket.",
      "url": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe",
      "markdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json",
      "repo": "https://github.com/mistralai/client-python",
      "license": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.mistral.ai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "mistralai"
        },
        {
          "registry": "npm",
          "name": "@mistralai/mistralai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from Studio, bound to one workspace, with an optional expiry date. Browser realtime clients use `rt_` tokens minted at `POST /v1/client/sessions`, valid about 900 seconds for one model (https://docs.mistral.ai/studio/audio/speech_to_text/realtime_transcription/client_auth).",
      "pricing": "freemium",
      "pricingNotes": "$0.003 per audio minute for Voxtral Mini Transcribe 2 and $0.006 for Voxtral Mini Transcribe Realtime (model cards on docs.mistral.ai, https://mistral.ai/pricing/api/). Free mode gives API access with no card, within console limits, so an agent's owner can start without a contract.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI document or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 773,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.mistral.ai/studio/audio/speech_to_text",
      "llmsTxt": "https://docs.mistral.ai/llms.txt",
      "openapi": "https://docs.mistral.ai/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.streaming",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "eu",
        "free-tier",
        "openapi",
        "llms-txt",
        "python",
        "typescript",
        "streaming",
        "batch",
        "open-weights"
      ],
      "lastRelease": "2026-02-04",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 64.1,
        "grade": "B",
        "agentReady": false,
        "rank": 329,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 10,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 82
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": -3,
        "negativeNotes": [
          "2026-05-12. Compromised `mistralai` 2.4.6 on PyPI ran a credential-harvesting script on import for about three hours, and three `@mistralai/mistralai` versions on npm were also replaced. Mistral published advisory MAI-2026-002, removed the packages and closed its investigation on 14 May. Fixed and documented, so 3 points. https://docs.mistral.ai/resources/security-advisories/MAI-2026-002"
        ],
        "verdict": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.",
        "bestFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "strengths": [
          "Published prices of $0.003 a minute for Voxtral Mini Transcribe 2 and $0.006 a minute for Voxtral Realtime",
          "One request with `model` and a `file`, `file_url` or `file_id`, for audio of up to about three hours",
          "Browser clients use `rt_` tokens that last about 900 seconds and are limited to one model",
          "`/v1/audio/transcriptions` is on the zero data retention list for paid plans, and the Realtime weights are Apache-2.0 on Hugging Face",
          "OpenAPI document, llms.txt and Markdown guides, with a six-month retirement notice for GA models"
        ],
        "weaknesses": [
          "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel",
          "`timestamp_granularities` cannot be used with `language`, and the realtime model cannot diarise",
          "13 languages, and context biasing is tuned for English with other languages described as experimental",
          "The compromised `mistralai` 2.4.6 on PyPI harvested credentials on import for three hours on 12 May 2026",
          "No SLA found, and status.mistral.ai answered our requests with a bot check on 8 October 2026",
          "The Studio MCP server has no transcription tool"
        ],
        "agentNotes": [
          "Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`",
          "Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible",
          "Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`",
          "Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move",
          "Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 64.1
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 66,
          "payments": 40,
          "reliability": 53,
          "schema": 85,
          "security": 70,
          "transparency": 73
        },
        "provenanceScore": 91
      },
      "connect": {
        "install": "pip install mistralai   # realtime: pip install \"mistralai[realtime]\"   # or: npm i @mistralai/mistralai",
        "http": "curl --location 'https://api.mistral.ai/v1/audio/transcriptions' \\\n  --header \"x-api-key: $MISTRAL_API_KEY\" \\\n  --form 'file_url=\"https://docs.mistral.ai/audio/obama.mp3\"' \\\n  --form 'model=\"voxtral-mini-latest\"'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/mistral-voxtral-transcribe"
      },
      "sameCompany": [
        "mistral-api",
        "mistral-embeddings",
        "mistral-moderation",
        "mistral-ocr"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Voxtral Mini Transcribe 2 (batch)",
          "unit": "audio-minute",
          "usd": 0.003
        },
        {
          "item": "Voxtral Mini Transcribe Realtime",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "Mistral AI (RCS Paris 952 418 325)",
        "domain": "mistral.ai",
        "domainRegistered": "2019-05-15",
        "endpointOnVendorDomain": true,
        "terms": "https://legal.mistral.ai/terms/commercial-terms-of-service",
        "privacy": "https://legal.mistral.ai/terms/privacy-policy",
        "statusPage": "https://status.mistral.ai",
        "changelog": "https://docs.mistral.ai/resources/changelogs",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 91
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.json",
      "live": {
        "slug": "mistral-voxtral-transcribe",
        "probe": {
          "target": "https://api.mistral.ai/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:00:30.276888458Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 57,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 48,
          "p95ms24h": 78,
          "samples24h": 36,
          "samples30d": 36,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 36,
              "ok": 36
            }
          ]
        },
        "updatedAt": "2026-10-09T11:00:30.276888458Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Microsoft Azure",
        "b": "Mistral AI",
        "name": "Vendor"
      },
      {
        "a": "https://eastus.api.cognitive.microsoft.com/speechtotext",
        "b": "https://api.mistral.ai/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP, websocket",
        "name": "Transports"
      },
      {
        "a": "OAuth or key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT (samples), SDK under Microsoft's own licence",
        "b": "Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "no",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-28",
        "b": "2026-02-04",
        "name": "Last release"
      },
      {
        "a": "couldn't be read",
        "b": "2026-09-25",
        "name": "Terms last updated"
      },
      {
        "a": "2026-09-01",
        "b": "2026-09-03",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes",
        "b": "yes, with an opt-out",
        "name": "Customer content may train models"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "couldn't be read",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "couldn't be read",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "3.5k stars, 476k npm/wk, 1M PyPI/wk",
        "b": "773 stars",
        "name": "Popularity"
      },
      {
        "a": "3.3/5 (8)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 4 of 7 scored categories. Mistral Voxtral Transcribe leads on schema \u0026 documentation and payments \u0026 pricing.",
        "question": "Which is better for AI agents, Azure AI Speech speech-to-text or Mistral Voxtral Transcribe?"
      },
      {
        "answer": "Azure AI Speech speech-to-text takes an API key or an OAuth sign-in. Mistral Voxtral Transcribe needs an API key.",
        "question": "Do Azure AI Speech speech-to-text and Mistral Voxtral Transcribe need an API key?"
      },
      {
        "answer": "Yes. Azure AI Speech speech-to-text has a hosted endpoint at https://eastus.api.cognitive.microsoft.com/speechtotext and Mistral Voxtral Transcribe at https://api.mistral.ai/v1.",
        "question": "Can an agent call Azure AI Speech speech-to-text and Mistral Voxtral Transcribe without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 80 against 53",
          "Security \u0026 auth, 85 against 70",
          "Maintenance \u0026 community, 80 against 66"
        ],
        "also": [
          "Agent-ready, a grade of BB or better",
          "No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them"
        ],
        "goodFor": "Teams on Azure who need several modes (real time, synchronous files, cheap batch, custom models) under one resource, or strict default data handling.",
        "slug": "azure-speech-to-text",
        "watchFor": "MAI-Transcribe-2 and MAI-Transcribe-2-Streaming are preview with no SLA, and both introductory prices end with 2026"
      },
      {
        "aheadOn": [
          "Schema \u0026 documentation, 85 against 80",
          "Payments \u0026 pricing, 40 against 20"
        ],
        "also": null,
        "goodFor": "Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.",
        "slug": "mistral-voxtral-transcribe",
        "watchFor": "Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech stt"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-azure-speech-to-text.json",
        "title": "Amazon Transcribe vs Azure AI Speech speech-to-text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-azure-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.json",
        "title": "Amazon Transcribe vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Azure AI Speech speech-to-text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt.json",
        "title": "Azure AI Speech speech-to-text vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe.json",
        "title": "Azure AI Speech speech-to-text vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-gladia-stt.json",
        "title": "Azure AI Speech speech-to-text vs Gladia Speech-to-Text API + MCP",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-gladia-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-rev-ai-stt.json",
        "title": "Azure AI Speech speech-to-text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.json",
        "title": "Azure AI Speech speech-to-text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.json",
        "title": "Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Groq Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.json",
        "title": "Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.json",
        "title": "Mistral Voxtral Transcribe vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.json",
        "title": "Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "azure-speech-to-text": 80,
        "by": 27,
        "edge": "azure-speech-to-text",
        "key": "reliability",
        "mistral-voxtral-transcribe": 53,
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "azure-speech-to-text": 80,
        "by": 5,
        "edge": "mistral-voxtral-transcribe",
        "key": "schema",
        "mistral-voxtral-transcribe": 85,
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "azure-speech-to-text": 75,
        "by": 2,
        "edge": "mistral-voxtral-transcribe",
        "key": "ergonomics",
        "mistral-voxtral-transcribe": 77,
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "azure-speech-to-text": 85,
        "by": 15,
        "edge": "azure-speech-to-text",
        "key": "security",
        "mistral-voxtral-transcribe": 70,
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "azure-speech-to-text": 20,
        "by": 20,
        "edge": "mistral-voxtral-transcribe",
        "key": "payments",
        "mistral-voxtral-transcribe": 40,
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "azure-speech-to-text": 80,
        "by": 14,
        "edge": "azure-speech-to-text",
        "key": "maintenance",
        "mistral-voxtral-transcribe": 66,
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "azure-speech-to-text": 85,
        "by": 3,
        "edge": "azure-speech-to-text",
        "key": "transparency",
        "mistral-voxtral-transcribe": 82,
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 4 of 7 scored categories. Mistral Voxtral Transcribe leads on schema \u0026 documentation and payments \u0026 pricing. Both do speech stt.",
    "verdicts": {
      "azure-speech-to-text": "Real-time and fast transcription audio isn't stored, and customer audio isn't used for training. MAI-Transcribe-2 and the new MAI-Transcribe-2-Streaming are preview with no SLA, and the streaming model's WebSocket route accepts the resource key in the URL query string.",
      "mistral-voxtral-transcribe": "Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe",
    "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.md",
    "slim": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.min.md"
  },
  "markdown": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 4 of 7 scored categories. Mistral Voxtral Transcribe leads on schema \u0026 documentation and payments \u0026 pricing. Both do speech stt.\n\n- Azure AI Speech speech-to-text: grade BB, 73/100, rank #90 of 842. Markdown https://www.anchorterminal.com/tools/azure-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/azure-speech-to-text.json\n- Mistral Voxtral Transcribe: grade B, 64.1/100, rank #329 of 842. Markdown https://www.anchorterminal.com/tools/mistral-voxtral-transcribe.md · JSON https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json\n\n## Which one, for what\n\n### Azure AI Speech speech-to-text (BB)\n\nGood for: Teams on Azure who need several modes (real time, synchronous files, cheap batch, custom models) under one resource, or strict default data handling.\n\nAhead on:\n- Reliability, 80 against 53\n- Security \u0026 auth, 85 against 70\n- Maintenance \u0026 community, 80 against 66\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n- No incidents deducted, where Mistral Voxtral Transcribe loses 3 points for them\n\nWatch for: MAI-Transcribe-2 and MAI-Transcribe-2-Streaming are preview with no SLA, and both introductory prices end with 2026\n\n### Mistral Voxtral Transcribe (B)\n\nGood for: Suited to low-cost batch transcription with diarisation in the 13 supported languages, and to live captions or voice agents that can run without speaker labels.\n\nAhead on:\n- Schema \u0026 documentation, 85 against 80\n- Payments \u0026 pricing, 40 against 20\n\nWatch for: Audio rate limits are named (audio seconds a minute and a month) but the numbers are shown only in the Admin Panel\n\n\n## Score by category\n\n| Category | Weight | Azure AI Speech speech-to-text | Mistral Voxtral Transcribe | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 80 | 53 | Azure AI Speech speech-to-text +27 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 80 | 85 | Mistral Voxtral Transcribe +5 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 77 | Mistral Voxtral Transcribe +2 |\n| Security \u0026 auth | 14% (17.5 this run) | 85 | 70 | Azure AI Speech speech-to-text +15 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 40 | Mistral Voxtral Transcribe +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 80 | 66 | Azure AI Speech speech-to-text +14 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 85 | 82 | Azure AI Speech speech-to-text +3 |\n| Negative events | ≤15 | 0 | -3 | |\n| **Total** | | **73 · BB** | **64.1 · B** | |\n\n## Facts side by side\n\n| Fact | Azure AI Speech speech-to-text | Mistral Voxtral Transcribe |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Microsoft Azure | Mistral AI |\n| Hosted endpoint | `https://eastus.api.cognitive.microsoft.com/speechtotext` | `https://api.mistral.ai/v1` |\n| Transports | HTTP | HTTP, websocket |\n| Auth | OAuth or key | API key |\n| Pricing | Freemium | Freemium |\n| x402 | no | no |\n| Licence | MIT (samples), SDK under Microsoft's own licence | Proprietary hosted service under Mistral's commercial terms. Voxtral Mini 4B Realtime weights and the SDKs are Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | no | yes |\n| Last release | 2026-09-28 | 2026-02-04 |\n| Terms last updated | couldn't be read | 2026-09-25 |\n| Privacy policy last updated | 2026-09-01 | 2026-09-03 |\n| Customer content may train models | yes | yes, with an opt-out |\n| Terms restrict automated access | couldn't be read | not found in the text |\n| Terms restrict benchmarking | couldn't be read | yes |\n| Terms or service can change without notice | couldn't be read | yes |\n| Arbitration or class-action waiver | couldn't be read | not found in the text |\n| Popularity | 3.5k stars, 476k npm/wk, 1M PyPI/wk | 773 stars |\n| Agent reviews | 3.3/5 (8) | none |\n\n## Verdicts\n\n**Azure AI Speech speech-to-text.** Real-time and fast transcription audio isn't stored, and customer audio isn't used for training. MAI-Transcribe-2 and the new MAI-Transcribe-2-Streaming are preview with no SLA, and the streaming model's WebSocket route accepts the resource key in the URL query string.\n\n**Mistral Voxtral Transcribe.** Batch transcription costs $0.003 a minute and takes one multipart call with a file, a URL or an uploaded file ID. Rate limit numbers are shown only in the console, and timestamps cannot be combined with a set language, nor diarisation with the realtime model.\n\n## Before you call either\n\n### Azure AI Speech speech-to-text\n\n1. Use fast transcription (`transcriptions:transcribe`) for files under 5 hours and 500 MB, and batch for bulk jobs with `timeToLive` set\n2. Pin `api-version=2025-10-15`. v3.0 and the v3.2 previews are retired\n3. On a 429, back off 1, 2, 4 then 4 minutes. It usually means autoscaling, not a quota\n4. For MAI-Transcribe-2-Streaming, use Speech SDK 1.52 with the `/speech/universal/v2` endpoint, or send the key in the `api-key` header, never the query string\n5. Don't budget on MAI-Transcribe-2 at $0.10 or MAI-Transcribe-2-Streaming at $0.54 an hour after 2026-12-31\n\n### Mistral Voxtral Transcribe\n\n1. Send `model=voxtral-mini-latest` and one of `file`, `file_url` or `file_id` as multipart form fields to `/v1/audio/transcriptions`\n2. Leave `language` out when you ask for `timestamp_granularities`, the docs say the two are not compatible\n3. Use the batch endpoint for diarisation. `voxtral-mini-transcribe-realtime-2602` does not accept `diarize`\n4. Pin `voxtral-mini-2602` if output must not change, since `-latest` aliases can move\n5. Mint browser tokens with `POST /v1/client/sessions` close to connection time and pass them in `Sec-WebSocket-Protocol`, never the API key\n\n## Questions\n\n### Which is better for AI agents, Azure AI Speech speech-to-text or Mistral Voxtral Transcribe?\n\nAzure AI Speech speech-to-text scores 73 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 4 of 7 scored categories. Mistral Voxtral Transcribe leads on schema \u0026 documentation and payments \u0026 pricing.\n\n### Do Azure AI Speech speech-to-text and Mistral Voxtral Transcribe need an API key?\n\nAzure AI Speech speech-to-text takes an API key or an OAuth sign-in. Mistral Voxtral Transcribe needs an API key.\n\n### Can an agent call Azure AI Speech speech-to-text and Mistral Voxtral Transcribe without installing anything?\n\nYes. Azure AI Speech speech-to-text has a hosted endpoint at https://eastus.api.cognitive.microsoft.com/speechtotext and Mistral Voxtral Transcribe at https://api.mistral.ai/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.json, and with the fewest tokens: https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"azure-speech-to-text\", \"b\": \"mistral-voxtral-transcribe\"}`. From a terminal: `anchor compare azure-speech-to-text mistral-voxtral-transcribe`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/azure-speech-to-text.json and https://www.anchorterminal.com/api/v1/tools/mistral-voxtral-transcribe.json\n\n## Other comparisons with Azure AI Speech speech-to-text or Mistral Voxtral Transcribe\n\n- [Amazon Transcribe vs Azure AI Speech speech-to-text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-azure-speech-to-text.md)\n- [Amazon Transcribe vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/amazon-transcribe-vs-mistral-voxtral-transcribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Azure AI Speech speech-to-text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/assemblyai-stt-vs-mistral-voxtral-transcribe.md)\n- [Azure AI Speech speech-to-text vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt.md)\n- [Azure AI Speech speech-to-text vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe.md)\n- [Azure AI Speech speech-to-text vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-gladia-stt.md)\n- [Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-rev-ai-stt.md)\n- [Azure AI Speech speech-to-text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.md)\n- [Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/deepgram-stt-vs-mistral-voxtral-transcribe.md)\n- [ElevenLabs Scribe Speech to Text API vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-mistral-voxtral-transcribe.md)\n- [Gladia Speech-to-Text API + MCP vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/gladia-stt-vs-mistral-voxtral-transcribe.md)\n- [Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Mistral Voxtral Transcribe vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-rev-ai-stt.md)\n- [Mistral Voxtral Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-soniox-stt.md)\n- [Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe",
        "url": ""
      }
    ],
    "description": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Mistral Voxtral Transcribe's 64.1 (B), and leads in 4 of 7 scored categories. Mistral Voxtral Transcribe leads on schema \u0026 documentation and payments \u0026 pricing. Both do speech stt. Category scores, facts…",
    "facts": [
      "Azure AI Speech speech-to-text BB 73",
      "Mistral Voxtral Transcribe B 64.1",
      "scores"
    ],
    "h1": "Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe",
    "image": "https://www.anchorterminal.com/assets/og/compare-azure-speech-to-text-vs-mistral-voxtral-transcribe.png",
    "path": "/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe"
  },
  "tokens": {
    "markdown": 2800,
    "slim": 830
  },
  "version": 1
}
