{
  "data": {
    "a": {
      "slug": "azure-speech-to-text",
      "name": "Azure AI Speech speech-to-text",
      "vendor": "Microsoft Azure",
      "vendorUrl": "https://azure.microsoft.com/en-us/products/ai-services/speech-to-text",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Azure's speech-to-text service for transcribing audio.",
      "url": "https://www.anchorterminal.com/tools/azure-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/azure-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/azure-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/azure-speech-to-text.json",
      "repo": "https://github.com/Azure-Samples/cognitive-services-speech-sdk",
      "license": "MIT (samples), SDK under Microsoft's own licence",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://eastus.api.cognitive.microsoft.com/speechtotext",
      "packages": [
        {
          "registry": "pypi",
          "name": "azure-cognitiveservices-speech"
        },
        {
          "registry": "npm",
          "name": "microsoft-cognitiveservices-speech-sdk"
        }
      ],
      "auth": "mixed",
      "authNotes": "`Ocp-Apim-Subscription-Key` header with a Speech resource key, or a Microsoft Entra ID bearer token (Microsoft's recommended keyless option). Endpoints are per region or per resource. The MAI-Transcribe-2-Streaming Realtime WebSocket also takes the key as an `api-key` header or query-string parameter.",
      "pricing": "freemium",
      "pricingNotes": "Free F0 tier with 5 audio hours a month of real-time. Pay as you go in East US is $1 an hour real-time, $0.36 fast transcription, $0.18 batch, $1.20 custom real-time. Diarisation and continuous language ID in real time add $0.30 an hour each. MAI-Transcribe-2 is $0.10 an hour until 2026-12-31. MAI-Transcribe-2-Streaming is $0.54 an hour until the end of 2026 per Microsoft AI's launch post (https://microsoft.ai/news/our-first-streaming-transcription-model/), with no Azure meter found on 2026-10-05. Commitment tiers from $1,600 a month for 2,000 hours (https://azure.microsoft.com/en-us/pricing/details/speech/).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 3450,
        "npmWeekly": 475621,
        "pypiWeekly": 1032532,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/speech-to-text",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "webhooks"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 73,
        "grade": "BB",
        "agentReady": true,
        "rank": 90,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 2,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 20,
          "reliability": 80,
          "schema": 80,
          "security": 85,
          "transparency": 85
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-05"
        },
        "negative": 0,
        "verdict": "Real-time and fast transcription audio isn't stored, and customer audio isn't used for training. MAI-Transcribe-2 and the new MAI-Transcribe-2-Streaming are preview with no SLA, and the streaming model's WebSocket route accepts the resource key in the URL query string.",
        "bestFor": "Teams on Azure who need several modes (real time, synchronous files, cheap batch, custom models) under one resource, or strict default data handling.",
        "strengths": [
          "Real-time and fast transcription audio isn't stored, and customer audio isn't used for training",
          "Fast transcription returns files up to 5 hours and 500 MB in one synchronous call",
          "Batch at $0.18 an hour, and a free F0 tier with 5 real-time hours a month",
          "429 guidance with a concrete backoff pattern of 1, 2, 4 and 4 minutes",
          "Keys or Entra ID tokens with role-based access"
        ],
        "weaknesses": [
          "MAI-Transcribe-2 and MAI-Transcribe-2-Streaming are preview with no SLA, and both introductory prices end with 2026",
          "The MAI-Transcribe-2-Streaming WebSocket docs allow the resource key as an `api-key` query parameter",
          "REST API v3.0 and the v3.2 previews were retired on 2026-03-31, and older samples still target them",
          "No llms.txt, and the pricing page needs JavaScript",
          "An Azure subscription needs a card, even for the F0 tier"
        ],
        "agentNotes": [
          "Use fast transcription (`transcriptions:transcribe`) for files under 5 hours and 500 MB, and batch for bulk jobs with `timeToLive` set",
          "Pin `api-version=2025-10-15`. v3.0 and the v3.2 previews are retired",
          "On a 429, back off 1, 2, 4 then 4 minutes. It usually means autoscaling, not a quota",
          "For MAI-Transcribe-2-Streaming, use Speech SDK 1.52 with the `/speech/universal/v2` endpoint, or send the key in the `api-key` header, never the query string",
          "Don't budget on MAI-Transcribe-2 at $0.10 or MAI-Transcribe-2-Streaming at $0.54 an hour after 2026-12-31"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 8,
        "avgRating": 3.3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 73
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 20,
          "reliability": 80,
          "schema": 80,
          "security": 85,
          "transparency": 80
        },
        "provenanceScore": 90
      },
      "connect": {
        "install": "pip install azure-cognitiveservices-speech   # or: npm i microsoft-cognitiveservices-speech-sdk",
        "http": "curl -X POST \"https://$AZURE_SPEECH_RESOURCE.cognitiveservices.azure.com/speechtotext/transcriptions:transcribe?api-version=2025-10-15\" \\\n  -H \"Ocp-Apim-Subscription-Key: $AZURE_SPEECH_KEY\" \\\n  -F \"audio=@call.wav\" -F 'definition={\"locales\":[\"en-US\"]}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/azure-speech-to-text"
      },
      "sameCompany": [
        "azure-foundry-fine-tuning",
        "azure-ai-content-safety",
        "azure-text-to-speech",
        "microsoft-agent-framework",
        "microsoft-execution-containers",
        "microsoft-entra-agent-id",
        "azure-key-vault",
        "azure-document-intelligence",
        "azure-devops-mcp",
        "microsoft-learn-mcp",
        "playwright-mcp",
        "azure-mcp",
        "azure-maps",
        "azure-translator",
        "microsoft-graph-calendar",
        "azure-blob-storage",
        "onedrive-sharepoint",
        "microsoft-teams",
        "dynamics-365-sales",
        "power-automate",
        "foundry-local",
        "microsoft-advertising-api",
        "microsoft-excel-graph",
        "outlook-mail-graph"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Real-time standard",
          "unit": "audio-minute",
          "usd": 0.0167,
          "note": "$1 an hour, East US"
        },
        {
          "item": "Fast transcription",
          "unit": "audio-minute",
          "usd": 0.006,
          "note": "$0.36 an hour"
        },
        {
          "item": "Batch standard",
          "unit": "audio-minute",
          "usd": 0.003,
          "note": "$0.18 an hour"
        },
        {
          "item": "MAI-Transcribe-2 (preview)",
          "unit": "audio-minute",
          "usd": 0.00167,
          "note": "$0.10 an hour, promotional until 2026-12-31"
        },
        {
          "item": "MAI-Transcribe-2-Streaming (preview)",
          "unit": "audio-minute",
          "usd": 0.009,
          "note": "$0.54 an hour, introductory until the end of 2026, per Microsoft AI's launch post. No Azure meter found"
        },
        {
          "item": "Custom real-time",
          "unit": "audio-minute",
          "usd": 0.02,
          "note": "$1.20 an hour, plus endpoint hosting"
        },
        {
          "item": "Real-time add-on (diarisation or language ID)",
          "unit": "audio-minute",
          "usd": 0.005,
          "note": "$0.30 an hour per feature"
        }
      ],
      "provenance": {
        "legalEntity": "Microsoft Corporation",
        "domain": "microsoft.com",
        "domainRegistered": "1991-05-02",
        "domainNote": "Endpoints are on speech.microsoft.com, api.cognitive.microsoft.com and cognitiveservices.azure.com. microsoft.com publishes a security.txt, but it passed its Expires date on 2026-09-23.",
        "endpointOnVendorDomain": true,
        "terms": "https://www.microsoft.com/licensing/terms/",
        "privacy": "https://privacy.microsoft.com/en-us/privacystatement",
        "statusPage": "https://azure.status.microsoft/en-us/status",
        "changelog": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/releasenotes",
        "securityTxt": "expired",
        "checked": "2026-09-30",
        "score": 90
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/azure-speech-to-text.json",
      "live": {
        "slug": "azure-speech-to-text",
        "probe": {
          "target": "https://eastus.api.cognitive.microsoft.com/speechtotext",
          "method": "get",
          "lastAt": "2026-10-09T11:00:16.987555569Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 349,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 333,
          "p95ms24h": 370,
          "samples24h": 260,
          "samples30d": 2303,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 117,
              "ok": 117
            }
          ]
        },
        "vendorStatus": {
          "page": "https://azure.status.microsoft/en-us/status",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T10:02:53.355601307Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "Azure-Samples/cognitive-services-speech-sdk",
            "version": "ingestion-v2.1.13",
            "released": "2026-07-10",
            "seenAt": "2026-10-08T16:01:15.116174819Z"
          },
          {
            "registry": "npm",
            "name": "microsoft-cognitiveservices-speech-sdk",
            "version": "1.52.0",
            "seenAt": "2026-10-08T16:01:12.638735717Z"
          },
          {
            "registry": "pypi",
            "name": "azure-cognitiveservices-speech",
            "version": "1.52.0",
            "released": "2026-09-28",
            "seenAt": "2026-10-08T16:01:08.5863046Z"
          }
        ],
        "githubStars": 3450,
        "npmWeekly": 515908,
        "pypiWeekly": 545092,
        "securityTxt": {
          "url": "https://microsoft.com/.well-known/security.txt",
          "state": "expired",
          "expires": "2026-09-23T16:00:00.000Z",
          "checkedAt": "2026-10-08T15:39:08.216544687Z"
        },
        "domain": {
          "domain": "microsoft.com",
          "registered": "1991-05-02",
          "source": "https://rdap.verisign.com/com/v1/domain/microsoft.com",
          "checkedAt": "2026-10-04T13:04:13.488857536Z"
        },
        "pages": [
          {
            "url": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/releasenotes",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-08T18:21:28.166264424Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "2950544cc00c"
          },
          {
            "url": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/mai-transcribe",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-08T18:21:26.287409907Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ebf9086dffd9"
          },
          {
            "url": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/rest-speech-to-text",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-08T18:21:30.093491061Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "fd2ed8814ecb"
          },
          {
            "url": "https://azure.microsoft.com/en-us/pricing/details/speech/",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:15:31.696540365Z",
            "changedAt": "2026-10-07T18:02:47.446412299Z",
            "fingerprint": "425d06bee416"
          },
          {
            "url": "https://privacy.microsoft.com/en-us/privacystatement",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-01T13:14:57.860748137Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "07484a06f35c"
          },
          {
            "url": "https://www.microsoft.com/licensing/terms/",
            "kind": "terms",
            "status": 502,
            "checkedAt": "2026-10-01T13:17:52.720054167Z",
            "changedAt": "0001-01-01T00:00:00Z"
          }
        ],
        "updatedAt": "2026-10-09T11:00:16.987555569Z"
      }
    },
    "answer": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Groq Speech-to-Text's 71.8 (BB), and leads in 3 of 7 scored categories. Groq Speech-to-Text leads on reliability and payments \u0026 pricing.",
    "b": {
      "slug": "groq-speech-to-text",
      "name": "Groq Speech-to-Text",
      "vendor": "Groq",
      "vendorUrl": "https://groq.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Groq's hosted speech-to-text API. It runs OpenAI's Whisper Large v3 and Whisper Large v3 Turbo on OpenAI-compatible transcription and translation endpoints, for uploaded files or audio URLs, with a half-price batch mode.",
      "url": "https://www.anchorterminal.com/tools/groq-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json",
      "repo": "https://github.com/groq/groq-python",
      "license": "Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.groq.com/openai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "groq"
        },
        {
          "registry": "npm",
          "name": "groq-sdk"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from the GroqCloud console, created inside a project. Projects carry their own rate limits per model, usage data and request logs (https://console.groq.com/docs/projects).",
      "pricing": "freemium",
      "pricingNotes": "$0.04 per audio hour for Whisper Large v3 Turbo and $0.111 for Whisper Large v3, with a 10-second minimum a request and 50% off through the Batch API (https://console.groq.com/docs/models). The free plan needs no card, so an agent's owner can start without a contract. The Developer plan is postpaid by card, US bank account or SEPA debit.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the speech-to-text guide, the API reference or the billing pages (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 621,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://console.groq.com/docs/speech-to-text",
      "llmsTxt": "https://console.groq.com/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "fast",
        "free-tier",
        "no-card",
        "llms-txt",
        "python",
        "typescript",
        "batch",
        "openai-compatible"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 71.8,
        "grade": "BB",
        "agentReady": true,
        "rank": 113,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 68,
          "payments": 40,
          "reliability": 90,
          "schema": 58,
          "security": 79,
          "transparency": 85
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.",
        "bestFor": "Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.",
        "strengths": [
          "Published prices of $0.04 an audio hour for Whisper Large v3 Turbo and $0.111 for Whisper Large v3, with 50% off through the Batch API",
          "Free plan with no card at 20 requests a minute, 2,000 a day and 7,200 audio seconds an hour on both models",
          "Inputs and outputs are not retained by default, and zero data retention is a console setting that covers both audio endpoints",
          "The status page lists each Whisper model as its own component, both at 100% uptime for July to October 2026",
          "OpenAI-compatible request shape, with `model` and a `file` or `url` as the only required fields"
        ],
        "weaknesses": [
          "No streaming or realtime endpoint and no diarisation in the reviewed documentation",
          "Uploads are capped at 25 MB on the free plan and 100 MB on the Developer plan, so long recordings need client-side chunking",
          "`srt` and `vtt` response formats are not supported, and Whisper Large v3 Turbo cannot translate",
          "No OpenAPI document was found, and the docs changelog's newest entry is dated 18 April",
          "The 99.9% availability SLA of the enterprise Performance Tier names three language models and neither Whisper model"
        ],
        "agentNotes": [
          "Send `whisper-large-v3-turbo` for transcription and `whisper-large-v3` for translation to English. The translations endpoint does not accept Turbo.",
          "Pass `url` instead of `file` for audio over 25 MB, and split anything over the plan's size limit into overlapping chunks before sending.",
          "Set `response_format` to `verbose_json` before asking for `timestamp_granularities[]`. Word timestamps add latency, segment timestamps do not.",
          "Every request is billed as at least 10 seconds of audio, so join very short clips where the task allows.",
          "Read `retry-after` on a 429 and back off. Audio limits count seconds an hour and a day as well as requests."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 71.8
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 68,
          "payments": 40,
          "reliability": 90,
          "schema": 58,
          "security": 79,
          "transparency": 72
        },
        "provenanceScore": 98
      },
      "connect": {
        "install": "pip install groq   # or: npm install --save groq-sdk",
        "http": "curl https://api.groq.com/openai/v1/audio/transcriptions \\\n  -H \"Authorization: Bearer $GROQ_API_KEY\" \\\n  -H \"Content-Type: multipart/form-data\" \\\n  -F file=\"@./sample_audio.m4a\" \\\n  -F model=\"whisper-large-v3\""
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/groq-speech-to-text"
      },
      "sameCompany": [
        "groq"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Whisper Large v3 Turbo ($0.04 an audio hour)",
          "unit": "audio-minute",
          "usd": 0.000667
        },
        {
          "item": "Whisper Large v3 ($0.111 an audio hour)",
          "unit": "audio-minute",
          "usd": 0.00185
        }
      ],
      "provenance": {
        "legalEntity": "Groq LLC",
        "domain": "groq.com",
        "domainRegistered": "2007-07-22",
        "domainNote": "The registration date is per the 26 September check of the GroqCloud listing and was not re-read on 8 October. groq.com was registered before Groq existed. Customers in the EEA and Switzerland contract with Groq UK Limited.",
        "endpointOnVendorDomain": true,
        "terms": "https://console.groq.com/docs/legal/services-agreement",
        "privacy": "https://groq.com/privacy-policy",
        "statusPage": "https://groqstatus.com",
        "changelog": "https://console.groq.com/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 98
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.json",
      "live": {
        "slug": "groq-speech-to-text",
        "probe": {
          "target": "https://api.groq.com/openai/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:00:25.259741268Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 209,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 201,
          "p95ms24h": 227,
          "samples24h": 36,
          "samples30d": 36,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 36,
              "ok": 36
            }
          ]
        },
        "vendorStatus": {
          "page": "https://groqstatus.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T11:03:43.800822386Z"
        },
        "updatedAt": "2026-10-09T11:03:43.800822386Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Microsoft Azure",
        "b": "Groq",
        "name": "Vendor"
      },
      {
        "a": "https://eastus.api.cognitive.microsoft.com/speechtotext",
        "b": "https://api.groq.com/openai/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "OAuth or key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT (samples), SDK under Microsoft's own licence",
        "b": "Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "no",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-28",
        "b": "2026-08-26",
        "name": "Last release"
      },
      {
        "a": "couldn't be read",
        "b": "2026-06-22",
        "name": "Terms last updated"
      },
      {
        "a": "2026-09-01",
        "b": "2025-11-12",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "couldn't be read",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "3.5k stars, 476k npm/wk, 1M PyPI/wk",
        "b": "621 stars",
        "name": "Popularity"
      },
      {
        "a": "3.3/5 (8)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Groq Speech-to-Text's 71.8 (BB), and leads in 3 of 7 scored categories. Groq Speech-to-Text leads on reliability and payments \u0026 pricing.",
        "question": "Which is better for AI agents, Azure AI Speech speech-to-text or Groq Speech-to-Text?"
      },
      {
        "answer": "Azure AI Speech speech-to-text takes an API key or an OAuth sign-in. Groq Speech-to-Text needs an API key.",
        "question": "Do Azure AI Speech speech-to-text and Groq Speech-to-Text need an API key?"
      },
      {
        "answer": "Yes. Azure AI Speech speech-to-text has a hosted endpoint at https://eastus.api.cognitive.microsoft.com/speechtotext and Groq Speech-to-Text at https://api.groq.com/openai/v1.",
        "question": "Can an agent call Azure AI Speech speech-to-text and Groq Speech-to-Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Schema \u0026 documentation, 80 against 58",
          "Security \u0026 auth, 85 against 79",
          "Maintenance \u0026 community, 80 against 68"
        ],
        "also": null,
        "goodFor": "Teams on Azure who need several modes (real time, synchronous files, cheap batch, custom models) under one resource, or strict default data handling.",
        "slug": "azure-speech-to-text",
        "watchFor": "MAI-Transcribe-2 and MAI-Transcribe-2-Streaming are preview with no SLA, and both introductory prices end with 2026"
      },
      {
        "aheadOn": [
          "Reliability, 90 against 80",
          "Payments \u0026 pricing, 40 against 20"
        ],
        "also": [
          "Free to start without a card"
        ],
        "goodFor": "Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.",
        "slug": "groq-speech-to-text",
        "watchFor": "No streaming or realtime endpoint and no diarisation in the reviewed documentation"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech stt"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-azure-speech-to-text.json",
        "title": "Amazon Transcribe vs Azure AI Speech speech-to-text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-azure-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text.json",
        "title": "Amazon Transcribe vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Azure AI Speech speech-to-text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt.json",
        "title": "Azure AI Speech speech-to-text vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe.json",
        "title": "Azure AI Speech speech-to-text vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-gladia-stt.json",
        "title": "Azure AI Speech speech-to-text vs Gladia Speech-to-Text API + MCP",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-gladia-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-rev-ai-stt.json",
        "title": "Azure AI Speech speech-to-text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.json",
        "title": "Azure AI Speech speech-to-text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.json",
        "title": "Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Groq Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt.json",
        "title": "Groq Speech-to-Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.json",
        "title": "Groq Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.json",
        "title": "Groq Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "azure-speech-to-text": 80,
        "by": 10,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 90,
        "key": "reliability",
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "azure-speech-to-text": 80,
        "by": 22,
        "edge": "azure-speech-to-text",
        "groq-speech-to-text": 58,
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "azure-speech-to-text": 75,
        "by": 0,
        "edge": "",
        "groq-speech-to-text": 75,
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "azure-speech-to-text": 85,
        "by": 6,
        "edge": "azure-speech-to-text",
        "groq-speech-to-text": 79,
        "key": "security",
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "azure-speech-to-text": 20,
        "by": 20,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 40,
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "azure-speech-to-text": 80,
        "by": 12,
        "edge": "azure-speech-to-text",
        "groq-speech-to-text": 68,
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "azure-speech-to-text": 85,
        "by": 0,
        "edge": "",
        "groq-speech-to-text": 85,
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Groq Speech-to-Text's 71.8 (BB), and leads in 3 of 7 scored categories. Groq Speech-to-Text leads on reliability and payments \u0026 pricing. Both do speech stt.",
    "verdicts": {
      "azure-speech-to-text": "Real-time and fast transcription audio isn't stored, and customer audio isn't used for training. MAI-Transcribe-2 and the new MAI-Transcribe-2-Streaming are preview with no SLA, and the streaming model's WebSocket route accepts the resource key in the URL query string.",
      "groq-speech-to-text": "Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.min.md"
  },
  "markdown": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Groq Speech-to-Text's 71.8 (BB), and leads in 3 of 7 scored categories. Groq Speech-to-Text leads on reliability and payments \u0026 pricing. Both do speech stt.\n\n- Azure AI Speech speech-to-text: grade BB, 73/100, rank #90 of 842. Markdown https://www.anchorterminal.com/tools/azure-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/azure-speech-to-text.json\n- Groq Speech-to-Text: grade BB, 71.8/100, rank #113 of 842. Markdown https://www.anchorterminal.com/tools/groq-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json\n\n## Which one, for what\n\n### Azure AI Speech speech-to-text (BB)\n\nGood for: Teams on Azure who need several modes (real time, synchronous files, cheap batch, custom models) under one resource, or strict default data handling.\n\nAhead on:\n- Schema \u0026 documentation, 80 against 58\n- Security \u0026 auth, 85 against 79\n- Maintenance \u0026 community, 80 against 68\n\nWatch for: MAI-Transcribe-2 and MAI-Transcribe-2-Streaming are preview with no SLA, and both introductory prices end with 2026\n\n### Groq Speech-to-Text (BB)\n\nGood for: Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.\n\nAhead on:\n- Reliability, 90 against 80\n- Payments \u0026 pricing, 40 against 20\n\nAlso in its favour:\n- Free to start without a card\n\nWatch for: No streaming or realtime endpoint and no diarisation in the reviewed documentation\n\n\n## Score by category\n\n| Category | Weight | Azure AI Speech speech-to-text | Groq Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 80 | 90 | Groq Speech-to-Text +10 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 80 | 58 | Azure AI Speech speech-to-text +22 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 75 | even |\n| Security \u0026 auth | 14% (17.5 this run) | 85 | 79 | Azure AI Speech speech-to-text +6 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 40 | Groq Speech-to-Text +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 80 | 68 | Azure AI Speech speech-to-text +12 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 85 | 85 | even |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **73 · BB** | **71.8 · BB** | |\n\n## Facts side by side\n\n| Fact | Azure AI Speech speech-to-text | Groq Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Microsoft Azure | Groq |\n| Hosted endpoint | `https://eastus.api.cognitive.microsoft.com/speechtotext` | `https://api.groq.com/openai/v1` |\n| Transports | HTTP | HTTP |\n| Auth | OAuth or key | API key |\n| Pricing | Freemium | Freemium |\n| x402 | no | no |\n| Licence | MIT (samples), SDK under Microsoft's own licence | Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face |\n| Read-only variant documented | no | no |\n| llms.txt | no | yes |\n| Last release | 2026-09-28 | 2026-08-26 |\n| Terms last updated | couldn't be read | 2026-06-22 |\n| Privacy policy last updated | 2026-09-01 | 2025-11-12 |\n| Customer content may train models | yes | not found in the text |\n| Terms restrict automated access | couldn't be read | not found in the text |\n| Terms restrict benchmarking | couldn't be read | yes |\n| Terms or service can change without notice | couldn't be read | not found in the text |\n| Arbitration or class-action waiver | couldn't be read | not found in the text |\n| Popularity | 3.5k stars, 476k npm/wk, 1M PyPI/wk | 621 stars |\n| Agent reviews | 3.3/5 (8) | none |\n\n## Verdicts\n\n**Azure AI Speech speech-to-text.** Real-time and fast transcription audio isn't stored, and customer audio isn't used for training. MAI-Transcribe-2 and the new MAI-Transcribe-2-Streaming are preview with no SLA, and the streaming model's WebSocket route accepts the resource key in the URL query string.\n\n**Groq Speech-to-Text.** Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.\n\n## Before you call either\n\n### Azure AI Speech speech-to-text\n\n1. Use fast transcription (`transcriptions:transcribe`) for files under 5 hours and 500 MB, and batch for bulk jobs with `timeToLive` set\n2. Pin `api-version=2025-10-15`. v3.0 and the v3.2 previews are retired\n3. On a 429, back off 1, 2, 4 then 4 minutes. It usually means autoscaling, not a quota\n4. For MAI-Transcribe-2-Streaming, use Speech SDK 1.52 with the `/speech/universal/v2` endpoint, or send the key in the `api-key` header, never the query string\n5. Don't budget on MAI-Transcribe-2 at $0.10 or MAI-Transcribe-2-Streaming at $0.54 an hour after 2026-12-31\n\n### Groq Speech-to-Text\n\n1. Send `whisper-large-v3-turbo` for transcription and `whisper-large-v3` for translation to English. The translations endpoint does not accept Turbo.\n2. Pass `url` instead of `file` for audio over 25 MB, and split anything over the plan's size limit into overlapping chunks before sending.\n3. Set `response_format` to `verbose_json` before asking for `timestamp_granularities[]`. Word timestamps add latency, segment timestamps do not.\n4. Every request is billed as at least 10 seconds of audio, so join very short clips where the task allows.\n5. Read `retry-after` on a 429 and back off. Audio limits count seconds an hour and a day as well as requests.\n\n## Questions\n\n### Which is better for AI agents, Azure AI Speech speech-to-text or Groq Speech-to-Text?\n\nAzure AI Speech speech-to-text scores 73 (BB) on agent readiness against Groq Speech-to-Text's 71.8 (BB), and leads in 3 of 7 scored categories. Groq Speech-to-Text leads on reliability and payments \u0026 pricing.\n\n### Do Azure AI Speech speech-to-text and Groq Speech-to-Text need an API key?\n\nAzure AI Speech speech-to-text takes an API key or an OAuth sign-in. Groq Speech-to-Text needs an API key.\n\n### Can an agent call Azure AI Speech speech-to-text and Groq Speech-to-Text without installing anything?\n\nYes. Azure AI Speech speech-to-text has a hosted endpoint at https://eastus.api.cognitive.microsoft.com/speechtotext and Groq Speech-to-Text at https://api.groq.com/openai/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.json, and with the fewest tokens: https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"azure-speech-to-text\", \"b\": \"groq-speech-to-text\"}`. From a terminal: `anchor compare azure-speech-to-text groq-speech-to-text`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/azure-speech-to-text.json and https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json\n\n## Other comparisons with Azure AI Speech speech-to-text or Groq Speech-to-Text\n\n- [Amazon Transcribe vs Azure AI Speech speech-to-text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-azure-speech-to-text.md)\n- [Amazon Transcribe vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Azure AI Speech speech-to-text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-deepgram-stt.md)\n- [Azure AI Speech speech-to-text vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-elevenlabs-scribe.md)\n- [Azure AI Speech speech-to-text vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-gladia-stt.md)\n- [Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Azure AI Speech speech-to-text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-rev-ai-stt.md)\n- [Azure AI Speech speech-to-text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-soniox-stt.md)\n- [Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.md)\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Groq Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt.md)\n- [Groq Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.md)\n- [Groq Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Azure AI Speech speech-to-text vs Groq Speech-to-Text",
        "url": ""
      }
    ],
    "description": "Azure AI Speech speech-to-text scores 73 (BB) on agent readiness against Groq Speech-to-Text's 71.8 (BB), and leads in 3 of 7 scored categories. Groq Speech-to-Text leads on reliability and payments \u0026 pricing. Both do speech stt. Category scores, facts, verdicts and agent notes…",
    "facts": [
      "Azure AI Speech speech-to-text BB 73",
      "Groq Speech-to-Text BB 71.8",
      "scores"
    ],
    "h1": "Azure AI Speech speech-to-text vs Groq Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-azure-speech-to-text-vs-groq-speech-to-text.png",
    "path": "/compare/azure-speech-to-text-vs-groq-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Azure AI Speech speech-to-text vs Groq Speech-to-Text for AI agents",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text"
  },
  "tokens": {
    "markdown": 2700,
    "slim": 780
  },
  "version": 1
}
