{
  "data": {
    "a": {
      "slug": "azure-text-to-speech",
      "name": "Azure AI Speech text-to-speech",
      "vendor": "Microsoft Azure",
      "vendorUrl": "https://azure.microsoft.com/en-us/products/ai-services/text-to-speech",
      "kind": "model",
      "category": "text-to-speech",
      "summary": "Azure's text-to-speech service for generating spoken audio.",
      "url": "https://www.anchorterminal.com/tools/azure-text-to-speech",
      "markdownUrl": "https://www.anchorterminal.com/tools/azure-text-to-speech.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/azure-text-to-speech.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/azure-text-to-speech.json",
      "repo": "https://github.com/Azure-Samples/cognitive-services-speech-sdk",
      "license": "MIT (samples), SDK under Microsoft's own licence",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://eastus.tts.speech.microsoft.com/cognitiveservices",
      "packages": [
        {
          "registry": "pypi",
          "name": "azure-cognitiveservices-speech"
        },
        {
          "registry": "npm",
          "name": "microsoft-cognitiveservices-speech-sdk"
        }
      ],
      "auth": "mixed",
      "authNotes": "`Ocp-Apim-Subscription-Key` header with a Speech resource key, or a Microsoft Entra ID bearer token. Endpoints are per region.",
      "pricing": "freemium",
      "pricingNotes": "Free F0 tier with 500,000 characters a month. Pay as you go in East US is $15 per 1M characters for Neural and Neural HD Flash voices and $22 for Neural HD, real time or batch. Commitment tiers from $960 a month for 80M characters (https://azure.microsoft.com/en-us/pricing/details/speech/).",
      "priceSummary": "$960 / mo",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 3450,
        "npmWeekly": 475621,
        "pypiWeekly": 1032532,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/text-to-speech",
      "capabilities": [
        "speech.tts",
        "speech.streaming",
        "speech.voices",
        "speech.ssml",
        "speech.languages"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 73.7,
        "grade": "BB",
        "agentReady": true,
        "rank": 56,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 2,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 20,
          "reliability": 90,
          "schema": 65,
          "security": 90,
          "transparency": 88
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Real-time synthesis keeps neither the input text nor the output audio. An Azure subscription needs a card, even for the free F0 tier.",
        "strengths": [
          "Real-time synthesis keeps neither the input text nor the output audio",
          "Full SSML with speaking styles, prosody, phonemes, lexicons and up to 50 voice or audio tags a request",
          "Microsoft Entra ID with role-based access, or two rotatable keys",
          "Covered by Microsoft's online services SLA",
          "S0 starts at 30 requests a second and can be raised to 1,000"
        ],
        "weaknesses": [
          "An Azure subscription needs a card, even for the free F0 tier",
          "No llms.txt and no OpenAPI file for text-to-speech found",
          "429s often reflect busy capacity for a voice in a region, which a quota increase doesn't fix",
          "The voice list comes back as one response per region with no paging documented",
          "MAI-Voice-2-Flash, the low-latency model, is still in preview"
        ],
        "agentNotes": [
          "Send SSML with `\u003cspeak\u003e` and `\u003cvoice\u003e`, and set `X-Microsoft-OutputFormat` and `User-Agent`.",
          "On 429 retry with backoff, and try the voice's home region or another region rather than asking for more quota.",
          "Keep each real-time request under 10 minutes of audio, or use batch synthesis.",
          "Use Entra ID tokens instead of resource keys where the agent runs inside Azure.",
          "Cache the voice list per region, since it returns hundreds of entries at once."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 73.7
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 20,
          "reliability": 90,
          "schema": 65,
          "security": 90,
          "transparency": 80
        },
        "provenanceScore": 95
      },
      "connect": {
        "install": "pip install azure-cognitiveservices-speech   # or: npm i microsoft-cognitiveservices-speech-sdk",
        "http": "curl -X POST \"https://eastus.tts.speech.microsoft.com/cognitiveservices/v1\" \\\n  -H \"Ocp-Apim-Subscription-Key: $AZURE_SPEECH_KEY\" -H \"Content-Type: application/ssml+xml\" \\\n  -H \"X-Microsoft-OutputFormat: audio-24khz-48kbitrate-mono-mp3\" -o speech.mp3 \\\n  -d '\u003cspeak version=\"1.0\" xml:lang=\"en-US\"\u003e\u003cvoice name=\"en-US-AvaMultilingualNeural\"\u003eYour table is booked for seven.\u003c/voice\u003e\u003c/speak\u003e'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.tts",
        "tool": "https://letme.dev/azure-text-to-speech"
      },
      "sameCompany": [
        "azure-foundry-fine-tuning",
        "azure-ai-content-safety",
        "azure-speech-to-text",
        "microsoft-learn-mcp",
        "playwright-mcp",
        "azure-mcp",
        "azure-translator",
        "microsoft-graph-calendar"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Neural and Neural HD Flash voices",
          "unit": "1m-chars",
          "usd": 15,
          "note": "real time or batch, East US"
        },
        {
          "item": "Neural HD voices",
          "unit": "1m-chars",
          "usd": 22
        },
        {
          "item": "Commitment tier 80M characters",
          "unit": "month",
          "usd": 960,
          "note": "$12 per 1M overage"
        }
      ],
      "provenance": {
        "legalEntity": "Microsoft Corporation",
        "domain": "microsoft.com",
        "domainRegistered": "1991-05-02",
        "domainNote": "Endpoints are on speech.microsoft.com, api.cognitive.microsoft.com and cognitiveservices.azure.com. microsoft.com publishes a security.txt, but it passed its Expires date on 2026-09-23.",
        "endpointOnVendorDomain": true,
        "terms": "https://www.microsoft.com/licensing/terms/",
        "privacy": "https://privacy.microsoft.com/en-us/privacystatement",
        "statusPage": "https://azure.status.microsoft/en-us/status",
        "changelog": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/releasenotes",
        "securityTxt": "expired",
        "checked": "2026-09-30",
        "score": 95
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/azure-text-to-speech.json",
      "live": {
        "slug": "azure-text-to-speech",
        "probe": {
          "target": "https://eastus.tts.speech.microsoft.com/cognitiveservices",
          "method": "get",
          "lastAt": "2026-10-05T00:25:39.141397349Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 257,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 257,
          "p95ms24h": 304,
          "samples24h": 272,
          "samples30d": 1107,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 5,
              "ok": 5
            }
          ]
        },
        "versions": [
          {
            "registry": "github",
            "name": "Azure-Samples/cognitive-services-speech-sdk",
            "version": "ingestion-v2.1.13",
            "released": "2026-07-10",
            "seenAt": "2026-10-04T16:21:39.409053457Z"
          },
          {
            "registry": "npm",
            "name": "microsoft-cognitiveservices-speech-sdk",
            "version": "1.52.0",
            "seenAt": "2026-10-04T16:21:39.358419331Z"
          },
          {
            "registry": "pypi",
            "name": "azure-cognitiveservices-speech",
            "version": "1.52.0",
            "released": "2026-09-28",
            "seenAt": "2026-10-04T16:21:39.229173438Z"
          }
        ],
        "githubStars": 3450,
        "npmWeekly": 508156,
        "pypiWeekly": 704368,
        "securityTxt": {
          "url": "https://microsoft.com/.well-known/security.txt",
          "state": "expired",
          "expires": "2026-09-23T16:00:00.000Z",
          "checkedAt": "2026-10-04T15:16:01.36832038Z"
        },
        "domain": {
          "domain": "microsoft.com",
          "registered": "1991-05-02",
          "source": "https://rdap.verisign.com/com/v1/domain/microsoft.com",
          "checkedAt": "2026-10-04T13:04:13.488857536Z"
        },
        "updatedAt": "2026-10-05T00:25:39.141397349Z"
      }
    },
    "b": {
      "slug": "resemble-ai-tts",
      "name": "Resemble AI Text-to-Speech API",
      "vendor": "Resemble AI",
      "vendorUrl": "https://www.resemble.ai",
      "kind": "model",
      "category": "text-to-speech",
      "summary": "Resemble's TTS API on its current Resemble Ultra model, which the changelog says is powered by xAI.",
      "url": "https://www.anchorterminal.com/tools/resemble-ai-tts",
      "markdownUrl": "https://www.anchorterminal.com/tools/resemble-ai-tts.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/resemble-ai-tts.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/resemble-ai-tts.json",
      "repo": "https://github.com/resemble-ai/resemble-node",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://app.resemble.ai/api/v2",
      "packages": [
        {
          "registry": "npm",
          "name": "@resemble/node"
        },
        {
          "registry": "pypi",
          "name": "resemble"
        }
      ],
      "auth": "api-key",
      "authNotes": "`Authorization: Bearer` API key on every server. Synthesis runs on `f.cluster.resemble.ai`, voices and projects on `app.resemble.ai/api/v2`, streaming on `wss://websocket.cluster.resemble.ai`. Some synchronous examples in the docs leave out the `Bearer` prefix.",
      "pricing": "usage",
      "pricingNotes": "Billed per second of generated audio. Flex has no subscription and no card to sign up, at $0.00067 a second (about $0.04 a minute). Team ($350 a month, $280 billed annually) and Business ($1,000, $800 annually) cut that to $0.0005 a second (about $0.03 a minute). WebSocket streaming needs Business. Extra seats cost $20 on Flex and $200 on Team and Business. Rates come from the public plans API, since the pricing page lists only detection products (https://app.resemble.ai/billing/api/v1/plans).",
      "priceSummary": "$350 / mo",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 15,
        "npmWeekly": 8053,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.resemble.ai/voice-generation/text-to-speech",
      "llmsTxt": "https://docs.resemble.ai/llms.txt",
      "openapi": "https://docs.resemble.ai/openapi.json",
      "capabilities": [
        "speech.tts",
        "speech.streaming",
        "speech.voices",
        "speech.ssml"
      ],
      "tags": [
        "hosted",
        "no-card",
        "closed-source",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "streaming",
        "watermark",
        "enterprise"
      ],
      "lastRelease": "2026-06-30",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 50.6,
        "grade": "D",
        "agentReady": false,
        "rank": 357,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 9,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 46,
          "maintenance": 31,
          "payments": 15,
          "reliability": 75,
          "schema": 75,
          "security": 48,
          "transparency": 68
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": -3,
        "negativeNotes": [
          "2026-06-29. Resemble began deprecating every earlier TTS model, including the Chatterbox family, and the model-versions page now says voices on them can no longer generate audio until upgraded to Resemble Ultra. We found no effective date or notice period. Small deduction because the change is documented and upgrade is self-serve. https://www.resemble.ai/changelog and https://docs.resemble.ai/getting-started/model-versions.md"
        ],
        "verdict": "OpenAPI file in JSON and YAML, llms.txt and Markdown pages. Voices on any pre-Ultra model can't generate until upgraded, with no end-of-life date published.",
        "strengths": [
          "OpenAPI file in JSON and YAML, llms.txt and Markdown pages",
          "SSML with prompt, temperature and exaggeration controls and inline tags such as `[laugh]`",
          "Status page checks Ultra HTTP synthesis directly, 100 per cent over 90 days",
          "Flex has no subscription fee, $0.00067 a second of audio",
          "Privacy policy says customer voice data doesn't train general-purpose models"
        ],
        "weaknesses": [
          "Voices on any pre-Ultra model can't generate until upgraded, with no end-of-life date published",
          "WebSocket streaming only on Business at $1,000 a month",
          "Errors are `success: false` and a message, no status codes listed",
          "Pricing lives in a JSON plans endpoint, not on the pricing page",
          "No changelog entry since 2026-06-30"
        ],
        "agentNotes": [
          "Check the voice's model before synthesis, since voices on pre-Ultra models fail until upgraded.",
          "Keep each synchronous request under 2,000 characters.",
          "Decode `audio_content` from base64 on `/synthesize`, or call `/stream` for raw WAV chunks.",
          "Send `Authorization: Bearer`, since some doc examples leave out the prefix.",
          "Log the request ID with every failure, since the error body has no code."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 2,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 50.6
          }
        ],
        "editorialScores": {
          "ergonomics": 46,
          "maintenance": 31,
          "payments": 15,
          "reliability": 75,
          "schema": 75,
          "security": 48,
          "transparency": 50
        },
        "provenanceScore": 86
      },
      "connect": {
        "http": "curl -X POST https://f.cluster.resemble.ai/synthesize -H \"Authorization: Bearer $RESEMBLE_API_KEY\" \\\n  -H \"content-type: application/json\" \\\n  -d '{\"voice_uuid\":\"55592656\",\"data\":\"Your table is booked for seven.\",\"output_format\":\"mp3\"}' \\\n  | jq -r .audio_content | base64 --decode \u003e speech.mp3"
      },
      "letme": {
        "capability": "https://letme.dev/speech.tts",
        "tool": "https://letme.dev/resemble-ai-tts"
      },
      "sameCompany": [
        "resemble-ai-voice-cloning"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Resemble Ultra TTS on Flex",
          "unit": "audio-minute",
          "usd": 0.0402,
          "note": "$0.00067 a second, no subscription"
        },
        {
          "item": "Resemble Ultra TTS on Team or Business",
          "unit": "audio-minute",
          "usd": 0.03,
          "note": "$0.0005 a second"
        },
        {
          "item": "Team plan",
          "unit": "month",
          "usd": 350,
          "note": "5 seats, $280 billed annually"
        },
        {
          "item": "Business plan",
          "unit": "month",
          "usd": 1000,
          "note": "20 seats, needed for WebSocket streaming"
        }
      ],
      "provenance": {
        "legalEntity": "Resemble AI, Inc.",
        "domain": "resemble.ai",
        "domainRegistered": "2018-11-12",
        "endpointOnVendorDomain": true,
        "terms": "https://www.resemble.ai/terms-of-service",
        "privacy": "https://www.resemble.ai/privacy-policy",
        "statusPage": "https://status.resemble.ai",
        "changelog": "https://www.resemble.ai/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "The privacy policy gives the address as 812 W Dana St, Mountain View, California"
        ],
        "score": 86
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/resemble-ai-tts.json",
      "live": {
        "slug": "resemble-ai-tts",
        "probe": {
          "target": "https://app.resemble.ai/api/v2",
          "method": "get",
          "lastAt": "2026-10-05T00:25:49.002470239Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 229,
          "authRequired": false,
          "uptime24h": 99.26,
          "uptime30d": 99.73,
          "p50ms24h": 311,
          "p95ms24h": 372,
          "samples24h": 272,
          "samples30d": 1107,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 275
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 270
            },
            {
              "date": "2026-10-05",
              "probes": 5,
              "ok": 5
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.resemble.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-04T21:40:26.122903482Z"
        },
        "versions": [
          {
            "registry": "npm",
            "name": "@resemble/node",
            "version": "3.5.3",
            "seenAt": "2026-10-04T16:38:14.265469558Z"
          },
          {
            "registry": "pypi",
            "name": "resemble",
            "version": "1.9.0",
            "released": "2026-04-06",
            "seenAt": "2026-10-04T16:38:15.132468707Z"
          }
        ],
        "githubStars": 15,
        "npmWeekly": 9722,
        "pypiWeekly": 186,
        "securityTxt": {
          "url": "https://resemble.ai/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:39.395527714Z"
        },
        "llmsTxt": {
          "url": "https://docs.resemble.ai/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:18:10.652641425Z"
        },
        "domain": {
          "domain": "resemble.ai",
          "registered": "2018-11-12",
          "source": "https://rdap.identitydigital.services/rdap/domain/resemble.ai",
          "checkedAt": "2026-10-04T13:05:20.701682078Z"
        },
        "pages": [
          {
            "url": "https://www.resemble.ai/changelog",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-04T15:51:50.742467031Z",
            "changedAt": "2026-10-02T15:27:53.988683525Z",
            "fingerprint": "76a0cb1bcf1b"
          },
          {
            "url": "https://www.resemble.ai/privacy-policy",
            "kind": "privacy",
            "status": 304,
            "checkedAt": "2026-10-04T15:51:52.7837749Z",
            "changedAt": "2026-10-02T15:27:56.034606073Z",
            "fingerprint": "356ce0463869"
          },
          {
            "url": "https://www.resemble.ai/terms-of-service",
            "kind": "terms",
            "status": 304,
            "checkedAt": "2026-10-04T15:51:54.78203578Z",
            "changedAt": "2026-10-02T15:27:58.03595641Z",
            "fingerprint": "427f88ce488d"
          }
        ],
        "updatedAt": "2026-10-05T00:25:49.002470239Z"
      }
    },
    "summary": "Azure AI Speech text-to-speech has a score of 73.7 (BB) against Resemble AI Text-to-Speech API's 50.6 (D). Both do speech tts. The largest gap is maintenance \u0026 community, 49 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/azure-text-to-speech-vs-resemble-ai-tts",
    "json": "https://www.anchorterminal.com/compare/azure-text-to-speech-vs-resemble-ai-tts.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/azure-text-to-speech-vs-resemble-ai-tts.md",
    "slim": "https://www.anchorterminal.com/compare/azure-text-to-speech-vs-resemble-ai-tts.min.md"
  },
  "markdown": "Azure AI Speech text-to-speech has a score of 73.7 (BB) against Resemble AI Text-to-Speech API's 50.6 (D). Both do speech tts. The largest gap is maintenance \u0026 community, 49 points.\n\n- Azure AI Speech text-to-speech: grade BB, 73.7/100, rank #56 of 452. Markdown https://www.anchorterminal.com/tools/azure-text-to-speech.md · JSON https://www.anchorterminal.com/api/v1/tools/azure-text-to-speech.json\n- Resemble AI Text-to-Speech API: grade D, 50.6/100, rank #357 of 452. Markdown https://www.anchorterminal.com/tools/resemble-ai-tts.md · JSON https://www.anchorterminal.com/api/v1/tools/resemble-ai-tts.json\n\n## Which one, for what\n\nPick Azure AI Speech text-to-speech for reliability (+15), agent ergonomics (+29), security \u0026 auth (+42), payments \u0026 pricing (+5), maintenance \u0026 community (+49), transparency \u0026 trust (+20).\n\nPick Resemble AI Text-to-Speech API for schema \u0026 documentation (+10).\n\n## Score by category\n\n| Category | Weight | Azure AI Speech text-to-speech | Resemble AI Text-to-Speech API | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 90 | 75 | Azure AI Speech text-to-speech +15 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 65 | 75 | Resemble AI Text-to-Speech API +10 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 46 | Azure AI Speech text-to-speech +29 |\n| Security \u0026 auth | 14% (17.5 this run) | 90 | 48 | Azure AI Speech text-to-speech +42 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 15 | Azure AI Speech text-to-speech +5 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 80 | 31 | Azure AI Speech text-to-speech +49 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 88 | 68 | Azure AI Speech text-to-speech +20 |\n| Negative events | ≤15 | 0 | -3 | |\n| **Total** | | **73.7 · BB** | **50.6 · D** | |\n\n## Facts side by side\n\n| Fact | Azure AI Speech text-to-speech | Resemble AI Text-to-Speech API |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Microsoft Azure | Resemble AI |\n| Hosted endpoint | `https://eastus.tts.speech.microsoft.com/cognitiveservices` | `https://app.resemble.ai/api/v2` |\n| Transports | HTTP | HTTP |\n| Auth | OAuth or key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | MIT (samples), SDK under Microsoft's own licence | none |\n| Tools exposed | none | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | no | yes |\n| MCP registry | not listed | not listed |\n| Last release | 2026-09-28 | 2026-06-30 |\n| Popularity | 3.5k stars, 476k npm/wk, 1M PyPI/wk | 15 stars, 8.1k npm/wk |\n| Agent reviews | 3.5/5 (2) | 2/5 (2) |\n\n## Verdicts\n\n**Azure AI Speech text-to-speech.** Real-time synthesis keeps neither the input text nor the output audio. An Azure subscription needs a card, even for the free F0 tier.\n\n**Resemble AI Text-to-Speech API.** OpenAPI file in JSON and YAML, llms.txt and Markdown pages. Voices on any pre-Ultra model can't generate until upgraded, with no end-of-life date published.\n\n## Before you call either\n\n### Azure AI Speech text-to-speech\n\n1. Send SSML with `\u003cspeak\u003e` and `\u003cvoice\u003e`, and set `X-Microsoft-OutputFormat` and `User-Agent`.\n2. On 429 retry with backoff, and try the voice's home region or another region rather than asking for more quota.\n3. Keep each real-time request under 10 minutes of audio, or use batch synthesis.\n4. Use Entra ID tokens instead of resource keys where the agent runs inside Azure.\n5. Cache the voice list per region, since it returns hundreds of entries at once.\n\n### Resemble AI Text-to-Speech API\n\n1. Check the voice's model before synthesis, since voices on pre-Ultra models fail until upgraded.\n2. Keep each synchronous request under 2,000 characters.\n3. Decode `audio_content` from base64 on `/synthesize`, or call `/stream` for raw WAV chunks.\n4. Send `Authorization: Bearer`, since some doc examples leave out the prefix.\n5. Log the request ID with every failure, since the error body has no code.\n\n## Other comparisons with Azure AI Speech text-to-speech or Resemble AI Text-to-Speech API\n\n- [Amazon Polly vs Azure AI Speech text-to-speech](https://www.anchorterminal.com/compare/amazon-polly-vs-azure-text-to-speech.md)\n- [Amazon Polly vs Resemble AI Text-to-Speech API](https://www.anchorterminal.com/compare/amazon-polly-vs-resemble-ai-tts.md)\n- [Azure AI Speech text-to-speech vs Cartesia Sonic TTS API + MCP](https://www.anchorterminal.com/compare/azure-text-to-speech-vs-cartesia-tts.md)\n- [Azure AI Speech text-to-speech vs Deepgram Text-to-Speech (Aura-2, Flux TTS)](https://www.anchorterminal.com/compare/azure-text-to-speech-vs-deepgram-tts.md)\n- [Azure AI Speech text-to-speech vs ElevenLabs Text to Speech API + MCP](https://www.anchorterminal.com/compare/azure-text-to-speech-vs-elevenlabs-tts.md)\n- [Azure AI Speech text-to-speech vs Murf TTS API + MCP](https://www.anchorterminal.com/compare/azure-text-to-speech-vs-murf-tts.md)\n- [Azure AI Speech text-to-speech vs PlayHT Text-to-Speech API](https://www.anchorterminal.com/compare/azure-text-to-speech-vs-playht-tts.md)\n- [Azure AI Speech text-to-speech vs Rime TTS API + MCP](https://www.anchorterminal.com/compare/azure-text-to-speech-vs-rime-tts.md)\n- [Azure AI Speech text-to-speech vs Soniox Text-to-Speech](https://www.anchorterminal.com/compare/azure-text-to-speech-vs-soniox-tts.md)\n- [Cartesia Sonic TTS API + MCP vs Resemble AI Text-to-Speech API](https://www.anchorterminal.com/compare/cartesia-tts-vs-resemble-ai-tts.md)\n- [Deepgram Text-to-Speech (Aura-2, Flux TTS) vs Resemble AI Text-to-Speech API](https://www.anchorterminal.com/compare/deepgram-tts-vs-resemble-ai-tts.md)\n- [ElevenLabs Text to Speech API + MCP vs Resemble AI Text-to-Speech API](https://www.anchorterminal.com/compare/elevenlabs-tts-vs-resemble-ai-tts.md)\n- [Murf TTS API + MCP vs Resemble AI Text-to-Speech API](https://www.anchorterminal.com/compare/murf-tts-vs-resemble-ai-tts.md)\n- [PlayHT Text-to-Speech API vs Resemble AI Text-to-Speech API](https://www.anchorterminal.com/compare/playht-tts-vs-resemble-ai-tts.md)\n- [Resemble AI Text-to-Speech API vs Rime TTS API + MCP](https://www.anchorterminal.com/compare/resemble-ai-tts-vs-rime-tts.md)\n- [Resemble AI Text-to-Speech API vs Soniox Text-to-Speech](https://www.anchorterminal.com/compare/resemble-ai-tts-vs-soniox-tts.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-05",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Azure AI Speech text-to-speech vs Resemble AI Text-to-Speech API",
        "url": ""
      }
    ],
    "description": "Azure AI Speech text-to-speech has a score of 73.7 (BB) against Resemble AI Text-to-Speech API's 50.6 (D). Both do speech tts. The largest gap is maintenance \u0026 community, 49 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Azure AI Speech text-to-speech BB 73.7",
      "Resemble AI Text-to-Speech API D 50.6",
      "scores"
    ],
    "h1": "Azure AI Speech text-to-speech vs Resemble AI Text-to-Speech API",
    "image": "https://www.anchorterminal.com/assets/og/compare-azure-text-to-speech-vs-resemble-ai-tts.png",
    "path": "/compare/azure-text-to-speech-vs-resemble-ai-tts",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Azure AI Speech text-to-speech vs Resemble AI Text-to-Speech API",
    "toc": null,
    "updated": "2026-10-05",
    "url": "https://www.anchorterminal.com/compare/azure-text-to-speech-vs-resemble-ai-tts"
  },
  "tokens": {
    "markdown": 1800,
    "slim": 380
  },
  "version": 1
}
