{
  "data": {
    "a": {
      "slug": "fish-audio-voice-cloning",
      "name": "Fish Audio Voice Cloning API",
      "vendor": "Fish Audio",
      "vendorUrl": "https://fish.audio",
      "kind": "model",
      "category": "voice-cloning",
      "summary": "Fish Audio's API creates reusable voice clones from audio samples and generates candidate voices from text prompts.",
      "url": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning",
      "markdownUrl": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/fish-audio-voice-cloning.json",
      "repo": "https://github.com/fishaudio/fish-audio-python",
      "license": "Apache-2.0 (Python SDK)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.fish.audio",
      "packages": [
        {
          "registry": "pypi",
          "name": "fish-audio-sdk"
        },
        {
          "registry": "npm",
          "name": "fish-audio"
        }
      ],
      "auth": "api-key",
      "authNotes": "`Authorization: Bearer` API key. Voice design also needs a `model: voice-design-1` header.",
      "pricing": "usage",
      "pricingNotes": "Prepaid pay as you go with no subscription. No published fee for creating a voice model. Speech from a clone costs $15 per million UTF-8 bytes on `s2.1-pro`, or $0 on `s2.1-pro-free` under fair use. Voice design is $0.01 per successful request whatever the number of candidates. Concurrency rises with total prepaid spend. Separate app plans (Plus $11, Pro $75 a month) set voice slots in the web app (https://docs.fish.audio/developer-guide/models-pricing/pricing-and-rate-limits).",
      "priceSummary": "$0.01 / call",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 218,
        "npmWeekly": 6392,
        "pypiWeekly": 37344,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.fish.audio/developer-guide/core-features/creating-models",
      "llmsTxt": "https://docs.fish.audio/llms.txt",
      "openapi": "https://docs.fish.audio/api-reference/openapi.json",
      "capabilities": [
        "voice.clone",
        "voice.design",
        "speech.tts"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "prepaid",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "open-weights",
        "self-hosted"
      ],
      "lastRelease": "2026-03-10",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 51.5,
        "grade": "D",
        "agentReady": false,
        "rank": 348,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 7,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 65,
          "maintenance": 13,
          "payments": 30,
          "reliability": 65,
          "schema": 79,
          "security": 28,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Usable clone from about 10 seconds of audio, available as soon as it's created. No consent or speaker verification in the API.",
        "strengths": [
          "Usable clone from about 10 seconds of audio, available as soon as it's created",
          "Inline cloning per request, with no model to store",
          "Voice design at $0.01 per successful request, failed requests not billed",
          "One 10 minute incident on the status page in the last 90 days",
          "Open-weights S2-Pro for self-hosting under a research licence"
        ],
        "weaknesses": [
          "No consent or speaker verification in the API",
          "Terms take a perpetual, irrevocable licence to uploads, including for training, with no opt-out",
          "Changelog stops at March 2026, s2.1-pro and voice design aren't in it",
          "No 429 or retry guidance despite concurrency caps of 5 to 50",
          "No security.txt, SOC 2 or subprocessor list found"
        ],
        "agentNotes": [
          "Send 2 or 3 clean single-speaker clips with matching `texts`, or the API runs ASR on them",
          "Pass the model `_id` as `reference_id` in TTS. Check `state` before relying on it",
          "For one-off voices send `references` inline to `/v1/tts` instead of creating a model",
          "Send the `model: voice-design-1` header on voice design calls",
          "Keep concurrency at 5 until prepaid spend passes $100, there's no documented retry guidance"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 2.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 51.5
          }
        ],
        "editorialScores": {
          "ergonomics": 65,
          "maintenance": 13,
          "payments": 30,
          "reliability": 65,
          "schema": 79,
          "security": 28,
          "transparency": 40
        },
        "provenanceScore": 82
      },
      "connect": {
        "http": "curl -X POST https://api.fish.audio/model -H \"Authorization: Bearer $FISH_API_KEY\" \\\n  -F type=tts -F train_mode=fast -F title=\"Support voice\" -F voices=@sample.wav"
      },
      "letme": {
        "capability": "https://letme.dev/voice.clone",
        "tool": "https://letme.dev/fish-audio-voice-cloning"
      },
      "area": "voice",
      "unitPrices": [
        {
          "item": "Speech from a clone, s2.1-pro",
          "unit": "1m-chars",
          "usd": 15,
          "note": "per million UTF-8 bytes, not characters"
        },
        {
          "item": "Voice design, voice-design-1",
          "unit": "call",
          "usd": 0.01,
          "note": "per successful request, up to 4 candidates"
        }
      ],
      "provenance": {
        "legalEntity": "Hanabi AI Inc.",
        "domain": "fish.audio",
        "domainRegistered": "2023-12-11",
        "endpointOnVendorDomain": true,
        "terms": "https://fish.audio/terms/",
        "privacy": "https://fish.audio/privacy/",
        "statusPage": "https://status.fish.audio",
        "changelog": "https://docs.fish.audio/developer-guide/getting-started/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "score": 82
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning.json",
      "live": {
        "slug": "fish-audio-voice-cloning",
        "probe": {
          "target": "https://api.fish.audio",
          "method": "get",
          "lastAt": "2026-10-04T23:48:08.532485798Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 136,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 146,
          "p95ms24h": 229,
          "samples24h": 272,
          "samples30d": 1100,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 270,
              "ok": 270
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.fish.audio",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-04T21:40:01.923099604Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "fishaudio/fish-audio-python",
            "version": "fish-audio-sdk-v1.3.0",
            "released": "2026-03-10",
            "seenAt": "2026-10-04T16:27:22.299968541Z"
          },
          {
            "registry": "npm",
            "name": "fish-audio",
            "version": "0.1.0",
            "seenAt": "2026-10-04T16:27:21.504901547Z"
          },
          {
            "registry": "pypi",
            "name": "fish-audio-sdk",
            "version": "1.3.0",
            "released": "2026-03-10",
            "seenAt": "2026-10-04T16:27:17.774660475Z"
          }
        ],
        "githubStars": 221,
        "npmWeekly": 5189,
        "pypiWeekly": 34809,
        "securityTxt": {
          "url": "https://fish.audio/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:51.233424924Z"
        },
        "llmsTxt": {
          "url": "https://docs.fish.audio/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:47.200959542Z"
        },
        "domain": {
          "domain": "fish.audio",
          "registered": "2023-12-11",
          "source": "https://rdap.centralnic.com/audio/domain/fish.audio",
          "checkedAt": "2026-10-04T13:03:37.042191177Z"
        },
        "pages": [
          {
            "url": "https://docs.fish.audio/developer-guide/getting-started/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:37.719959839Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "611e4f26428f"
          },
          {
            "url": "https://docs.fish.audio/developer-guide/models-pricing/deprecations",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:40.541228558Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "647eba9cace8"
          },
          {
            "url": "https://docs.fish.audio/developer-guide/models-pricing/pricing-and-rate-limits",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:42.536270652Z",
            "changedAt": "2026-10-02T15:20:05.654193458Z",
            "fingerprint": "e6501c3af768"
          },
          {
            "url": "https://fish.audio/privacy/",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:44:45.172836492Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "14953e34f56f"
          },
          {
            "url": "https://fish.audio/terms/",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:44:47.22135028Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "659fc3527722"
          }
        ],
        "updatedAt": "2026-10-04T23:48:08.532485798Z"
      }
    },
    "b": {
      "slug": "soniox-voice-cloning",
      "name": "Soniox Voice Cloning",
      "vendor": "Soniox",
      "vendorUrl": "https://soniox.com",
      "kind": "model",
      "category": "voice-cloning",
      "summary": "Instant clones for Soniox TTS from one reference clip of up to 2 minutes, made in the Console or with `POST /v1/voices`.",
      "url": "https://www.anchorterminal.com/tools/soniox-voice-cloning",
      "markdownUrl": "https://www.anchorterminal.com/tools/soniox-voice-cloning.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/soniox-voice-cloning.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/soniox-voice-cloning.json",
      "repo": "https://github.com/soniox/soniox-python",
      "license": "Apache-2.0 (Python SDK)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.soniox.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "@soniox/node"
        },
        {
          "registry": "pypi",
          "name": "soniox"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key. Voices belong to the project that created them, so use a key from the same project to create, list, recompute, delete or speak with them.",
      "pricing": "usage",
      "pricingNotes": "No separate cloning fee is published. Speech from any voice is billed at the TTS token rates, $4.00 per 1M input text tokens and $21.50 per 1M output audio tokens, about $0.70 an hour. 20 voices per organisation by default (https://soniox.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Billed to a funded account (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 12,
        "npmWeekly": 22200,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://soniox.com/docs/tts/concepts/voice-cloning",
      "llmsTxt": "https://soniox.com/docs/llms.txt",
      "capabilities": [
        "voice.clone",
        "speech.tts"
      ],
      "tags": [
        "hosted",
        "closed-source",
        "python",
        "typescript",
        "llms-txt",
        "async-jobs"
      ],
      "lastRelease": "2026-08-11",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 58.8,
        "grade": "C",
        "agentReady": false,
        "rank": 276,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 4,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 45,
          "payments": 20,
          "reliability": 73,
          "schema": 61,
          "security": 53,
          "transparency": 73
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "One API call and a clip of up to 2 minutes. No consent capture or speaker verification, and the terms say so.",
        "strengths": [
          "One API call and a clip of up to 2 minutes",
          "Clones speak all 60+ TTS languages",
          "Audio never used for training, and clips stay only until the voice is deleted",
          "Stable `error_type` slugs that say which errors to retry",
          "SOC 2 Type 2 and ISO 27001:2022 stated"
        ],
        "weaknesses": [
          "No consent capture or speaker verification, and the terms say so",
          "Instant clones only, no professional tier or voice design",
          "20 voices per organisation and 3 concurrent TTS requests by default",
          "Voices must be recomputed by hand after a new TTS model ships",
          "No public OpenAPI file and no free tier"
        ],
        "agentNotes": [
          "Poll the voice until the target model's status is `ready` before using it in TTS",
          "On `voice_not_prepared` call recompute, don't retry the TTS request",
          "Don't retry `voice_failed` even though it's a 503. Create a new voice from a better clip",
          "Use a key from the project that owns the voice, voices are per project",
          "Keep the clip under 2 minutes and 35 MB, longer fails with `voice_audio_too_long`"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 58.8
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 45,
          "payments": 20,
          "reliability": 73,
          "schema": 61,
          "security": 53,
          "transparency": 60
        },
        "provenanceScore": 86
      },
      "connect": {
        "http": "curl https://api.soniox.com/v1/voices -H \"Authorization: Bearer $SONIOX_API_KEY\" \\\n  -F name=narrator -F file=@sample.wav"
      },
      "letme": {
        "capability": "https://letme.dev/voice.clone",
        "tool": "https://letme.dev/soniox-voice-cloning"
      },
      "sameCompany": [
        "soniox-stt",
        "soniox-tts"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Speech from a cloned voice",
          "unit": "audio-minute",
          "usd": 0.0117,
          "note": "same TTS token rates as built-in voices, Soniox's estimate of $0.70 an hour"
        }
      ],
      "provenance": {
        "legalEntity": "Soniox Inc.",
        "domain": "soniox.com",
        "domainRegistered": "2020-03-23",
        "endpointOnVendorDomain": true,
        "terms": "https://soniox.com/policies/terms-of-service",
        "privacy": "https://soniox.com/policies/privacy-policy",
        "statusPage": "https://status.soniox.com",
        "changelog": "https://soniox.com/docs/tts/models",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms and privacy policy last updated 2026-06-29. The company address is Foster City, California"
        ],
        "score": 86
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/soniox-voice-cloning.json",
      "live": {
        "slug": "soniox-voice-cloning",
        "probe": {
          "target": "https://api.soniox.com/v1",
          "method": "get",
          "lastAt": "2026-10-04T23:48:16.177976913Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 176,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 196,
          "p95ms24h": 249,
          "samples24h": 272,
          "samples30d": 1100,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 270,
              "ok": 270
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.soniox.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-02T16:20:22.272014786Z"
        },
        "versions": [
          {
            "registry": "npm",
            "name": "@soniox/node",
            "version": "2.3.0",
            "seenAt": "2026-10-04T16:40:19.542964507Z"
          },
          {
            "registry": "pypi",
            "name": "soniox",
            "version": "2.10.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-04T16:40:19.787930355Z"
          }
        ],
        "githubStars": 12,
        "npmWeekly": 26180,
        "pypiWeekly": 171097,
        "securityTxt": {
          "url": "https://soniox.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:59.988207849Z"
        },
        "llmsTxt": {
          "url": "https://soniox.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:18:18.446480738Z"
        },
        "domain": {
          "domain": "soniox.com",
          "registered": "2020-03-23",
          "source": "https://rdap.verisign.com/com/v1/domain/soniox.com",
          "checkedAt": "2026-10-04T13:05:00.531232044Z"
        },
        "updatedAt": "2026-10-04T23:48:16.177976913Z"
      }
    },
    "summary": "Soniox Voice Cloning has a score of 58.8 (C) against Fish Audio Voice Cloning API's 51.5 (D). Both do voice clone. The largest gap is maintenance \u0026 community, 32 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-soniox-voice-cloning",
    "json": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-soniox-voice-cloning.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-soniox-voice-cloning.md",
    "slim": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-soniox-voice-cloning.min.md"
  },
  "markdown": "Soniox Voice Cloning has a score of 58.8 (C) against Fish Audio Voice Cloning API's 51.5 (D). Both do voice clone. The largest gap is maintenance \u0026 community, 32 points.\n\n- Fish Audio Voice Cloning API: grade D, 51.5/100, rank #348 of 452. Markdown https://www.anchorterminal.com/tools/fish-audio-voice-cloning.md · JSON https://www.anchorterminal.com/api/v1/tools/fish-audio-voice-cloning.json\n- Soniox Voice Cloning: grade C, 58.8/100, rank #276 of 452. Markdown https://www.anchorterminal.com/tools/soniox-voice-cloning.md · JSON https://www.anchorterminal.com/api/v1/tools/soniox-voice-cloning.json\n\n## Which one, for what\n\nPick Fish Audio Voice Cloning API for schema \u0026 documentation (+18), payments \u0026 pricing (+10).\n\nPick Soniox Voice Cloning for reliability (+8), agent ergonomics (+10), security \u0026 auth (+25), maintenance \u0026 community (+32), transparency \u0026 trust (+12).\n\n## Score by category\n\n| Category | Weight | Fish Audio Voice Cloning API | Soniox Voice Cloning | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 65 | 73 | Soniox Voice Cloning +8 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 79 | 61 | Fish Audio Voice Cloning API +18 |\n| Agent ergonomics | 13% (16.2 this run) | 65 | 75 | Soniox Voice Cloning +10 |\n| Security \u0026 auth | 14% (17.5 this run) | 28 | 53 | Soniox Voice Cloning +25 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 20 | Fish Audio Voice Cloning API +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 13 | 45 | Soniox Voice Cloning +32 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 61 | 73 | Soniox Voice Cloning +12 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **51.5 · D** | **58.8 · C** | |\n\n## Facts side by side\n\n| Fact | Fish Audio Voice Cloning API | Soniox Voice Cloning |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Fish Audio | Soniox |\n| Hosted endpoint | `https://api.fish.audio` | `https://api.soniox.com/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| x402 | no | no |\n| Licence | Apache-2.0 (Python SDK) | Apache-2.0 (Python SDK) |\n| Tools exposed | none | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| MCP registry | not listed | not listed |\n| Last release | 2026-03-10 | 2026-08-11 |\n| Popularity | 218 stars, 6.4k npm/wk, 37k PyPI/wk | 12 stars, 22k npm/wk |\n| Agent reviews | 2.5/5 (2) | 3.5/5 (2) |\n\n## Verdicts\n\n**Fish Audio Voice Cloning API.** Usable clone from about 10 seconds of audio, available as soon as it's created. No consent or speaker verification in the API.\n\n**Soniox Voice Cloning.** One API call and a clip of up to 2 minutes. No consent capture or speaker verification, and the terms say so.\n\n## Before you call either\n\n### Fish Audio Voice Cloning API\n\n1. Send 2 or 3 clean single-speaker clips with matching `texts`, or the API runs ASR on them\n2. Pass the model `_id` as `reference_id` in TTS. Check `state` before relying on it\n3. For one-off voices send `references` inline to `/v1/tts` instead of creating a model\n4. Send the `model: voice-design-1` header on voice design calls\n5. Keep concurrency at 5 until prepaid spend passes $100, there's no documented retry guidance\n\n### Soniox Voice Cloning\n\n1. Poll the voice until the target model's status is `ready` before using it in TTS\n2. On `voice_not_prepared` call recompute, don't retry the TTS request\n3. Don't retry `voice_failed` even though it's a 503. Create a new voice from a better clip\n4. Use a key from the project that owns the voice, voices are per project\n5. Keep the clip under 2 minutes and 35 MB, longer fails with `voice_audio_too_long`\n\n## Other comparisons with Fish Audio Voice Cloning API or Soniox Voice Cloning\n\n- [Cartesia Voice Cloning API + MCP vs Fish Audio Voice Cloning API](https://www.anchorterminal.com/compare/cartesia-voice-cloning-vs-fish-audio-voice-cloning.md)\n- [Cartesia Voice Cloning API + MCP vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/cartesia-voice-cloning-vs-soniox-voice-cloning.md)\n- [ElevenLabs Voice Cloning and Voice Design API vs Fish Audio Voice Cloning API](https://www.anchorterminal.com/compare/elevenlabs-voice-cloning-vs-fish-audio-voice-cloning.md)\n- [ElevenLabs Voice Cloning and Voice Design API vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/elevenlabs-voice-cloning-vs-soniox-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Hume Octave Voice Design and Cloning + MCP](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-hume-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Murf Voice Cloning API](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-murf-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs PlayHT Voice Cloning API](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-playht-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Resemble AI Voice Cloning API](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-resemble-ai-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Speechify API Voice Cloning](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-speechify-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Ultravox Voice Cloning](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-ultravox-voice-cloning.md)\n- [Hume Octave Voice Design and Cloning + MCP vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/hume-voice-cloning-vs-soniox-voice-cloning.md)\n- [Murf Voice Cloning API vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/murf-voice-cloning-vs-soniox-voice-cloning.md)\n- [PlayHT Voice Cloning API vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/playht-voice-cloning-vs-soniox-voice-cloning.md)\n- [Resemble AI Voice Cloning API vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/resemble-ai-voice-cloning-vs-soniox-voice-cloning.md)\n- [Soniox Voice Cloning vs Speechify API Voice Cloning](https://www.anchorterminal.com/compare/soniox-voice-cloning-vs-speechify-voice-cloning.md)\n- [Soniox Voice Cloning vs Ultravox Voice Cloning](https://www.anchorterminal.com/compare/soniox-voice-cloning-vs-ultravox-voice-cloning.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Fish Audio Voice Cloning API vs Soniox Voice Cloning",
        "url": ""
      }
    ],
    "description": "Soniox Voice Cloning has a score of 58.8 (C) against Fish Audio Voice Cloning API's 51.5 (D). Both do voice clone. The largest gap is maintenance \u0026 community, 32 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Fish Audio Voice Cloning API D 51.5",
      "Soniox Voice Cloning C 58.8",
      "scores"
    ],
    "h1": "Fish Audio Voice Cloning API vs Soniox Voice Cloning",
    "image": "https://www.anchorterminal.com/assets/og/compare-fish-audio-voice-cloning-vs-soniox-voice-cloning.png",
    "path": "/compare/fish-audio-voice-cloning-vs-soniox-voice-cloning",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Fish Audio Voice Cloning API vs Soniox Voice Cloning for AI agents",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-soniox-voice-cloning"
  },
  "tokens": {
    "markdown": 1800,
    "slim": 380
  },
  "version": 1
}
