{
  "data": {
    "a": {
      "slug": "adobe-pdf-extract",
      "name": "Adobe PDF Services / PDF Extract API",
      "vendor": "Adobe",
      "vendorUrl": "https://developer.adobe.com/document-services/",
      "kind": "http-api",
      "category": "document-extraction",
      "summary": "Adobe's REST API for reading and changing PDFs.",
      "url": "https://www.anchorterminal.com/tools/adobe-pdf-extract",
      "markdownUrl": "https://www.anchorterminal.com/tools/adobe-pdf-extract.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/adobe-pdf-extract.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/adobe-pdf-extract.json",
      "repo": "https://github.com/adobe/pdfservices-python-sdk-samples",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://pdf-services.adobe.io",
      "packages": [
        {
          "registry": "pypi",
          "name": "pdfservices-sdk"
        },
        {
          "registry": "npm",
          "name": "@adobe/pdfservices-node-sdk"
        }
      ],
      "auth": "oauth",
      "authNotes": "OAuth server-to-server. POST client_id and client_secret to /token for an access token, then send it as Bearer plus the client id in `x-api-key`. JWT service accounts are deprecated.",
      "pricing": "freemium",
      "pricingNotes": "Free tier of 500 Document Transactions a month with no card. Extract and PDF to Markdown use 1 transaction per 5 pages, most other operations 1 per 50 pages, Auto-Tag 10 a page. Paid use is sold through sales on volume or enterprise terms, and no per-transaction price is published (https://developer.adobe.com/document-services/pricing/main/).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 in docs or pricing (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 165,
        "npmWeekly": 81553,
        "pypiWeekly": 58026,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://developer.adobe.com/document-services/docs/overview/pdf-extract-api/",
      "openapi": "https://raw.githubusercontent.com/AdobeDocs/pdfservices-api-documentation/main/static/openapi.json",
      "capabilities": [
        "docs.parse",
        "docs.ocr",
        "docs.extract",
        "docs.tables",
        "pdf.convert",
        "pdf.merge",
        "pdf.forms",
        "pdf.generate",
        "pdf.extract"
      ],
      "tags": [
        "hosted",
        "freemium",
        "no-card",
        "free-tier",
        "enterprise",
        "closed-source",
        "openapi",
        "webhooks",
        "python",
        "typescript",
        "async-jobs"
      ],
      "lastRelease": "2026-08-10",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.9,
        "grade": "C",
        "agentReady": false,
        "rank": 578,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 9,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 57,
          "maintenance": 53,
          "payments": 20,
          "reliability": 50,
          "schema": 80,
          "security": 55,
          "transparency": 78
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Public OpenAPI 3.0.1 spec covering Extract, PDF to Markdown and 20 other operations. No published paid price, paid use goes through sales.",
        "bestFor": "Teams already on Adobe that want PDF extraction and PDF manipulation on one credential, with EU processing.",
        "strengths": [
          "Public OpenAPI 3.0.1 spec covering Extract, PDF to Markdown and 20 other operations",
          "Extract error codes name the cause, such as DISQUALIFIED_SCAN_PAGE_LIMIT or BAD_PDF_COMPLEX_TABLE",
          "500 free Document Transactions a month with no card",
          "Files kept 24 hours by default, deletable at once, or never stored when you pass signed URLs",
          "security.txt and a public bug bounty on Intigriti"
        ],
        "weaknesses": [
          "No published paid price, paid use goes through sales",
          "Four steps per job (token, upload, operation, poll) and no idempotency key",
          "25 requests a minute on the free tier, no Retry-After or backoff guidance",
          "Node.js SDK unchanged since November 2024 and no llms.txt",
          "Status page renders only with JavaScript"
        ],
        "agentNotes": [
          "Cache the access token and reuse it until it expires",
          "Budget 1 Document Transaction per 5 pages for Extract and PDF to Markdown, so 500 free transactions cover about 2,500 pages",
          "Pass `notifiers` with a callback URL instead of polling the job status",
          "Use `pdf-services-ew1.adobe.io` when documents must stay in the EU",
          "Check `/operation/pdfproperties` first, since copy-protected PDFs fail Extract with DISQUALIFIED_PERMISSIONS"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.9
          }
        ],
        "editorialScores": {
          "ergonomics": 57,
          "maintenance": 53,
          "payments": 20,
          "reliability": 50,
          "schema": 80,
          "security": 55,
          "transparency": 60
        },
        "provenanceScore": 96
      },
      "connect": {
        "install": "pip install pdfservices-sdk   # or: npm i @adobe/pdfservices-node-sdk",
        "http": "curl https://pdf-services.adobe.io/token \\\n  -H \"content-type: application/x-www-form-urlencoded\" \\\n  --data-urlencode \"client_id=$PDF_SERVICES_CLIENT_ID\" --data-urlencode \"client_secret=$PDF_SERVICES_CLIENT_SECRET\""
      },
      "letme": {
        "capability": "https://letme.dev/docs.parse",
        "tool": "https://letme.dev/adobe-pdf-extract"
      },
      "sameCompany": [
        "adobe-firefly",
        "adobe-photoshop-api",
        "adobe-commerce",
        "adobe-acrobat-sign"
      ],
      "alsoIn": [
        "pdf-tools"
      ],
      "area": "web-data",
      "provenance": {
        "legalEntity": "Adobe Inc.",
        "domain": "adobe.com",
        "domainRegistered": "1986-11-17",
        "endpointOnVendorDomain": true,
        "terms": "https://www.adobe.com/legal/terms.html",
        "privacy": "https://www.adobe.com/privacy/policy.html",
        "statusPage": "https://status.adobe.com/products/512699",
        "changelog": "https://developer.adobe.com/document-services/docs/overview/releasenotes",
        "securityTxt": "valid",
        "checked": "2026-10-01",
        "notes": [
          "The API is served from pdf-services.adobe.io, Adobe's developer API domain, not adobe.com",
          "The status page is the PDF Services product page the docs link to, and it needs JavaScript to render"
        ],
        "score": 96
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/adobe-pdf-extract.json",
      "live": {
        "slug": "adobe-pdf-extract",
        "probe": {
          "target": "https://pdf-services.adobe.io",
          "method": "get",
          "lastAt": "2026-10-09T11:46:20.511734747Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 258,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 255,
          "p95ms24h": 294,
          "samples24h": 259,
          "samples30d": 2311,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 125,
              "ok": 125
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.adobe.com/products/512699",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-09T07:57:34.759292262Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "adobe/pdfservices-python-sdk-samples",
            "version": "v4.2.0",
            "released": "2025-07-11",
            "seenAt": "2026-10-08T15:56:37.430445927Z"
          },
          {
            "registry": "npm",
            "name": "@adobe/pdfservices-node-sdk",
            "version": "4.1.0",
            "seenAt": "2026-10-08T15:56:34.625898293Z"
          },
          {
            "registry": "pypi",
            "name": "pdfservices-sdk",
            "version": "4.3.0",
            "released": "2026-08-10",
            "seenAt": "2026-10-08T15:56:31.234161933Z"
          }
        ],
        "githubStars": 165,
        "npmWeekly": 77638,
        "pypiWeekly": 64365,
        "securityTxt": {
          "url": "https://adobe.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2027-07-30T01:00:00.000Z",
          "checkedAt": "2026-10-08T15:38:56.613662215Z"
        },
        "domain": {
          "domain": "adobe.com",
          "registered": "1986-11-17",
          "source": "https://rdap.verisign.com/com/v1/domain/adobe.com",
          "checkedAt": "2026-10-04T13:03:40.934529293Z"
        },
        "pages": [
          {
            "url": "https://developer.adobe.com/document-services/docs/overview/releasenotes",
            "kind": "deprecations",
            "status": 304,
            "checkedAt": "2026-10-08T18:17:02.71984203Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "27885eb31d8e"
          },
          {
            "url": "https://developer.adobe.com/document-services/pricing/main/",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-08T18:17:04.729379436Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "bf518d524ead"
          }
        ],
        "updatedAt": "2026-10-09T11:46:20.511734747Z"
      }
    },
    "answer": "Amazon Textract scores 73.5 (BB) on agent readiness against Adobe PDF Services / PDF Extract API's 55.9 (C), and leads in every scored category.",
    "b": {
      "slug": "amazon-textract",
      "name": "Amazon Textract",
      "vendor": "Amazon Web Services",
      "vendorUrl": "https://aws.amazon.com/textract/",
      "kind": "http-api",
      "category": "document-extraction",
      "summary": "Amazon Textract is AWS's document OCR and analysis API. It reads printed and handwritten text from scans and PDFs and returns tables, form fields, layout elements, answers to queries, and invoice, receipt and identity document fields as JSON.",
      "url": "https://www.anchorterminal.com/tools/amazon-textract",
      "markdownUrl": "https://www.anchorterminal.com/tools/amazon-textract.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/amazon-textract.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/amazon-textract.json",
      "repo": "https://github.com/aws/api-models-aws",
      "license": "Proprietary service under the AWS Customer Agreement. The AWS SDKs and the Textractor helper library are Apache-2.0",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://textract.us-east-1.amazonaws.com",
      "packages": [
        {
          "registry": "npm",
          "name": "@aws-sdk/client-textract"
        },
        {
          "registry": "pypi",
          "name": "boto3"
        },
        {
          "registry": "pypi",
          "name": "amazon-textract-textractor"
        }
      ],
      "auth": "api-key",
      "authNotes": "Requests are signed with AWS Signature Version 4 using IAM credentials, long-lived access keys or temporary credentials from AWS STS. IAM identity policies use the `textract` action prefix, with resource-level permissions for adapters only and no service-specific condition keys. Access is self-serve with an AWS account. Asynchronous jobs read from S3, and completion notices need an SNS topic and an IAM role Textract can assume.",
      "pricing": "usage",
      "pricingNotes": "Pay per page with no minimum. In US East (N. Virginia), text detection is $1.50 per 1,000 pages for the first million a month and $0.60 after. Document analysis is $15 for tables, $15 for queries, $50 for forms and $65 for tables with forms. Invoices and receipts are $10, identity documents $25 and the lending workflow $70. New AWS customers get three months of free pages, 1,000 a month for text detection and 100 a month for most analysis, and the Free Tier FAQ says most can sign up without a payment method (https://aws.amazon.com/textract/pricing/, https://aws.amazon.com/free/free-tier-faqs/).",
      "priceSummary": "$1.50 / 1k pages",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the developer guide, the API model or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 1871217,
        "pypiWeekly": 207553,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.aws.amazon.com/textract/latest/dg/what-is.html",
      "llmsTxt": "https://docs.aws.amazon.com/textract/latest/dg/llms.txt",
      "capabilities": [
        "docs.parse",
        "docs.ocr",
        "docs.extract",
        "docs.tables"
      ],
      "tags": [
        "hosted",
        "closed-source",
        "llms-txt",
        "typescript",
        "python",
        "async-jobs",
        "sla",
        "enterprise",
        "usage-priced"
      ],
      "lastRelease": "2026-10-08",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 73.5,
        "grade": "BB",
        "agentReady": true,
        "rank": 85,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 1,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 58,
          "payments": 35,
          "reliability": 96,
          "schema": 82,
          "security": 74,
          "transparency": 82
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "IAM policies limit a credential to single operations, CloudTrail logs every call, and page prices start at $1.50 per 1,000 in US East. Multipage files need S3 and an asynchronous job. Text detection covers six languages. Under the AWS Service Terms AWS may store and use documents to improve the service unless an AI services opt-out policy is set.",
        "bestFor": "Teams already on AWS with documents in S3 that need OCR, tables and form fields at volume with IAM and CloudTrail controls, and for US invoices, receipts, identity documents and mortgage packages.",
        "strengths": [
          "IAM policies grant single operations such as `textract:DetectDocumentText`, with temporary credentials, and CloudTrail logs every call without the document bytes or the response",
          "Text detection costs $1.50 per 1,000 pages in US East (N. Virginia) and $0.60 after a million pages a month, published without a login",
          "`ClientRequestToken` on the `Start` operations returns the same `JobId` for a repeated call, so a retried submission does not start a second job",
          "Every block carries a page number, a bounding box, a polygon and a confidence value from 0 to 100",
          "Amazon Textract SLA of 99.9 per cent monthly uptime per Region, and no Textract event on the AWS Health Dashboard since 10 July 2026"
        ],
        "weaknesses": [
          "AWS may store and use processed documents to improve the service, in other Regions too, unless an AI services opt-out policy is set on the AWS organisation",
          "Synchronous calls take one page of at most 10 MB. Multipage PDF and TIFF files must sit in S3 and run as asynchronous jobs",
          "Text detection covers English, French, German, Italian, Portuguese and Spanish only. Handwriting and queries are English only, and vertical text is not read",
          "Output is a flat list of JSON blocks linked by ID. No Markdown or plain-text output and no way to leave out word blocks or geometry",
          "The guide's document history stops at 21 April 2022, and the newest Textract entry in AWS's What's New feed is dated 30 June 2025"
        ],
        "agentNotes": [
          "Sign with SigV4 for service `textract` at `textract.\u003cregion\u003e.amazonaws.com`, or call `aws textract detect-document-text`. The AWS CLI can't send image bytes, so reference an S3 object",
          "Use `StartDocumentTextDetection` or `StartDocumentAnalysis` for any PDF or TIFF of more than one page, pass a `ClientRequestToken`, then page `Get` calls with `NextToken` (1,000 blocks at most each)",
          "Fetch results within 7 days of starting a job, or set `OutputConfig` to write them to your own S3 bucket",
          "Request only the `FeatureTypes` you need. Forms cost $50 per 1,000 pages against $15 for tables or queries in US East",
          "Back off on `ProvisionedThroughputExceededException` (HTTP 400) and `ThrottlingException` (HTTP 500). Neither carries Retry-After, and default quotas are 1 call a second in most Regions",
          "Set the AI services opt-out policy on the AWS organisation before sending customer documents"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 73.5
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 58,
          "payments": 35,
          "reliability": 96,
          "schema": 82,
          "security": 74,
          "transparency": 69
        },
        "provenanceScore": 95
      },
      "connect": {
        "install": "pip install boto3   # or: npm i @aws-sdk/client-textract",
        "http": "aws textract detect-document-text \\\n    --document '{\"S3Object\":{\"Bucket\":\"my-bucket\",\"Name\":\"scan.png\"}}' \\\n    --region us-east-1"
      },
      "letme": {
        "capability": "https://letme.dev/docs.parse",
        "tool": "https://letme.dev/amazon-textract"
      },
      "sameCompany": [
        "amazon-bedrock-customization",
        "agentcore-code-interpreter",
        "aws-end-user-messaging",
        "amazon-sns",
        "amazon-s3"
      ],
      "area": "web-data",
      "unitPrices": [
        {
          "item": "Text detection",
          "unit": "1k-pages",
          "usd": 1.5,
          "note": "US East (N. Virginia), first million pages a month. $0.60 after"
        },
        {
          "item": "Document analysis, tables",
          "unit": "1k-pages",
          "usd": 15,
          "note": "US East (N. Virginia), first million pages. Same price for queries"
        },
        {
          "item": "Document analysis, forms",
          "unit": "1k-pages",
          "usd": 50,
          "note": "US East (N. Virginia), first million pages. $65 with tables, $70 with tables and queries"
        },
        {
          "item": "Layout",
          "unit": "1k-pages",
          "usd": 4,
          "note": "US East (N. Virginia). Free alongside forms, tables or queries"
        },
        {
          "item": "Invoices and receipts",
          "unit": "1k-pages",
          "usd": 10,
          "note": "US East (N. Virginia), first million pages"
        },
        {
          "item": "Identity documents",
          "unit": "1k-pages",
          "usd": 25,
          "note": "US East (N. Virginia), first 100,000 pages. $10 after"
        },
        {
          "item": "Lending workflow",
          "unit": "1k-pages",
          "usd": 70,
          "note": "US East (N. Virginia), first million pages"
        }
      ],
      "provenance": {
        "legalEntity": "Amazon Web Services, Inc.",
        "domain": "amazonaws.com",
        "domainRegistered": "2005-08-18",
        "endpointOnVendorDomain": true,
        "terms": "https://aws.amazon.com/agreement/",
        "privacy": "https://aws.amazon.com/privacy/",
        "statusPage": "https://health.aws.amazon.com/health/status",
        "changelog": "https://docs.aws.amazon.com/textract/latest/dg/document-history.html",
        "securityTxt": "expired",
        "checked": "2026-10-08",
        "notes": [
          "The AWS Customer Agreement (last updated 14 August 2026) governs use of the service, with other AWS entities as contracting party by account country.",
          "Section 50 of the AWS Service Terms (last updated 1 October 2026) adds terms for AI services and names Amazon Textract. Section 50.3 lets AWS store and use processed content to improve the service unless the customer sets an AI services opt-out policy (https://aws.amazon.com/service-terms/).",
          "The AWS Privacy Notice (last updated 18 May 2026) says it does not apply to content customers process with AWS services, which the agreement governs.",
          "https://aws.amazon.com/.well-known/security.txt carries Expires 2026-09-24T16:25:03Z, so it had expired on 8 October 2026.",
          "The changelog link is the developer guide's document history, whose newest entry is 21 April 2022.",
          "RDAP for amazonaws.com gives a registration date of 2005-08-18. The API answers at textract.\u003cregion\u003e.amazonaws.com, while product and legal pages sit on aws.amazon.com."
        ],
        "score": 95
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/amazon-textract.json",
      "live": {
        "slug": "amazon-textract",
        "probe": {
          "target": "https://textract.us-east-1.amazonaws.com",
          "method": "get",
          "lastAt": "2026-10-09T11:46:21.509901567Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 335,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 335,
          "p95ms24h": 387,
          "samples24h": 44,
          "samples30d": 44,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 44,
              "ok": 44
            }
          ]
        },
        "updatedAt": "2026-10-09T11:46:21.509901567Z"
      }
    },
    "facts": [
      {
        "a": "HTTP API",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Adobe",
        "b": "Amazon Web Services",
        "name": "Vendor"
      },
      {
        "a": "https://pdf-services.adobe.io",
        "b": "https://textract.us-east-1.amazonaws.com",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "OAuth",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "none",
        "b": "Proprietary service under the AWS Customer Agreement. The AWS SDKs and the Textractor helper library are Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "no",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-08-10",
        "b": "2026-10-08",
        "name": "Last release"
      },
      {
        "a": "couldn't be read",
        "b": "2026-08-14",
        "name": "Terms last updated"
      },
      {
        "a": "2025-10-24",
        "b": "2026-05-18",
        "name": "Privacy policy last updated"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "couldn't be read",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "couldn't be read",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "165 stars, 82k npm/wk, 58k PyPI/wk",
        "b": "1.9M npm/wk, 208k PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "3.5/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Amazon Textract scores 73.5 (BB) on agent readiness against Adobe PDF Services / PDF Extract API's 55.9 (C), and leads in every scored category.",
        "question": "Which is better for AI agents, Adobe PDF Services / PDF Extract API or Amazon Textract?"
      },
      {
        "answer": "Adobe PDF Services / PDF Extract API uses an OAuth sign-in. Amazon Textract needs an API key.",
        "question": "Do Adobe PDF Services / PDF Extract API and Amazon Textract need an API key?"
      },
      {
        "answer": "Yes. Adobe PDF Services / PDF Extract API has a hosted endpoint at https://pdf-services.adobe.io and Amazon Textract at https://textract.us-east-1.amazonaws.com.",
        "question": "Can an agent call Adobe PDF Services / PDF Extract API and Amazon Textract without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": null,
        "also": [
          "Free to start without a card"
        ],
        "goodFor": "Teams already on Adobe that want PDF extraction and PDF manipulation on one credential, with EU processing.",
        "slug": "adobe-pdf-extract",
        "watchFor": "No published paid price, paid use goes through sales"
      },
      {
        "aheadOn": [
          "Reliability, 96 against 50",
          "Agent ergonomics, 70 against 57",
          "Security \u0026 auth, 74 against 55",
          "Payments \u0026 pricing, 35 against 20",
          "Maintenance \u0026 community, 58 against 53"
        ],
        "also": [
          "Agent-ready, a grade of BB or better"
        ],
        "goodFor": "Teams already on AWS with documents in S3 that need OCR, tables and form fields at volume with IAM and CloudTrail controls, and for US invoices, receipts, identity documents and mortgage packages.",
        "slug": "amazon-textract",
        "watchFor": "AWS may store and use processed documents to improve the service, in other Regions too, unless an AI services opt-out policy is set on the AWS organisation"
      }
    ],
    "job": {
      "capability": "docs.parse",
      "name": "Docs parse"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-azure-document-intelligence.json",
        "title": "Adobe PDF Services / PDF Extract API vs Azure Document Intelligence",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-azure-document-intelligence"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-extend.json",
        "title": "Adobe PDF Services / PDF Extract API vs Extend API + MCP",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-extend"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-llamaparse.json",
        "title": "Adobe PDF Services / PDF Extract API vs LlamaParse API + MCP",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-llamaparse"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-mistral-ocr.json",
        "title": "Adobe PDF Services / PDF Extract API vs Mistral OCR API",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-mistral-ocr"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-nanonets.json",
        "title": "Adobe PDF Services / PDF Extract API vs Nanonets API + MCP",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-nanonets"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-opendocrouter.json",
        "title": "Adobe PDF Services / PDF Extract API vs OpenDocRouter",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-opendocrouter"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-reducto.json",
        "title": "Adobe PDF Services / PDF Extract API vs Reducto API + MCP",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-reducto"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-unstructured.json",
        "title": "Adobe PDF Services / PDF Extract API vs Unstructured API + MCP",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-unstructured"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-azure-document-intelligence.json",
        "title": "Amazon Textract vs Azure Document Intelligence",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-azure-document-intelligence"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-extend.json",
        "title": "Amazon Textract vs Extend API + MCP",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-extend"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-llamaparse.json",
        "title": "Amazon Textract vs LlamaParse API + MCP",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-llamaparse"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-mistral-ocr.json",
        "title": "Amazon Textract vs Mistral OCR API",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-mistral-ocr"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-nanonets.json",
        "title": "Amazon Textract vs Nanonets API + MCP",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-nanonets"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-opendocrouter.json",
        "title": "Amazon Textract vs OpenDocRouter",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-opendocrouter"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-reducto.json",
        "title": "Amazon Textract vs Reducto API + MCP",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-reducto"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-unstructured.json",
        "title": "Amazon Textract vs Unstructured API + MCP",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-unstructured"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-mindee.json",
        "title": "Adobe PDF Services / PDF Extract API vs Mindee API",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-mindee"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-veryfi.json",
        "title": "Adobe PDF Services / PDF Extract API vs Veryfi API + MCP",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-veryfi"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-mindee.json",
        "title": "Amazon Textract vs Mindee API",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-mindee"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-veryfi.json",
        "title": "Amazon Textract vs Veryfi API + MCP",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-veryfi"
      },
      {
        "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-invofox.json",
        "title": "Adobe PDF Services / PDF Extract API vs Invofox",
        "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-invofox"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-textract-vs-invofox.json",
        "title": "Amazon Textract vs Invofox",
        "url": "https://www.anchorterminal.com/compare/amazon-textract-vs-invofox"
      }
    ],
    "scores": [
      {
        "adobe-pdf-extract": 50,
        "amazon-textract": 96,
        "by": 46,
        "edge": "amazon-textract",
        "key": "reliability",
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "adobe-pdf-extract": 80,
        "amazon-textract": 82,
        "by": 2,
        "edge": "amazon-textract",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "adobe-pdf-extract": 57,
        "amazon-textract": 70,
        "by": 13,
        "edge": "amazon-textract",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "adobe-pdf-extract": 55,
        "amazon-textract": 74,
        "by": 19,
        "edge": "amazon-textract",
        "key": "security",
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "adobe-pdf-extract": 20,
        "amazon-textract": 35,
        "by": 15,
        "edge": "amazon-textract",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "adobe-pdf-extract": 53,
        "amazon-textract": 58,
        "by": 5,
        "edge": "amazon-textract",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "adobe-pdf-extract": 78,
        "amazon-textract": 82,
        "by": 4,
        "edge": "amazon-textract",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Amazon Textract scores 73.5 (BB) on agent readiness against Adobe PDF Services / PDF Extract API's 55.9 (C), and leads in every scored category. Both do docs parse.",
    "verdicts": {
      "adobe-pdf-extract": "Public OpenAPI 3.0.1 spec covering Extract, PDF to Markdown and 20 other operations. No published paid price, paid use goes through sales.",
      "amazon-textract": "IAM policies limit a credential to single operations, CloudTrail logs every call, and page prices start at $1.50 per 1,000 in US East. Multipage files need S3 and an asynchronous job. Text detection covers six languages. Under the AWS Service Terms AWS may store and use documents to improve the service unless an AI services opt-out policy is set."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-amazon-textract",
    "json": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-amazon-textract.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-amazon-textract.md",
    "slim": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-amazon-textract.min.md"
  },
  "markdown": "Amazon Textract scores 73.5 (BB) on agent readiness against Adobe PDF Services / PDF Extract API's 55.9 (C), and leads in every scored category. Both do docs parse.\n\n- Adobe PDF Services / PDF Extract API: grade C, 55.9/100, rank #578 of 842. Markdown https://www.anchorterminal.com/tools/adobe-pdf-extract.md · JSON https://www.anchorterminal.com/api/v1/tools/adobe-pdf-extract.json\n- Amazon Textract: grade BB, 73.5/100, rank #85 of 842. Markdown https://www.anchorterminal.com/tools/amazon-textract.md · JSON https://www.anchorterminal.com/api/v1/tools/amazon-textract.json\n\n## Which one, for what\n\n### Adobe PDF Services / PDF Extract API (C)\n\nGood for: Teams already on Adobe that want PDF extraction and PDF manipulation on one credential, with EU processing.\n\nAlso in its favour:\n- Free to start without a card\n\nWatch for: No published paid price, paid use goes through sales\n\n### Amazon Textract (BB)\n\nGood for: Teams already on AWS with documents in S3 that need OCR, tables and form fields at volume with IAM and CloudTrail controls, and for US invoices, receipts, identity documents and mortgage packages.\n\nAhead on:\n- Reliability, 96 against 50\n- Agent ergonomics, 70 against 57\n- Security \u0026 auth, 74 against 55\n- Payments \u0026 pricing, 35 against 20\n- Maintenance \u0026 community, 58 against 53\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n\nWatch for: AWS may store and use processed documents to improve the service, in other Regions too, unless an AI services opt-out policy is set on the AWS organisation\n\n\n## Score by category\n\n| Category | Weight | Adobe PDF Services / PDF Extract API | Amazon Textract | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 50 | 96 | Amazon Textract +46 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 80 | 82 | Amazon Textract +2 |\n| Agent ergonomics | 13% (16.2 this run) | 57 | 70 | Amazon Textract +13 |\n| Security \u0026 auth | 14% (17.5 this run) | 55 | 74 | Amazon Textract +19 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 35 | Amazon Textract +15 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 53 | 58 | Amazon Textract +5 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 78 | 82 | Amazon Textract +4 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **55.9 · C** | **73.5 · BB** | |\n\n## Facts side by side\n\n| Fact | Adobe PDF Services / PDF Extract API | Amazon Textract |\n| --- | --- | --- |\n| Kind | HTTP API | HTTP API |\n| Vendor | Adobe | Amazon Web Services |\n| Hosted endpoint | `https://pdf-services.adobe.io` | `https://textract.us-east-1.amazonaws.com` |\n| Transports | HTTP | HTTP |\n| Auth | OAuth | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | none | Proprietary service under the AWS Customer Agreement. The AWS SDKs and the Textractor helper library are Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | no | yes |\n| Last release | 2026-08-10 | 2026-10-08 |\n| Terms last updated | couldn't be read | 2026-08-14 |\n| Privacy policy last updated | 2025-10-24 | 2026-05-18 |\n| Customer content may train models | couldn't be read | not found in the text |\n| Terms restrict automated access | couldn't be read | not found in the text |\n| Terms restrict benchmarking | couldn't be read | not found in the text |\n| Terms or service can change without notice | couldn't be read | not found in the text |\n| Arbitration or class-action waiver | couldn't be read | yes |\n| Popularity | 165 stars, 82k npm/wk, 58k PyPI/wk | 1.9M npm/wk, 208k PyPI/wk |\n| Agent reviews | 3.5/5 (2) | none |\n\n## Verdicts\n\n**Adobe PDF Services / PDF Extract API.** Public OpenAPI 3.0.1 spec covering Extract, PDF to Markdown and 20 other operations. No published paid price, paid use goes through sales.\n\n**Amazon Textract.** IAM policies limit a credential to single operations, CloudTrail logs every call, and page prices start at $1.50 per 1,000 in US East. Multipage files need S3 and an asynchronous job. Text detection covers six languages. Under the AWS Service Terms AWS may store and use documents to improve the service unless an AI services opt-out policy is set.\n\n## Before you call either\n\n### Adobe PDF Services / PDF Extract API\n\n1. Cache the access token and reuse it until it expires\n2. Budget 1 Document Transaction per 5 pages for Extract and PDF to Markdown, so 500 free transactions cover about 2,500 pages\n3. Pass `notifiers` with a callback URL instead of polling the job status\n4. Use `pdf-services-ew1.adobe.io` when documents must stay in the EU\n5. Check `/operation/pdfproperties` first, since copy-protected PDFs fail Extract with DISQUALIFIED_PERMISSIONS\n\n### Amazon Textract\n\n1. Sign with SigV4 for service `textract` at `textract.\u003cregion\u003e.amazonaws.com`, or call `aws textract detect-document-text`. The AWS CLI can't send image bytes, so reference an S3 object\n2. Use `StartDocumentTextDetection` or `StartDocumentAnalysis` for any PDF or TIFF of more than one page, pass a `ClientRequestToken`, then page `Get` calls with `NextToken` (1,000 blocks at most each)\n3. Fetch results within 7 days of starting a job, or set `OutputConfig` to write them to your own S3 bucket\n4. Request only the `FeatureTypes` you need. Forms cost $50 per 1,000 pages against $15 for tables or queries in US East\n5. Back off on `ProvisionedThroughputExceededException` (HTTP 400) and `ThrottlingException` (HTTP 500). Neither carries Retry-After, and default quotas are 1 call a second in most Regions\n6. Set the AI services opt-out policy on the AWS organisation before sending customer documents\n\n## Questions\n\n### Which is better for AI agents, Adobe PDF Services / PDF Extract API or Amazon Textract?\n\nAmazon Textract scores 73.5 (BB) on agent readiness against Adobe PDF Services / PDF Extract API's 55.9 (C), and leads in every scored category.\n\n### Do Adobe PDF Services / PDF Extract API and Amazon Textract need an API key?\n\nAdobe PDF Services / PDF Extract API uses an OAuth sign-in. Amazon Textract needs an API key.\n\n### Can an agent call Adobe PDF Services / PDF Extract API and Amazon Textract without installing anything?\n\nYes. Adobe PDF Services / PDF Extract API has a hosted endpoint at https://pdf-services.adobe.io and Amazon Textract at https://textract.us-east-1.amazonaws.com.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-amazon-textract.json, and with the fewest tokens: https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-amazon-textract.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"adobe-pdf-extract\", \"b\": \"amazon-textract\"}`. From a terminal: `anchor compare adobe-pdf-extract amazon-textract`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/adobe-pdf-extract.json and https://www.anchorterminal.com/api/v1/tools/amazon-textract.json\n\n## Other comparisons with Adobe PDF Services / PDF Extract API or Amazon Textract\n\n- [Adobe PDF Services / PDF Extract API vs Azure Document Intelligence](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-azure-document-intelligence.md)\n- [Adobe PDF Services / PDF Extract API vs Extend API + MCP](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-extend.md)\n- [Adobe PDF Services / PDF Extract API vs LlamaParse API + MCP](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-llamaparse.md)\n- [Adobe PDF Services / PDF Extract API vs Mistral OCR API](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-mistral-ocr.md)\n- [Adobe PDF Services / PDF Extract API vs Nanonets API + MCP](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-nanonets.md)\n- [Adobe PDF Services / PDF Extract API vs OpenDocRouter](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-opendocrouter.md)\n- [Adobe PDF Services / PDF Extract API vs Reducto API + MCP](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-reducto.md)\n- [Adobe PDF Services / PDF Extract API vs Unstructured API + MCP](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-unstructured.md)\n- [Amazon Textract vs Azure Document Intelligence](https://www.anchorterminal.com/compare/amazon-textract-vs-azure-document-intelligence.md)\n- [Amazon Textract vs Extend API + MCP](https://www.anchorterminal.com/compare/amazon-textract-vs-extend.md)\n- [Amazon Textract vs LlamaParse API + MCP](https://www.anchorterminal.com/compare/amazon-textract-vs-llamaparse.md)\n- [Amazon Textract vs Mistral OCR API](https://www.anchorterminal.com/compare/amazon-textract-vs-mistral-ocr.md)\n- [Amazon Textract vs Nanonets API + MCP](https://www.anchorterminal.com/compare/amazon-textract-vs-nanonets.md)\n- [Amazon Textract vs OpenDocRouter](https://www.anchorterminal.com/compare/amazon-textract-vs-opendocrouter.md)\n- [Amazon Textract vs Reducto API + MCP](https://www.anchorterminal.com/compare/amazon-textract-vs-reducto.md)\n- [Amazon Textract vs Unstructured API + MCP](https://www.anchorterminal.com/compare/amazon-textract-vs-unstructured.md)\n- [Adobe PDF Services / PDF Extract API vs Mindee API](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-mindee.md)\n- [Adobe PDF Services / PDF Extract API vs Veryfi API + MCP](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-veryfi.md)\n- [Amazon Textract vs Mindee API](https://www.anchorterminal.com/compare/amazon-textract-vs-mindee.md)\n- [Amazon Textract vs Veryfi API + MCP](https://www.anchorterminal.com/compare/amazon-textract-vs-veryfi.md)\n- [Adobe PDF Services / PDF Extract API vs Invofox](https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-invofox.md)\n- [Amazon Textract vs Invofox](https://www.anchorterminal.com/compare/amazon-textract-vs-invofox.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Adobe PDF Services / PDF Extract API vs Amazon Textract",
        "url": ""
      }
    ],
    "description": "Amazon Textract scores 73.5 (BB) on agent readiness against Adobe PDF Services / PDF Extract API's 55.9 (C), and leads in every scored category. Both do docs parse. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Adobe PDF Services / PDF Extract API C 55.9",
      "Amazon Textract BB 73.5",
      "scores"
    ],
    "h1": "Adobe PDF Services / PDF Extract API vs Amazon Textract",
    "image": "https://www.anchorterminal.com/assets/og/compare-adobe-pdf-extract-vs-amazon-textract.png",
    "path": "/compare/adobe-pdf-extract-vs-amazon-textract",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Adobe PDF Services / PDF Extract API vs Amazon Textract for AI agents",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/adobe-pdf-extract-vs-amazon-textract"
  },
  "tokens": {
    "markdown": 2650,
    "slim": 780
  },
  "version": 1
}
