{
  "dataset": "free-llm-api-hub",
  "version": "2.9.0",
  "generated": "2026-08-14",
  "license": "MIT",
  "homepage": "https://freellmapihub.com/",
  "docs": "https://freellmapihub.com/api/",
  "description": "The editorial top 20 — hand-ranked free LLM APIs with the \"why\" for each pick, in rank order. Every pick carries its full verified profile.",
  "updated": "2026-08-14",
  "count": 20,
  "picks": [
    {
      "rank": 1,
      "why": "The rare package that is genuinely free forever (perpetual, not a trial) with a frontier-class model, OCR and real-time speech behind one no-card OpenAI-compatible key.",
      "tag": "Editor's pick",
      "slug": "typhoon",
      "name": "Typhoon (SCB 10X)",
      "category": "ongoing",
      "free_type": "perpetual",
      "free_tier": "Free to use research showcase API — all Typhoon models at $0",
      "rate_limits": "typhoon-asr-realtime: 100 reqs/minute (documented on the ASR page); LLM/OCR models: no published quotas (beta service)",
      "notes": "By SCB 10X, the venture arm of Siam Commercial Bank, focused on Thai-language models. The free catalog spans LLMs (typhoon-v2.5-30b-a3b-instruct), OCR (typhoon-ocr family) and realtime Thai ASR (typhoon-asr-realtime, typhoon-isan-asr-realtime) — the hosted API is OpenAI-compatible, including audio transcriptions. Beta, provided as-is with no formal support; usage data is collected to improve the model; SCB claims no rights in outputs. Sign up for a free API key at opentyphoon.ai.",
      "best_for": "Thai-language LLMs, OCR and speech via an OpenAI-compatible research API",
      "modalities": [
        "text",
        "audio"
      ],
      "models_free": [
        "typhoon-asr-realtime",
        "typhoon-isan-asr-realtime",
        "typhoon-ocr",
        "typhoon-ocr-preview",
        "typhoon-ocr-v1.5",
        "typhoon-v2.5-30b-a3b-instruct"
      ],
      "expires": null,
      "docs_url": "https://docs.opentyphoon.ai/en/faq/",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://api.opentyphoon.ai/v1",
      "env_key": "TYPHOON_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 2,
      "why": "The best free daily volume for a real side project — 10k requests a day with no card, and the quota renews.",
      "tag": "Best free quota",
      "slug": "cloudflare-workers-ai",
      "name": "Cloudflare Workers AI",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "10,000 Neurons/day, all account plans",
      "rate_limits": "30+ models: LLMs (Llama, Mistral, DeepSeek, Qwen...), embeddings, image, audio",
      "notes": "Resets daily at 00:00 UTC; overage on a Workers Paid plan bills at $0.011/1,000 Neurons. A few models (e.g. Kimi K2.6/K2.7-code, GLM-5.2) now require a Workers Paid plan",
      "best_for": "Highest free daily volume for a real side project",
      "modalities": [
        "text",
        "embeddings",
        "image",
        "audio"
      ],
      "models_free": [
        "@cf/meta/llama-3.3-70b-instruct-fp8-fast",
        "@cf/meta/llama-3.1-8b-instruct",
        "@cf/mistralai/mistral-small-3.1-24b-instruct",
        "@cf/qwen/qwen2.5-coder-32b-instruct",
        "@cf/deepseek-ai/deepseek-r1-distill-qwen-32b",
        "@cf/google/gemma-3-12b-it"
      ],
      "expires": null,
      "docs_url": "https://developers.cloudflare.com/workers-ai/platform/pricing/",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai/v1",
      "env_key": "CLOUDFLARE_API_TOKEN",
      "verified": true,
      "last_verified": "2026-08-02"
    },
    {
      "rank": 3,
      "why": "The only frontier-class model with a genuine, renewing free tier — the default when you want the smartest free model.",
      "tag": "Frontier quality",
      "slug": "google-gemini",
      "name": "Google Gemini API (AI Studio)",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "Gemini 2.5 Flash, 2.5 Flash-Lite, 2.5 Pro (limited), embeddings, TTS models",
      "rate_limits": "Varies by model: 5-30 RPM and 15-1,000 RPD (e.g. 2.5 Pro: 5 RPM/100 RPD; 2.5 Flash: 10/250; 2.5 Flash-Lite: 15/1,000; embeddings: 100 RPD; TTS: 15 RPD)",
      "notes": "Free-tier prompts/outputs may be used by Google to improve its products outside the UK/CH/EEA/EU. Since the 2026-03-23 terms, only Paid Services may serve API clients to end users in the EEA/CH/UK",
      "best_for": "The only frontier-class model with a genuine free tier",
      "modalities": [
        "text",
        "vision",
        "embeddings",
        "audio"
      ],
      "models_free": [
        "gemini-2.5-flash",
        "gemini-2.5-flash-lite",
        "gemini-2.5-pro"
      ],
      "expires": null,
      "docs_url": "https://ai.google.dev/gemini-api/docs/rate-limits",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://generativelanguage.googleapis.com/v1beta/openai/",
      "env_key": "GEMINI_API_KEY",
      "verified": true,
      "last_verified": "2026-08-02"
    },
    {
      "rank": 4,
      "why": "Cloud-hosted open models with the exact Ollama workflow and no card at signup — the easiest on-ramp for open-weight models.",
      "tag": "Best on-ramp",
      "slug": "ollama-cloud",
      "name": "Ollama Cloud",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "$0 Free plan: access to cloud-hosted open models (Qwen, GPT-OSS, DeepSeek, etc.) via API",
      "rate_limits": "Session limits reset every 5 hours and weekly limits every 7 days; 1 concurrent cloud model on the free plan (exact token caps not published)",
      "notes": "First-party — Ollama hosts the cloud models. Requires an ollama.com account + API key (`ollama signin`); the free plan is for light usage, Pro ($20/mo) raises limits. No card required.",
      "best_for": "Cloud-hosted open models with the Ollama workflow, no card",
      "modalities": [
        "text",
        "vision"
      ],
      "models_free": [
        "deepseek-v4-flash:0731",
        "deepseek-v4-flash:preview",
        "deepseek-v4-pro:preview",
        "gemma4:31b",
        "glm-5.1",
        "glm-5.2",
        "gpt-oss:120b",
        "gpt-oss:20b"
      ],
      "expires": null,
      "docs_url": "https://docs.ollama.com/cloud",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://ollama.com/v1",
      "env_key": "OLLAMA_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 5,
      "why": "One no-card key, every model: a wide rotating sample of :free models behind a single OpenAI-compatible endpoint.",
      "tag": "Best gateway",
      "slug": "openrouter",
      "name": "OpenRouter",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "A rotating set of models with a :free suffix (~14 today; count fluctuates), single API across many providers",
      "rate_limits": "20 req/min; 50 req/day under 10 credits purchased lifetime, 1000 req/day once 10+ credits purchased (one-time, not a subscription)",
      "notes": "ToS (Jul 2026) prohibits reselling API access or building a competing service — platform-wide, not just the free models; per-model terms still apply",
      "best_for": "Widest model selection behind one no-card key",
      "modalities": [
        "text",
        "vision"
      ],
      "models_free": [
        "cohere/north-mini-code:free",
        "google/gemma-4-26b-a4b-it:free",
        "google/gemma-4-31b-it:free",
        "liquid/lfm-2.5-2.6b:free",
        "nvidia/nemotron-3-nano-30b-a3b:free",
        "nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free",
        "openai/gpt-oss-20b:free",
        "poolside/laguna-s-2.1:free"
      ],
      "expires": null,
      "docs_url": "https://openrouter.ai/docs/api-reference/limits",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://openrouter.ai/api/v1",
      "env_key": "OPENROUTER_API_KEY",
      "verified": true,
      "last_verified": "2026-08-02"
    },
    {
      "rank": 6,
      "why": "Anonymous, no-signup and permanently free — the fallback that works when everything else asks for an account.",
      "tag": "No account needed",
      "slug": "ai-horde",
      "name": "AI Horde",
      "category": "ongoing",
      "free_type": "perpetual",
      "free_tier": "Free crowdsourced text & image generation; anonymous API key '0000000000' (no registration), or register to earn kudos for priority",
      "rate_limits": "Queue-based priority via kudos (no fixed quota); anonymous requests get lowest priority under load",
      "notes": "Community-powered volunteer network — model availability and speed vary with worker supply, so it is not a fixed-SLA service. No card, no phone. Kudos never expire and cannot be sold.",
      "best_for": "Anonymous, no-signup free generation (variable speed)",
      "modalities": [
        "text",
        "image"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://aihorde.net/",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": null,
      "openai_compatible": false,
      "openai_base_url": null,
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 7,
      "why": "The most generous free document-parsing tier for RAG builders (~10k pages a month), renewing every month.",
      "tag": "Best for RAG",
      "slug": "llamaparse",
      "name": "LlamaParse (LlamaCloud)",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "Free plan: 10,000 credits/month (~10,000 pages in balanced parse mode at 1 credit/page)",
      "rate_limits": "5 concurrent jobs; 1 project; 5 indexes on the free plan",
      "notes": "Credit-based (1,000 credits = $1.25); premium parse modes consume more credits per page. No card required. Document parsing for RAG (LlamaIndex).",
      "best_for": "Free document parsing for RAG (~10k pages/mo)",
      "modalities": [
        "ocr",
        "text"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://www.llamaindex.ai/pricing",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": false,
      "openai_base_url": null,
      "env_key": "LLAMACLOUD_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 8,
      "why": "Classic, reliable free OCR at 25k conversions a month with no card — the boring tool that just works.",
      "tag": "Reliable utility",
      "slug": "ocr-space",
      "name": "OCR.space",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "25,000 conversions/month (Engine 1 & 2) plus 2,500 Engine 3 conversions/month; max 1 MB file, PDFs up to 3 pages",
      "rate_limits": "500 requests per day per IP address",
      "notes": "Free searchable-PDF output carries a watermark (raw text extraction is unrestricted); the free key needs only an email, no card. Commercial use permitted.",
      "best_for": "Classic free OCR (25k conversions/mo, no card)",
      "modalities": [
        "ocr",
        "text"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://ocr.space/OCRAPI",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": false,
      "openai_base_url": null,
      "env_key": "OCRSPACE_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 9,
      "why": "Free commercial TTS at 50k characters a month — the pick when the output actually ships.",
      "tag": "Best for TTS",
      "slug": "speechify",
      "name": "Speechify API",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "50,000 characters/month TTS (hard cap) + 60 min/month voice agents",
      "rate_limits": "3 concurrent calls; hard cap pauses at the limit (no overages)",
      "notes": "Commercial use is allowed on the free tier. No credit card required. Hard monthly cap that pauses at the limit.",
      "best_for": "Free commercial TTS (50k chars/mo)",
      "modalities": [
        "audio"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://speechify.ai/pricing",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": false,
      "openai_base_url": null,
      "env_key": "SPEECHIFY_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 10,
      "why": "Free embeddings and rerank behind an OpenAI-compatible endpoint — the shortest path to a free RAG stack.",
      "tag": "Best for embeddings",
      "slug": "jina-ai",
      "name": "Jina AI",
      "category": "trial",
      "free_type": "trial-credit",
      "free_tier": "10M free tokens (one-time) across all models — embeddings, rerankers, classifier; plus a keyless Reader (r.jina.ai) for basic use",
      "rate_limits": "Free key: 100 RPM / 100k TPM for embeddings & reranker (2 concurrent); keyless Reader 20 RPM",
      "notes": "The 10M-token balance is a one-time grant that does not replenish; the keyless Reader is genuinely ongoing. Hosted API is commercial-OK and data is not used for training. No card required.",
      "best_for": "Free embeddings & rerank behind an OpenAI-compatible endpoint",
      "modalities": [
        "text",
        "image",
        "embeddings",
        "rerank"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://jina.ai/embeddings/",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://api.jina.ai/v1",
      "env_key": "JINA_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 11,
      "why": "The largest speech credit in the list — $200 for Nova STT and Aura TTS with no card and no expiration, plus an opt-out from model training.",
      "tag": "Best STT credit",
      "slug": "deepgram",
      "name": "Deepgram",
      "category": "trial",
      "free_type": "trial-credit",
      "free_tier": "$200 free credit on signup (no card, no expiration) — Nova speech-to-text and Aura text-to-speech at pay-as-you-go rates",
      "rate_limits": "STT pre-recorded up to 50 concurrent; STT streaming up to 150; TTS REST up to 15; TTS streaming up to 45",
      "notes": "No card and no expiration on the credit. Data catch: the Model Improvement Program is opt-OUT — send mip_opt_out=true per request to keep your data out of training. Native REST/WebSocket API, not OpenAI-compatible.",
      "best_for": "Biggest free credit here, for speech (STT/TTS)",
      "modalities": [
        "audio"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://deepgram.com/pricing",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": false,
      "openai_base_url": null,
      "env_key": "DEEPGRAM_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 12,
      "why": "Purpose-built LPU inference makes this the lowest-latency free tier for open-weight models — the only catch is phone verification at signup.",
      "tag": "Fastest inference",
      "slug": "groq",
      "name": "Groq",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "Open-weight models (Llama, Qwen, GPT-OSS) plus Whisper, no credit card required",
      "rate_limits": "e.g. llama-3.1-8b-instant: 30 RPM/14.4K RPD/6K TPM/500K TPD; llama-3.3-70b-versatile: 30 RPM/1K RPD/12K TPM/100K TPD; qwen/qwen3.6-27b: 30 RPM/1K RPD/8K TPM/200K TPD; similar for GPT-OSS and Whisper models",
      "notes": "Limits apply at the organization level, not per API key. Phone verification required at signup",
      "best_for": "Lowest-latency inference for open-weight models",
      "modalities": [
        "text",
        "audio"
      ],
      "models_free": [
        "llama-3.3-70b-versatile",
        "llama-3.1-8b-instant",
        "openai/gpt-oss-120b",
        "openai/gpt-oss-20b",
        "whisper-large-v3-turbo",
        "whisper-large-v3"
      ],
      "expires": null,
      "docs_url": "https://console.groq.com/docs/rate-limits",
      "phone_required": true,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://api.groq.com/openai/v1",
      "env_key": "GROQ_API_KEY",
      "verified": true,
      "last_verified": "2026-08-02"
    },
    {
      "rank": 13,
      "why": "A rate-limited free tier across all models with a commercial license and no card — fast open-model inference without a trial clock counting down.",
      "tag": "Fast, no card",
      "slug": "sambanova",
      "name": "SambaNova Cloud",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "Rate-limited free tier (applies when no payment method is linked) across all models",
      "rate_limits": "Free Tier: 20 RPM / 20 RPD / 200,000 TPD across all models; Developer Tier (card required): 60-240 RPM depending on model",
      "notes": "Free Tier applies when no payment method is linked to the account; SambaCloud ToS grants a commercial license (no evaluation-only clause). The previously-listed \"$5 / 3 months\" trial could not be re-confirmed on official pages (2026-07-30).",
      "best_for": "Fast open models on a rate-limited free tier",
      "modalities": [
        "text"
      ],
      "models_free": [
        "DeepSeek-V3.1",
        "Meta-Llama-3.3-70B-Instruct",
        "gpt-oss-120b"
      ],
      "expires": null,
      "docs_url": "https://cloud.sambanova.ai/plans",
      "phone_required": null,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://api.sambanova.ai/v1",
      "env_key": "SAMBANOVA_API_KEY",
      "verified": true,
      "last_verified": "2026-08-13"
    },
    {
      "rank": 14,
      "why": "$100/month of serverless inference credits covering frontier open models behind one OpenAI-compatible key, with Weave tracing on top.",
      "tag": "Best monthly frontier credits",
      "slug": "wandb-inference",
      "name": "W&B Inference",
      "category": "ongoing",
      "free_type": "recurring-credit",
      "free_tier": "$100/month of Serverless Inference credits on the Free plan (default spending cap; offer for a limited time)",
      "rate_limits": "No RPM/TPM published; default cap of $100/month on the Free tier; concurrency limits per project/user",
      "notes": "Serverless Inference credits come with Free, Pro and Academic plans for a limited time; when credits run out, Free accounts must activate pay-as-you-go on the Billing tab or upgrade. OpenAI-compatible endpoint at api.inference.wandb.ai/v1 with any W&B API key.",
      "best_for": "Open-weight frontier models (DeepSeek, Llama, Qwen, GLM, GPT-OSS) behind one unified API with Weave tracing",
      "modalities": [
        "text",
        "vision"
      ],
      "models_free": [
        "deepseek-ai/DeepSeek-V4-Flash",
        "meta-llama/Llama-3.3-70B-Instruct",
        "meta-llama/Llama-3.1-8B-Instruct",
        "MiniMaxAI/MiniMax-M3",
        "openai/gpt-oss-120b",
        "openai/gpt-oss-20b",
        "zai-org/GLM-5.2",
        "google/gemma-4-31B-it",
        "nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B"
      ],
      "expires": null,
      "docs_url": "https://docs.wandb.ai/inference/usage-limits",
      "phone_required": null,
      "card_required": false,
      "commercial_ok": true,
      "openai_compatible": true,
      "openai_base_url": "https://api.inference.wandb.ai/v1",
      "env_key": "WANDB_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 15,
      "why": "200M one-time tokens across best-in-class embedding and rerank models — the biggest free embedding balance here, no card required.",
      "tag": "Best embedding allotment",
      "slug": "voyage-ai",
      "name": "Voyage AI",
      "category": "trial",
      "free_type": "trial-credit",
      "free_tier": "200M free tokens on current embedding models (voyage-4-large, voyage-4, voyage-4-lite, voyage-context-4, voyage-code-4) and on rerankers (rerank-2.5 family); voyage-multimodal-3.5 and voyage-multimodal-3 get 200M text tokens + 150B pixels — a large one-time complimentary allotment per model",
      "rate_limits": "Standard per-model RPM/TPM limits apply; the free-token allotment is the practical ceiling",
      "notes": "The allotment is a one-time complimentary balance per model, not a renewing monthly quota. No credit card required to claim. Owned by MongoDB — a first-party model provider, not a proxy.",
      "best_for": "Best-in-class embeddings and reranking with a large free allotment",
      "modalities": [
        "text",
        "image",
        "embeddings",
        "rerank"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://docs.voyageai.com/docs/pricing",
      "phone_required": null,
      "card_required": false,
      "commercial_ok": null,
      "openai_compatible": false,
      "env_key": "VOYAGE_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14",
      "added": "2026-07-31"
    },
    {
      "rank": 16,
      "why": "5M embedding tokens and 500 rerank requests a month on a managed Starter plan — a real, renewing free tier for RAG.",
      "tag": "Best managed embeddings",
      "slug": "pinecone-inference",
      "name": "Pinecone Inference",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "Starter (free) plan: 5M tokens/mo for embedding models (llama-text-embed-v2, multilingual-e5-large) and 500 requests/mo for the bge-reranker-v2-m3 rerank model",
      "rate_limits": "5M embedding tokens/mo; 500 rerank requests/mo on Starter",
      "notes": "Free rerank limited to bge-reranker-v2-m3; overages are pay-as-you-go. Not OpenAI-compatible. No card required.",
      "best_for": "Free embeddings + rerank for RAG (5M tokens/mo)",
      "modalities": [
        "embeddings",
        "rerank"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://www.pinecone.io/pricing/",
      "phone_required": null,
      "card_required": false,
      "commercial_ok": null,
      "openai_compatible": false,
      "openai_base_url": null,
      "env_key": "PINECONE_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 17,
      "why": "Hosted Flux and Stable Diffusion behind a simple GET URL, anonymous to start — the easiest way to generate free images programmatically.",
      "tag": "Free images, no signup",
      "slug": "pollinations",
      "name": "Pollinations.ai",
      "category": "ongoing",
      "free_type": "perpetual",
      "free_tier": "Free hosted image models (Flux, Turbo, Stable Diffusion) via a simple GET URL; also text and audio. No signup required to start",
      "rate_limits": "Anonymous ~1 request / 15s; free registration (Seed tier) ~1 request / 5s",
      "notes": "Anonymous free images may carry a watermark (since 2025); free registration removes it via the nologo parameter. No card, no phone. Commercial use is not explicitly guaranteed in the docs.",
      "best_for": "No-signup free image generation (Flux/SD)",
      "modalities": [
        "image",
        "text",
        "audio"
      ],
      "models_free": [
        "openai-fast"
      ],
      "expires": null,
      "docs_url": "https://github.com/pollinations/pollinations/blob/master/APIDOCS.md",
      "phone_required": false,
      "card_required": false,
      "commercial_ok": null,
      "openai_compatible": false,
      "openai_base_url": null,
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 18,
      "why": "A $5/month recurring credit for the Moondream vision model — caption, VQA and detection behind an OpenAI-compatible endpoint.",
      "tag": "Best tiny vision",
      "slug": "moondream",
      "name": "Moondream Cloud",
      "category": "ongoing",
      "free_type": "recurring-credit",
      "free_tier": "$5/month usage credits in every workspace (Free plan) for the Moondream vision model — caption, query (VQA), detect, point",
      "rate_limits": "Bounded by the $5/month credit",
      "notes": "Recurring $5/month credit, no credit card required (stated on the pricing page/blog). Commercial terms not specified. OpenAI-compatible endpoint.",
      "best_for": "Tiny vision model, OpenAI-compatible, monthly free credit",
      "modalities": [
        "vision"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://moondream.ai/pricing",
      "phone_required": null,
      "card_required": false,
      "commercial_ok": null,
      "openai_compatible": true,
      "openai_base_url": "https://api.moondream.ai/v1",
      "env_key": "MOONDREAM_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14"
    },
    {
      "rank": 19,
      "why": "15,000 pages a month of document parsing and OCR with no card — the free path from messy PDFs to LLM-ready structured data.",
      "tag": "Best doc pipeline",
      "slug": "unstructured",
      "name": "Unstructured",
      "category": "ongoing",
      "free_type": "renewing-quota",
      "free_tier": "15,000 pages/month, resets monthly — document parsing/OCR across 50+ file types (layout, tables, generative OCR enrichment)",
      "rate_limits": "15,000 pages/month",
      "notes": "No credit card required. Purpose-built to turn documents into clean, structured input for RAG/LLM pipelines. Commercial terms not stated on the pricing page.",
      "best_for": "Turning messy documents into LLM-ready structured data, free every month",
      "modalities": [
        "ocr"
      ],
      "models_free": null,
      "expires": null,
      "docs_url": "https://unstructured.io/pricing",
      "phone_required": null,
      "card_required": false,
      "commercial_ok": null,
      "openai_compatible": false,
      "env_key": "UNSTRUCTURED_API_KEY",
      "verified": true,
      "last_verified": "2026-08-14",
      "added": "2026-07-31"
    },
    {
      "rank": 20,
      "why": "1M free tokens and 60 minutes of Whisper on EU-hosted open models — the sovereignty-friendly free tier for European workloads.",
      "tag": "Best EU open models",
      "slug": "scaleway",
      "name": "Scaleway Generative APIs",
      "category": "trial",
      "free_type": "trial-credit",
      "free_tier": "1,000,000 tokens free + 60 min Whisper transcription; billing starts at token 1,000,001",
      "rate_limits": "Gemma, Llama, Mistral, Qwen",
      "notes": "European provider (France). Free allowance is a one-time token bucket, not time-limited. The 1M free tokens need no card; adding a card + passing KYC unlocks the official rate limits.",
      "best_for": "EU-hosted open models with 1M free tokens",
      "modalities": [
        "text",
        "audio"
      ],
      "models_free": [
        "llama-3.3-70b-instruct",
        "qwen3-235b-a22b-instruct-2507",
        "mistral-small-3.2-24b-instruct-2506",
        "deepseek-r1-distill-llama-70b",
        "gemma-3-27b-it",
        "qwen2.5-coder-32b-instruct"
      ],
      "expires": null,
      "docs_url": "https://www.scaleway.com/en/pricing/model-as-a-service/",
      "phone_required": null,
      "card_required": false,
      "commercial_ok": null,
      "openai_compatible": true,
      "openai_base_url": "https://api.scaleway.ai/v1",
      "env_key": "SCALEWAY_API_KEY",
      "verified": true,
      "last_verified": "2026-08-13"
    }
  ]
}
