schema_version: 1
title: Free hosted model inference providers and routers
as_of: "2026-08-21"
checked_at: "2026-08-21T23:37:26-05:00"
timezone: America/Chicago

snapshot_counts:
  current_offer_records: 76
  headline_ongoing_or_recurring_records: 38
  conditional_delivery_current_records: 7
  finite_trial_or_promotion_records: 31
  conditional_or_needs_verification_records: 20
  historical_or_excluded_records: 37
  note: Current offers include trials and conditional delivery models; filter free_kind and headline_ongoing_free rather than treating 76 as a count of permanent free tiers.

scope:
  completeness: Best-effort exhaustive public-web audit, not a guarantee that every regional, private-beta, dashboard-only, or newly launched service has been found.
  primary_subject: Hosted foundation-model inference, routers, and serverless model-deployment platforms.
  modalities: [text_generation, embeddings, reranking, image, audio, speech, safety, custom_model_deployments]
  included:
    - Remotely callable hosted model inference with an API.
    - Ongoing free tiers, recurring credits, experimental access, and model-specific zero-price routes.
    - Time-limited trials and promotions when clearly labeled as such.
  excluded_from_current:
    - Open weights that require self-hosting.
    - Web chat without a remotely callable inference API.
    - A free gateway whose upstream model inference is still billed, unless it also has named zero-price models.
    - Offers found only in stale or third-party pages without current first-party confirmation.
  verification_levels:
    live_catalog: An official models endpoint was queried without credentials; this does not prove an authenticated generation succeeded.
    live_call: A credential-free inference call succeeded; this still does not prove authenticated-account behavior or billing.
    official_current: Current first-party pricing, documentation, or product pages explicitly support the claim.
    dashboard_only: The provider says the limit or catalog is visible only after sign-in.
    conflicting: Current first-party sources disagree; both claims are retained.
  authenticated_generation_policy: No provider credential was created and no authenticated generation was billed for this audit. A live_catalog label validates published server metadata, not end-to-end callability.
  important_note: Free catalogs and quotas are volatile. Re-fetch each discovery URL before relying on this snapshot.

website_filtering:
  default_current_list: Read records only from providers.
  headline_ongoing_free:
    omitted_or_true: Normal ongoing, recurring, restricted-research, or active experimental offer.
    conditional: Show with a community-capacity, user-pays, data-use, or sustainability warning.
    false: One-time trial, application credit, or limited promotion; never label as a permanent free tier.
  candidate_list: Read conditional_or_needs_verification separately and never merge it into verified counts.
  unavailable_list: historical_or_excluded contains retired, paid-only, non-API, nonfunctional, and self-hosted-only results.
  recommended_badges: [free_kind, status, confidence, payment_method_required, geography, modality, verification]

discovery_coverage:
  interpretation: Ecosystem registries are candidate generators, not proof that each upstream provider has a direct free API.
  openrouter_provider_registry:
    checked_at: "2026-08-21"
    discovery_url: https://openrouter.ai/api/v1/providers
    live_provider_count: 104
    caveat: The live API includes a synthetic-looking fake-provider entry; it is retained for exact snapshot fidelity and is not treated as a real provider candidate.
    provider_slugs:
      - ai21
      - aion-labs
      - akashml
      - alibaba
      - amazon-bedrock
      - amazon-nova
      - ambient
      - anthropic
      - arcee-ai
      - atlas-cloud
      - avian
      - azure
      - baidu
      - baseten
      - black-forest-labs
      - cerebras
      - chutes
      - cirrascale
      - clarifai
      - claude-on-aws
      - cloudflare
      - cohere
      - coreweave
      - crucible
      - crusoe
      - darkbloom
      - databricks
      - decart
      - deepgram
      - deepinfra
      - deepseek
      - dekallm
      - digitalocean
      - fake-provider
      - featherless
      - fireworks
      - fish-audio
      - friendli
      - gmicloud
      - google-ai-studio
      - google-vertex
      - groq
      - heygen
      - inception
      - inceptron
      - inferact-vllm
      - inference-net
      - infermatic
      - inflection
      - io-net
      - ionstream
      - krea
      - liquid
      - makora
      - mancer
      - mara
      - meta
      - minimax
      - mistral
      - modal
      - modelrun
      - modular
      - moonshotai
      - morph
      - ncompass
      - nebius
      - nex-agi
      - nextbit
      - novita
      - nvidia
      - open-inference
      - openai
      - parasail
      - perceptron
      - perplexity
      - phala
      - poolside
      - quiver
      - recraft
      - reka
      - relace
      - runpod
      - runway
      - sail-research
      - sakana
      - sambanova
      - seed
      - siliconflow
      - sourceful
      - stealth
      - stepfun
      - streamlake
      - switchpoint
      - tencent
      - tenstorrent
      - thinkingmachines
      - together
      - upstage
      - venice
      - voyageai
      - wafer
      - xai
      - xiaomi
      - z-ai
  huggingface_inference_provider_registry:
    checked_at: "2026-08-21"
    discovery_url: https://huggingface.co/docs/inference-providers/main/en/index
    documented_integration_count: 19
    provider_ids: [cerebras, cohere, deepinfra, fal_ai, featherless_ai, fireworks, groq, hf_inference, hyperbolic, novita, nscale, ovhcloud, public_ai, replicate, sambanova, scaleway, together, wavespeedai, zai]
  other_systematic_sources:
    - NVIDIA Build/NIM catalog, developer terms, model cards, and partner announcements.
    - Vercel AI Gateway, OrcaRouter, Kilo, FastRouter, OpenCode Zen, and community-router catalogs.
    - Official pricing, billing, rate-limit, changelog, status, GitHub, press-release, X/Twitter, and news searches for every named candidate.
  coverage_note: Every positive record below has provider-specific evidence. High-probability negative results are retained under historical_or_excluded to prevent rediscovery from stale announcements.

free_kinds:
  ongoing_free: Recurring or indefinite no-cost API usage under a quota.
  rotating_zero_price: Named models cost zero while they remain in a changing catalog.
  recurring_credit: A monetary allowance that refreshes on a stated schedule.
  experimental: Free only while a preview or experiment remains active.
  development_only: Free for prototyping or evaluation, not production use.
  promotional: A temporary model promotion or expiring allowance.
  trial: A one-time new-account credit or time window.
  community_capacity: No-cost inference supplied by volunteer or community-operated capacity with no guaranteed throughput.
  user_pays: The developer pays nothing because each end user authenticates and consumes that user's own allowance.

providers:
  - id: openrouter
    name: OpenRouter
    website: https://openrouter.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://openrouter.ai/api/v1
      compatibility: [openai_chat_completions, openai_completions, openai_models]
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      requests_per_minute: 20
      requests_per_day_without_purchase: 50
      requests_per_day_after_10_usd_credit_purchase: 1000
      scope: Total across free-model requests; failed requests count.
    free_router:
      model_id: openrouter/free
      price_usd: 0
      behavior: Randomly selects a compatible model from the current free pool.
    models:
      snapshot_at: "2026-08-21T21:13:22-05:00"
      verification: live_catalog
      discovery_url: https://openrouter.ai/api/v1/models
      count_including_router: 22
      free_model_ids:
        - stealth/ox-alpha
        - dots-studio/dots-3-note-preview:free
        - liquid/lfm-2.5-2.6b:free
        - nvidia/nemotron-3.5-lightning:free
        - thinkingmachines/inkling-small:free
        - poolside/laguna-s-2.1:free
        - thinkingmachines/inkling:free
        - poolside/laguna-xs-2.1:free
        - cohere/north-mini-code:free
        - z-ai/glm-5.2:free
        - nvidia/nemotron-3.5-content-safety:free
        - nvidia/nemotron-3-ultra-550b-a55b:free
        - nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free
        - google/gemma-4-26b-a4b-it:free
        - google/gemma-4-31b-it:free
        - google/lyria-3-pro-preview
        - google/lyria-3-clip-preview
        - nvidia/nemotron-3-super-120b-a12b:free
        - openrouter/free
        - nvidia/nemotron-3-nano-30b-a3b:free
        - nvidia/nemotron-nano-12b-v2-vl:free
        - nvidia/nemotron-nano-9b-v2:free
      expiring_models:
        dots-studio/dots-3-note-preview:free: "2026-09-30"
        nvidia/nemotron-3-nano-30b-a3b:free: "2026-08-24"
        nvidia/nemotron-nano-12b-v2-vl:free: "2026-08-24"
        nvidia/nemotron-nano-9b-v2:free: "2026-08-24"
      caveats:
        - Lyria preview entries are zero-priced despite lacking a :free suffix.
        - Marketing counts and the live API count differ; this snapshot uses zero prompt and completion prices from the API.
    history:
      - {date: "2025-07-10", event: "OpenRouter described how it would sustain its free tier after two upstream providers moved paid-only."}
      - {date: "2026-02-23", event: "The openrouter/free router was announced."}
    sources:
      - {type: official_docs, title: OpenRouter FAQ, url: https://openrouter.ai/docs/faq}
      - {type: live_catalog, title: Models endpoint, url: https://openrouter.ai/api/v1/models}
      - {type: official_docs, title: Free Models Router, url: https://openrouter.ai/docs/guides/routing/routers/free-router}
      - {type: announcement, title: Updates to Our Free Tier, url: https://openrouter.ai/blog/announcements/updates-to-our-free-tier-sustaining-accessible-ai-for-everyone/, published_at: "2025-07-10"}

  - id: nous_portal
    name: Nous Research / Nous Portal
    website: https://nousresearch.com/
    portal_url: https://portal.nousresearch.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    requested_name_note: User-confirmed resolution of “NUS Research” is Nous Research at https://nousresearch.com/.
    api:
      base_url: https://inference-api.nousresearch.com/v1
      compatibility: [openai_chat_completions, openai_models]
      authentication: Nous Portal OAuth or bearer credential
    eligibility:
      account_required: true
      plan: Free — $0/month
      payment_method_required: not_explicitly_documented
    limits:
      observed_portal_requests_per_minute: 50
      observed_portal_tokens_per_minute: 500000
      caveat: One current portal surface exposes these numbers while another says only “Standard rate limits”; account limits remain authoritative.
    models:
      snapshot_at: "2026-08-21T21:13:22-05:00"
      verification: live_catalog
      discovery_url: https://inference-api.nousresearch.com/v1/models
      count: 7
      free_model_ids:
        - stealth/ox-alpha
        - poolside/laguna-s-2.1:free
        - poolside/laguna-xs-2.1:free
        - tencent/hy3:free
        - stepfun/step-3.7-flash:free
        - upstage/solar-pro4:free
        - meituan/longcat-2.0:free
    history:
      - date: "2025-03-12"
        event: Initial inference API launch with a waitlist and $5 signup credit.
        initial_models: [Hermes 3 Llama 70B, DeepHermes 3 8B Preview]
        current_status: superseded_by_zero_price_catalog
    sources:
      - {type: official_website, title: Nous Research, url: https://nousresearch.com/}
      - {type: official_product, title: Nous Portal, url: https://portal.nousresearch.com/}
      - {type: live_catalog, title: Models endpoint, url: https://inference-api.nousresearch.com/v1/models}
      - {type: announcement, title: Announcing the Nous Portal, url: https://forum.nousresearch.com/t/announcing-the-nous-portal/62, published_at: "2025-03-12"}
      - {type: official_repository_docs, title: Hermes Agent Nous Portal integration, url: https://github.com/NousResearch/hermes-agent/blob/main/website/docs/integrations/nous-portal.md}

  - id: orcarouter
    name: OrcaRouter
    operator: Continuum AI Corp.
    website: https://www.orcarouter.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://api.orcarouter.ai/v1
      compatibility: [openai_chat_completions, openai_responses, anthropic_style, gemini_style]
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
      free_models_may_require_claim: true
    limits:
      published_general_numeric_limits: false
      exhausted_behavior: HTTP 429 or free-quota-exhausted; the wallet is not charged by the free router.
      recent_announcement_snapshot:
        requests_per_minute: 10
        requests_per_day: 50
        after_20_usd_cumulative_workspace_payments: {requests_per_minute: 20, requests_per_day: 1000}
      caveat: The launch announcement, localized offer page, and live catalog differed during the 2026-08-21 audit; prefer the signed-in offer/account state.
    free_router:
      model_id: orcarouter/free
      documented_routes:
        - {model: deepseek/deepseek-v4-flash, advertised_free_calls: 50}
        - {model: deepseek/deepseek-v4-pro, advertised_free_calls: 10}
      reset_period: unpublished
    models:
      snapshot_at: "2026-08-21T21:13:22-05:00"
      verification: live_catalog
      discovery_url: https://api.orcarouter.ai/v1/models
      free_models:
        - {id: deepseek/deepseek-v4-flash-free, context_tokens: 1048576, max_output_tokens: 384000}
        - {id: deepseek/deepseek-v4-pro-free, context_tokens: 1048576, max_output_tokens: 384000}
        - id: qwen/qwen3.8-27b-free
          context_tokens_live_api: 65536
          context_tokens_marketing_page: 262144
          conflict: Use the conservative live-API value until an authenticated call resolves the discrepancy.
        - id: tencent/hy3-free
          context_tokens: 262144
          callability: not_fully_proven
          caveat: The live entry has no supported endpoint type.
    important_distinction: The $0 Hacker gateway plan does not make paid upstream inference free; only the aliases and router above are zero-cost.
    sources:
      - {type: official_product, title: OrcaRouter offers, url: https://www.orcarouter.ai/offers}
      - {type: official_product, title: Free router, url: https://www.orcarouter.ai/models/orcarouter/free}
      - {type: live_catalog, title: Models endpoint, url: https://api.orcarouter.ai/v1/models}
      - {type: official_pricing, title: Pricing, url: https://www.orcarouter.ai/pricing}
      - {type: company_press_release, title: Hosted OrcaRouter launch, url: https://www.prnewswire.com/news-releases/orcarouter-launches-the-open-llm-api-router--zero-markup-mit-licensed-100-models-302766356.html, published_at: "2026-05"}

  - id: nvidia_build
    name: NVIDIA Build / NVIDIA-hosted NIM
    website: https://build.nvidia.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: development_only
    production_use_allowed: false
    api:
      base_url: https://integrate.api.nvidia.com/v1
      compatibility: LLM NIMs are OpenAI-compatible; non-LLM NIMs can use task-specific APIs.
      authentication: NVIDIA Developer API key
    eligibility:
      account_required: true
      developer_program_membership_required: true
      membership_cost_usd: 0
      payment_method_required: not_documented
    limits:
      public_numeric_limits: false
      official_rule: Varies by model and concurrency; inspect the signed-in account.
      current_marketing_phrase: Unlimited prototyping
      allowed_use: [prototyping, research, development, testing, learning]
    models:
      snapshot_at: "2026-08-21T21:13:22-05:00"
      verification: live_catalog
      discovery_url: https://integrate.api.nvidia.com/v1/models
      count: 102
      caveat: The catalog includes chat, vision, embedding, safety, detection, parsing, and translation models; not every ID accepts chat completions.
      model_ids:
        - 01-ai/yi-large
        - adept/fuyu-8b
        - ai21labs/jamba-1.5-large-instruct
        - aisingapore/sea-lion-7b-instruct
        - baai/bge-m3
        - bigcode/starcoder2-15b
        - databricks/dbrx-instruct
        - deepseek-ai/deepseek-coder-6.7b-instruct
        - deepseek-ai/deepseek-v4-flash-0731
        - google/codegemma-1.1-7b
        - google/codegemma-7b
        - google/deplot
        - google/diffusiongemma-26b-a4b-it
        - google/gemma-2b
        - google/gemma-3-12b-it
        - google/gemma-3-4b-it
        - google/gemma-4-31b-it
        - google/recurrentgemma-2b
        - ibm/granite-3.0-3b-a800m-instruct
        - ibm/granite-3.0-8b-instruct
        - ibm/granite-34b-code-instruct
        - ibm/granite-8b-code-instruct
        - meta/codellama-70b
        - meta/llama-3.1-70b-instruct
        - meta/llama-3.1-8b-instruct
        - meta/llama-3.2-11b-vision-instruct
        - meta/llama-3.2-1b-instruct
        - meta/llama-3.2-3b-instruct
        - meta/llama-3.2-90b-vision-instruct
        - meta/llama-3.3-70b-instruct
        - meta/llama-guard-4-12b
        - meta/llama2-70b
        - meta/muse-glimmer-30b
        - microsoft/kosmos-2
        - microsoft/phi-3-vision-128k-instruct
        - microsoft/phi-3.5-moe-instruct
        - minimaxai/minimax-m3
        - mistralai/codestral-22b-instruct-v0.1
        - mistralai/mistral-7b-instruct-v0.3
        - mistralai/mistral-large
        - mistralai/mistral-large-2-instruct
        - mistralai/mistral-nemotron
        - mistralai/mixtral-8x22b-v0.1
        - moonshotai/kimi-k2.6
        - moonshotai/kimi-k3
        - nv-mistralai/mistral-nemo-12b-instruct
        - nvidia/ai-synthetic-video-detector
        - nvidia/cosmos-reason2-8b
        - nvidia/embed-qa-4
        - nvidia/ising-calibration-1.5-31b
        - nvidia/llama-3.1-nemoguard-8b-content-safety
        - nvidia/llama-3.1-nemoguard-8b-topic-control
        - nvidia/llama-3.1-nemotron-51b-instruct
        - nvidia/llama-3.1-nemotron-70b-instruct
        - nvidia/llama-3.1-nemotron-nano-8b-v1
        - nvidia/llama-3.1-nemotron-nano-vl-8b-v1
        - nvidia/llama-3.1-nemotron-safety-guard-8b-v3
        - nvidia/llama-3.1-nemotron-ultra-253b-v1
        - nvidia/llama-3.2-nemoretriever-1b-vlm-embed-v1
        - nvidia/llama-3.2-nv-embedqa-1b-v1
        - nvidia/llama-3.3-nemotron-super-49b-v1
        - nvidia/llama-3.3-nemotron-super-49b-v1.5
        - nvidia/llama-nemotron-embed-1b-v2
        - nvidia/llama-nemotron-embed-vl-1b-v2
        - nvidia/llama3-chatqa-1.5-70b
        - nvidia/mistral-nemo-minitron-8b-8k-instruct
        - nvidia/nemoretriever-parse
        - nvidia/nemotron-3-embed-1b
        - nvidia/nemotron-3-nano-30b-a3b
        - nvidia/nemotron-3-nano-omni-30b-a3b-reasoning
        - nvidia/nemotron-3-super-120b-a12b
        - nvidia/nemotron-3-ultra-550b-a55b
        - nvidia/nemotron-3.5-content-safety
        - nvidia/nemotron-3.5-lightning-30b-a3b
        - nvidia/nemotron-4-340b-instruct
        - nvidia/nemotron-4-340b-reward
        - nvidia/nemotron-mini-4b-instruct
        - nvidia/nemotron-nano-12b-v2-vl
        - nvidia/nemotron-nano-3-30b-a3b
        - nvidia/nemotron-parse
        - nvidia/neva-22b
        - nvidia/nv-embed-v1
        - nvidia/nv-embedcode-7b-v1
        - nvidia/nv-embedqa-e5-v5
        - nvidia/nv-embedqa-mistral-7b-v2
        - nvidia/nvclip
        - nvidia/nvidia-nemotron-nano-9b-v2
        - nvidia/riva-translate-4b-instruct
        - nvidia/riva-translate-4b-instruct-v1.1
        - nvidia/riva-translate-4b-instruct-v2
        - nvidia/vila
        - openai/gpt-oss-120b
        - openai/gpt-oss-20b
        - poolside/laguna-xs-2.1
        - snowflake/arctic-embed-l
        - stepfun-ai/step-3.7-flash
        - thinkingmachines/inkling
        - writer/palmyra-creative-122b
        - writer/palmyra-fin-70b-32k
        - writer/palmyra-med-70b
        - writer/palmyra-med-70b-32k
        - zyphra/zamba2-7b-instruct
    history:
      - {date: "2024-07-29", event: "Free NIM access for Developer Program members announced."}
      - {date: "2024-09-04", event: "Older scheme documented 1,000 initial credits and up to 5,000; current 2026 docs instead describe prototyping access.", current_status: superseded_or_unconfirmed}
    sources:
      - {type: official_product, title: NVIDIA NIM for Developers, url: https://developer.nvidia.com/nim}
      - {type: official_docs, title: Run NIM Anywhere, url: https://docs.api.nvidia.com/nim/re/docs/run-anywhere}
      - {type: live_catalog, title: Models endpoint, url: https://integrate.api.nvidia.com/v1/models}
      - {type: announcement, title: Free NIM access announcement, url: "https://developer.nvidia.com/blog/?p=86238", published_at: "2024-07-29"}

  - id: hetzner_experiments
    name: Hetzner Experiments Inference API
    website: https://experiments.hetzner.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: experimental
    production_use_allowed: false
    api:
      base_url: https://inference.hetzner.com/api/v1
      compatibility: [openai_models, openai_completions, openai_chat_completions]
      authentication: Hetzner API token
    eligibility:
      account_required: true
      hetzner_customer_required: true
      payment_for_inference_required: false
    limits:
      window_seconds: 60
      requests: 10
      input_tokens: 4000000
      output_tokens: 100000
      scope: per_api_key
    models:
      snapshot_at: "2026-08-21T21:13:22-05:00"
      verification: official_current
      discovery_url: https://inference.hetzner.com/api/v1/models
      free_models:
        - {id: Qwen/Qwen3.6-35B-A3B-FP8, context_tokens: 262144, modalities: [text, image]}
        - {id: Qwen3.8-27B, context_tokens: 262144, modalities: [text, image]}
    duration: Free while experimental; Hetzner promises advance email notice if that changes.
    service_level: Best effort with no availability guarantee.
    data_handling: Usage metadata is retained, but prompt and response content is not retained absent legal compulsion.
    sources:
      - {type: official_docs, title: Inference API, url: https://docs.hetzner.com/general/company-and-policy/experiments/inference/, published_at: "2026-07-24"}
      - {type: official_docs, title: Experiments Platform, url: https://docs.hetzner.com/general/company-and-policy/experiments/experiments-platform/, published_at: "2026-07-24"}
      - {type: official_tutorial, title: OpenCode with Hetzner Inference API, url: https://community.hetzner.com/tutorials/opencode-with-hetzner-inference-api-systemd-sandbox/, published_at: "2026-07-15"}

  - id: groqcloud
    name: GroqCloud
    website: https://console.groq.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      base_url: https://api.groq.com/openai/v1
      compatibility: [openai_chat_completions, openai_responses, audio_transcriptions, audio_translations]
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      scope: organization
      caveat: The signed-in organization Limits page is authoritative if it differs from this public-table snapshot.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://console.groq.com/docs/models
      free_models_and_limits:
        - {id: canopylabs/orpheus-arabic-saudi, rpm: 10, rpd: 100, tpm: 1200, tpd: 3600}
        - {id: canopylabs/orpheus-v1-english, rpm: 10, rpd: 100, tpm: 1200, tpd: 3600}
        - {id: groq/compound, rpm: 30, rpd: 250, tpm: 70000}
        - {id: groq/compound-mini, rpm: 30, rpd: 250, tpm: 70000}
        - {id: meta-llama/llama-prompt-guard-2-22m, rpm: 30, rpd: 14400, tpm: 15000, tpd: 500000}
        - {id: meta-llama/llama-prompt-guard-2-86m, rpm: 30, rpd: 14400, tpm: 15000, tpd: 500000}
        - {id: openai/gpt-oss-120b, rpm: 30, rpd: 1000, tpm: 8000, tpd: 200000}
        - {id: openai/gpt-oss-20b, rpm: 30, rpd: 1000, tpm: 8000, tpd: 200000}
        - {id: openai/gpt-oss-safeguard-20b, rpm: 30, rpd: 1000, tpm: 8000, tpd: 200000}
        - {id: qwen/qwen3.6-27b, rpm: 30, rpd: 1000, tpm: 8000, tpd: 200000}
        - {id: whisper-large-v3, rpm: 20, rpd: 2000, audio_seconds_per_hour: 7200, audio_seconds_per_day: 28800}
        - {id: whisper-large-v3-turbo, rpm: 20, rpd: 2000, audio_seconds_per_hour: 7200, audio_seconds_per_day: 28800}
    history:
      - {date: "2024-03-01", event: GroqCloud launch}
      - {date: "2024-04-02", event: Early demand and free-access announcement}
    sources:
      - {type: official_docs, title: Rate limits, url: https://console.groq.com/docs/rate-limits}
      - {type: official_docs, title: Billing FAQ, url: https://console.groq.com/docs/billing-faqs}
      - {type: official_docs, title: Models, url: https://console.groq.com/docs/models}
      - {type: announcement, title: GroqCloud demand announcement, url: https://groq.com/newsroom/demand-for-real-time-ai-inference-from-groq-accelerates-week-over-week, published_at: "2024-04-02"}

  - id: google_gemini_api
    name: Google AI Studio / Gemini Developer API
    website: https://ai.google.dev/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      native_base_url: https://generativelanguage.googleapis.com
      openai_compatible_base_url: https://generativelanguage.googleapis.com/v1beta/openai/
      authentication: Google API key
    eligibility:
      account_required: true
      payment_method_required: false
      geographic_availability_applies: true
    privacy_caveat: Free-tier content may be used to improve Google products; paid-tier content is not used that way under the published terms.
    limits:
      dimensions: [requests_per_minute, tokens_per_minute, requests_per_day]
      public_fixed_numbers: null
      source_of_truth: AI Studio project rate-limit dashboard
      reset: Daily request quotas reset at midnight Pacific time.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://ai.google.dev/gemini-api/docs
      free_model_ids:
        - gemini-3.7-flash
        - gemini-3.6-flash
        - gemini-3.5-flash
        - gemini-3.5-flash-lite
        - gemini-3.1-flash-lite
        - gemini-2.5-flash
        - gemini-2.5-flash-lite
        - gemini-3.5-live-translate-preview
        - gemini-3.1-flash-live-preview
        - gemini-3.1-flash-tts-preview
        - gemini-2.5-flash-native-audio-preview-12-2025
      explicit_exclusions:
        - {id: gemini-3.1-pro-preview, reason: No free-token tier}
        - {id: gemini-3.1-flash-image, reason: Image generation is paid}
        - {id: gemini-2.0-flash, reason: Shut down on 2026-06-01 despite stale pricing rows}
        - {id: gemini-2.0-flash-lite, reason: Shut down on 2026-06-01 despite stale pricing rows}
    sources:
      - {type: official_pricing, title: Gemini API pricing, url: https://ai.google.dev/gemini-api/docs/pricing}
      - {type: official_docs, title: Rate limits, url: https://ai.google.dev/gemini-api/docs/rate-limits}
      - {type: official_docs, title: Deprecations, url: https://ai.google.dev/gemini-api/docs/deprecations}
      - {type: official_changelog, title: Gemini API changelog, url: https://ai.google.dev/gemini-api/docs/changelog}

  - id: mistral
    name: Mistral Studio / La Plateforme
    website: https://mistral.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: recurring_credit
    api:
      base_url: https://api.mistral.ai/v1
      compatibility: openai_style_rest
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
      intended_use: evaluation_and_prototyping
    limits:
      monthly_api_credit_usd: 10
      dimensions: [requests_per_second, tokens_per_minute, tokens_per_month]
      numeric_limits: dashboard_only
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://docs.mistral.ai/models
      always_zero_priced_model_ids: [mistral-moderation-2603, labs-leanstral-1-5]
      representative_ids_eligible_for_monthly_credit:
        - mistral-small-2603
        - mistral-medium-3-5
        - mistral-large-2512
        - ministral-14b-2512
        - ministral-8b-2512
        - ministral-3b-2512
      caveat: The $10 allowance is monetary, so usable token volume depends on model prices.
    history:
      - {date: "2023-12-11", event: La Plateforme announcement}
      - {date: "2024-09-17", event: Free API tier announcement}
    sources:
      - {type: official_pricing, title: Mistral pricing, url: https://mistral.ai/pricing}
      - {type: official_docs, title: Inference pricing, url: https://docs.mistral.ai/inference/pricing}
      - {type: official_docs, title: Model catalog, url: https://docs.mistral.ai/models}
      - {type: announcement, title: September 2024 release, url: https://mistral.ai/fr/news/september-24-release/, published_at: "2024-09-17"}

  - id: huggingface_inference_providers
    name: Hugging Face Inference Providers
    website: https://huggingface.co/inference/models
    status: current
    qualifies: true
    confidence: high
    free_kind: recurring_credit
    api:
      base_url: https://router.huggingface.co/v1
      compatibility: [openai_chat_completions, openai_responses, openai_models, huggingface_sdk]
      authentication: Hugging Face access token
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      monthly_routed_credit_usd: 0.10
      rate_limits: Provider and model specific.
      overage: Purchased credits required after the monthly allowance.
      caveat: BYOK calls do not consume Hugging Face monthly credits.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      coverage: Dynamic routed catalog of more than 200 models; the credit applies to eligible routed models.
      discovery_url: https://router.huggingface.co/v1/models
      human_catalog_url: https://huggingface.co/inference/models
    history:
      - {date: "2025-01-28", event: Inference Providers launched}
    sources:
      - {type: official_docs, title: Pricing, url: https://huggingface.co/docs/inference-providers/main/en/pricing}
      - {type: official_docs, title: Inference Providers overview, url: https://huggingface.co/docs/inference-providers/en/index}
      - {type: announcement, title: Inference Providers launch, url: https://huggingface.co/blog/inference-providers, published_at: "2025-01-28"}

  - id: cloudflare_workers_ai
    name: Cloudflare Workers AI
    website: https://developers.cloudflare.com/workers-ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      native_pattern: https://api.cloudflare.com/client/v4/accounts/{account_id}/ai/run/{model}
      compatibility: [openai_chat_completions, openai_responses, openai_embeddings]
      authentication: Cloudflare API token
    eligibility:
      account_required: true
      payment_method_required: false
      model_license_acceptance_may_be_required: true
    limits:
      neurons_per_day: 10000
      reset: "00:00 UTC"
      free_plan_exhausted_behavior: Requests fail until reset.
      paid_plan_overage: $0.011 per 1,000 neurons.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: All current Workers AI catalog entries except those marked as requiring Workers Paid.
      current_catalog_count: 84
      discovery_url: https://developers.cloudflare.com/workers-ai/models/
      representative_free_model_ids:
        - "@cf/google/gemma-4-26b-a4b-it"
        - "@cf/zai-org/glm-4.7-flash"
        - "@cf/nvidia/nemotron-3-120b-a12b"
        - "@cf/openai/gpt-oss-120b"
        - "@cf/meta/llama-3.1-8b-instruct"
        - "@cf/deepseek-ai/deepseek-r1-distill-qwen-32b"
        - "@cf/mistralai/mistral-small-3.1-24b-instruct"
      paid_only_model_ids:
        - "@cf/moonshotai/kimi-k2.6"
        - "@cf/moonshotai/kimi-k2.7-code"
        - "@cf/zai-org/glm-5.2"
        - "@cf/deepseek-ai/deepseek-v4-flash-0731"
        - "@cf/deepseek-ai/deepseek-v4-pro-0813"
    history:
      - {date: "2024-04", event: General availability with a daily free allocation}
      - {date: "2026-07-28", event: Cloudflare began marking selected premium models as Workers Paid-only}
    sources:
      - {type: official_pricing, title: Workers AI pricing, url: https://developers.cloudflare.com/workers-ai/platform/pricing/}
      - {type: official_catalog, title: Workers AI models, url: https://developers.cloudflare.com/workers-ai/models/}
      - {type: official_changelog, title: Models that require Workers Paid, url: https://developers.cloudflare.com/changelog/post/2026-07-28-models-require-workers-paid/, published_at: "2026-07-28"}
      - {type: announcement, title: Workers AI general availability, url: https://blog.cloudflare.com/workers-ai-ga-huggingface-loras-python-support/, published_at: "2024-04"}

  - id: cohere
    name: Cohere
    website: https://cohere.com/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: development_only
    production_use_allowed: false
    commercial_use_allowed: false
    api:
      base_url: https://api.cohere.ai/v2
      compatibility: cohere_v2_rest
      authentication: Trial API key
    eligibility:
      account_required: true
      payment_method_required: false
      allowed_use: evaluation_and_prototyping
    limits:
      calls_per_month: 1000
      chat_rpm: 20
      rerank_rpm: 10
      embed_inputs_per_minute: 2000
      embed_image_inputs_per_minute: 5
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Trial keys can evaluate the current Cohere catalog under endpoint limits.
      current_chat_families:
        - Command A+
        - Command A Reasoning
        - Command A Translate
        - Command A Vision
        - Command A
        - Command R+
        - Command R
        - Command R7B
        - North Mini Code
      exact_verified_id: command-a-plus-05-2026
    sources:
      - {type: official_pricing, title: Cohere pricing, url: https://cohere.com/pricing}
      - {type: official_docs, title: Rate limits, url: https://docs.cohere.com/v2/docs/rate-limits}
      - {type: official_docs, title: Chat API, url: https://docs.cohere.com/v2/docs/chat-api}
      - {type: official_changelog, title: Trial key pricing update, url: https://docs.cohere.com/v1/changelog/pricing-update-and-new-dashboard-ui, published_at: "2022-10-18"}

  - id: vercel_ai_gateway
    name: Vercel AI Gateway
    website: https://vercel.com/ai-gateway
    status: current
    qualifies: true
    confidence: high
    free_kind: recurring_credit
    api:
      base_url: https://ai-gateway.vercel.sh/v1
      compatibility: [openai_chat_completions, openai_responses, anthropic_and_provider_sdks]
      authentication: Vercel AI Gateway key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      included_credit_usd_per_month: 5
      important_transition: After a team purchases credits it moves to paid status and no longer receives the monthly free credit.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      coverage: The $5 monthly credit can pay for any available catalog model until exhausted.
      discovery_url: https://ai-gateway.vercel.sh/v1/models
      native_zero_price_model_ids:
        - poolside/laguna-s-2.1-free
      caveat: The native $0 model is volatile; query the discovery URL for current prices.
    history:
      - {date: "2025-08-21", event: AI Gateway became generally available}
    sources:
      - {type: official_pricing, title: AI Gateway pricing, url: https://vercel.com/docs/ai-gateway/pricing}
      - {type: live_catalog, title: Models API, url: https://ai-gateway.vercel.sh/v1/models}
      - {type: official_catalog, title: AI Gateway models, url: https://vercel.com/ai-gateway/models}
      - {type: announcement, title: AI Gateway general availability, url: https://vercel.com/changelog/ai-gateway-is-now-generally-available, published_at: "2025-08-21"}

  - id: ibm_watsonx_ai_runtime
    name: IBM watsonx.ai Runtime
    website: https://www.ibm.com/products/watsonx-ai
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      base_url_pattern: https://{region}.ml.cloud.ibm.com/ml/v1
      compatibility: ibm_watsonx_rest_and_sdks
      authentication: IBM Cloud IAM
    eligibility:
      account_required: true
      payment_method_required_for_new_accounts: true
      caveat: New IBM Cloud accounts require a card for identity verification even when using Lite services.
    limits:
      foundation_tokens_per_month: 300000
      requests_per_second: 2
      capacity_unit_hours: 20
      extraction_pages: 100
      plan: Lite
      inactivity: A Lite service can be deleted after 30 days of inactivity.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Region-available foundation models eligible under the Lite token quota.
      discovery_url_pattern: GET https://{region}.ml.cloud.ibm.com/ml/v1/foundation_model_specs
      representative_current_model_ids:
        - granite-4-h-small
        - granite-4-h-tiny
        - granite-4-h-micro
        - granite-3-1-8b-base
      caveat: Model availability is regional and dynamic; use the API rather than older static catalog pages.
    history:
      - {date: "2023-05-09", event: IBM announced the watsonx platform}
    sources:
      - {type: official_docs, title: watsonx.ai Runtime plans, url: "https://www.ibm.com/docs/en/watsonx/saas?topic=cloud-watsonxai-runtime-plans"}
      - {type: official_catalog, title: watsonx.ai Runtime service, url: https://cloud.ibm.com/catalog/services/pm-20}
      - {type: official_docs, title: List foundation models programmatically, url: "https://dataplatform.cloud.ibm.com/docs/content/wsj/analyze-data/fm-prompt-notebook-list-models.html?context=wx"}
      - {type: announcement, title: IBM unveils watsonx, url: https://newsroom.ibm.com/2023-05-09-IBM-Unveils-the-Watsonx-Platform-to-Power-Next-Generation-Foundation-Models-for-Business, published_at: "2023-05-09"}

  - id: opencode_zen
    name: OpenCode Zen
    website: https://opencode.ai/zen
    status: current_promotional
    qualifies: true
    headline_ongoing_free: false
    confidence: medium_high
    free_kind: promotional
    api:
      base_url: https://opencode.ai/zen/v1
      compatibility: openai_chat_completions
      authentication: Zen API key; account and billing setup details may vary.
    limits:
      published_numeric_limits: false
      duration: Limited-time free models with no common public expiration stated.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://opencode.ai/zen/v1/models
      free_model_ids:
        - big-pickle
        - deepseek-v4-flash-free
        - x-preview-f-free
        - muse-spark-1.2-contributor-free
        - mimo-v2.5-free
        - hy3-free
        - nemotron-3-ultra-free
        - nemotron-3.5-lightning-free
        - laguna-s-2.1-free
      caveat: Model objects did not expose price fields; free status was cross-checked against the first-party Zen documentation and labels.
    sources:
      - {type: official_docs, title: Zen documentation, url: https://opencode.ai/docs/zen}
      - {type: live_catalog, title: Zen models API, url: https://opencode.ai/zen/v1/models}

  - id: zai
    name: Z.AI Model API
    website: https://z.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://api.z.ai/api/paas/v4
      compatibility: openai_style_chat_completions
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: not_documented
    limits:
      published_numeric_limits: false
      source_of_truth: Signed-in rate-limit dashboard.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://docs.z.ai/guides/overview/pricing
      free_models:
        - {id: glm-4.7-flash, context_tokens: 200000}
        - {id: glm-4.6v-flash, context_tokens: 128000, modality: vision_and_text}
        - id: glm-4.5-flash
          context_tokens: 200000
          caveat: Global pricing still lists it, while Chinese documentation says it was to route to 4.7 after a 2026-01-30 shutdown; treat as an alias/conflict.
    sources:
      - {type: official_pricing, title: Z.AI pricing, url: https://docs.z.ai/guides/overview/pricing}
      - {type: official_docs, title: Models overview, url: https://docs.z.ai/guides/overview/overview}
      - {type: official_api_reference, title: Chat completion API, url: https://docs.z.ai/api-reference/llm/chat-completion}

  - id: llmapi_ai
    name: LLM.API
    website: https://llmapi.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://api.llmapi.ai/v1
      compatibility: openai_compatible
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      without_purchased_credits: 5 requests per 10 minutes
      after_adding_credits: 20 requests per minute for free models
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.llmapi.ai/v1/models
      free_models:
        - {id: zaya1-8b, context_tokens: 131072, modalities: text_to_text}
    sources:
      - {type: official_docs, title: API resources and limits, url: https://docs.llmapi.ai/resources}
      - {type: live_catalog, title: Models API, url: https://api.llmapi.ai/v1/models}
      - {type: official_catalog, title: ZAYA1-8B model page, url: https://llmapi.ai/models/}

  - id: api_airforce
    name: Api.Airforce
    website: https://api.airforce/
    status: current
    qualifies: true
    confidence: medium_high
    free_kind: rotating_zero_price
    api:
      base_url: https://api.airforce/v1
      compatibility: openai_compatible
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      requests_per_minute: 1
      requests_per_day: 1000
      per_model_daily_token_cap: true
      per_model_daily_token_cap_amount: unpublished
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://api.airforce/v1/models
      coverage: Exact free variants are identified by :free aliases in the authenticated dynamic catalog.
      verified_examples:
        - gpt-oss-120b
        - gpt-oss-20b
        - qwen3-30b-a3b-fp8
        - glm-4.7-flash
      caveat: The unauthenticated catalog timed out during the audit, so this is not asserted as a complete snapshot.
    sources:
      - {type: official_pricing, title: Pricing, url: https://api.airforce/pricing/}
      - {type: official_docs, title: Quickstart, url: https://api.airforce/docs/quickstart/}
      - {type: official_docs, title: Models API, url: https://api.airforce/docs/api/models/}
      - {type: official_catalog, title: GPT OSS 120B, url: https://api.airforce/models/gpt-oss-120b/}

  - id: llm7
    name: LLM7
    website: https://llm7.io/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      base_url: https://api.llm7.io/v1
      compatibility: openai_compatible
      authentication: Optional for anonymous access; free token raises quota.
    eligibility:
      account_required: false
      payment_method_required: false
    limits:
      anonymous:
        tokens_per_day: 500000
        requests_per_hour: 60
        requests_per_minute: 10
        requests_per_second: 1
      free_account:
        tokens_per_day: 1000000
        requests_per_hour: 250
        requests_per_minute: 60
        requests_per_second: 2
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://docs.llm7.io/guides/models
      free_router_ids: [default, fast]
      paid_router_ids: [pro]
      caveat: Concrete model IDs are being phased out; default and fast dynamically select free routes.
    sources:
      - {type: official_product, title: LLM7 home and quota table, url: https://llm7.io/}
      - {type: official_docs, title: Models, url: https://docs.llm7.io/guides/models}

  - id: modelscope_inference
    name: ModelScope API-Inference
    website: https://modelscope.cn/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: development_only
    commercial_use_allowed: false
    api:
      base_url: https://api-inference.modelscope.cn/v1/
      compatibility: openai_compatible
      authentication: ModelScope token
    eligibility:
      account_required: true
      alibaba_cloud_link_required: true
      real_name_verification_required: true
      allowed_use: noncommercial_and_nonprofit
    limits:
      calls_per_day_per_account: 2000
      calls_per_model_per_day: 200
      selected_expensive_models_calls_per_day: 100
      concurrency: Dynamic and model-specific.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Models carrying the API-Inference badge in the live ModelScope catalog.
      discovery_url: https://modelscope.cn/models?filter=inference_type&page=1
      current_example: Qwen/Qwen3.5-35B-A3B
      lower_quota_examples: [DeepSeek-R1-0528, DeepSeek-V3.2-Exp]
    history:
      - {date: "2024-12-06", event: Early free API-Inference announcement with a Qwen 2.5-era catalog, now superseded}
      - {date: "2026-08-14", event: Current free API-Inference resource page updated}
    sources:
      - {type: official_docs, title: API-Inference limits, url: https://modelscope.cn/docs/model-service/API-Inference/limits}
      - {type: official_docs, title: API-Inference introduction, url: https://modelscope.cn/docs/model-service/API-Inference/intro}
      - {type: official_product, title: Free API-Inference resources, url: https://www.modelscope.cn/learn/1409, updated_at: "2026-08-14"}

  - id: awanllm
    name: AwanLLM
    website: https://www.awanllm.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      base_url: https://api.awanllm.com/v1
      compatibility: openai_compatible
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      tokens: unlimited
      requests_per_minute: 20
      small_model_requests_per_day: 200
      medium_model_requests_per_day: 10
      large_model_requests_per_day: 10
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://www.awanllm.com/models
      free_model_ids:
        - Meta-Llama-3.1-8B-Instruct
        - Meta-Llama-3-8B-Instruct
        - Awanllm-Llama-3-8B-Dolfin
        - Awanllm-Llama-3-8B-Cumulus
        - Meta-Llama-3.1-70B-Instruct
        - Meta-Llama-3-70B-Instruct
    sources:
      - {type: official_pricing, title: Pricing, url: https://www.awanllm.com/pricing}
      - {type: official_catalog, title: Models, url: https://www.awanllm.com/models}
      - {type: official_docs, title: Quick start, url: https://www.awanllm.com/quick-start}

  - id: arliai
    name: Arli AI
    website: https://www.arliai.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      base_url: https://api.arliai.com/v1
      compatibility: openai_compatible
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      requests_per_model_per_two_days: 5
      max_context_tokens: 12000
      concurrent_requests: 1
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      coverage: All models in the dynamic text-generation catalog are trialable under the per-model quota.
      discovery_url: https://api.arliai.com/model/all
    sources:
      - {type: official_pricing, title: Pricing, url: https://www.arliai.com/pricing}
      - {type: official_docs, title: Text generation limits, url: https://www.arliai.com/docs/textgen}
      - {type: live_catalog, title: Public model catalog, url: https://api.arliai.com/model/all}

  - id: freeinference_org
    name: FreeInference.org
    website: https://freeinference.org/
    operator: Harvard SEAS MadSys Lab
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    api:
      base_url: https://freeinference.org/v1
      compatibility: [openai_compatible, anthropic_compatible]
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
      intended_use: research_and_education
    limits:
      public_numeric_limits: false
      official_text: Generous quota
    privacy_caveat: Prompts and responses are logged; anonymized derivatives may be open-sourced.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://doc.freeinference.org/models
      free_chat_model_ids:
        - glm-5.1
        - minimax-m2.5
        - minimax-m3
        - qwen3.6-35b
        - diffusiongemma
        - deepseek-v4-flash
      free_embedding_model_ids: [bge-m3]
      paid_or_pro_exclusions: [glm-5.2, glm-5.3, kimi-k2.7-code]
    sources:
      - {type: official_product, title: FreeInference home, url: https://freeinference.org/}
      - {type: official_docs, title: Models, url: https://doc.freeinference.org/models}

  - id: fastrouter
    name: FastRouter
    website: https://fastrouter.ai/
    status: current
    qualifies: true
    confidence: medium_high
    free_kind: rotating_zero_price
    api:
      base_url: https://api.fastrouter.ai/v1
      compatibility: openai_compatible
      authentication: API key
    limits:
      published_numeric_limits: false
      caveat: The free gateway plan alone is BYOK/pass-through; only the model-specific free routes below are zero-priced inference.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://fastrouter.ai/models
      free_model_ids:
        - openai/gpt-oss-120b:free
        - google/gemma-4-26b-a4b-it
        - nvidia/nemotron-3-nano-30b:free
        - nvidia/nemotron-3-super:free
        - sarvam/sarvam-105b:free
    sources:
      - {type: official_catalog, title: Models, url: https://fastrouter.ai/models}
      - {type: official_pricing, title: Pricing, url: https://fastrouter.ai/pricing}

  - id: kilo_ai_gateway
    name: Kilo AI Gateway
    website: https://kilo.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://api.kilo.ai/api/gateway
      compatibility: openai_compatible
      authentication: Optional for free models; authenticated use supported.
    eligibility:
      account_required: false
      payment_method_required: false
    limits:
      requests_per_hour_per_ip: 200
      price_usd: 0
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.kilo.ai/api/gateway/models
      free_model_ids:
        - cohere/north-mini-code:free
        - dots-studio/dots-3-note-preview:free
        - kilo-auto/free
        - liquid/lfm-2.5-2.6b:free
        - nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free
        - nvidia/nemotron-3-super-120b-a12b:free
        - nvidia/nemotron-3-ultra-550b-a55b:free
        - nvidia/nemotron-3.5-content-safety:free
        - nvidia/nemotron-3.5-lightning:free
        - poolside/laguna-s-2.1:free
        - poolside/laguna-xs-2.1:free
        - stepfun/step-3.7-flash:free
        - tencent/hy3:free
        - thinkingmachines/inkling-small:free
        - thinkingmachines/inkling:free
    sources:
      - {type: official_docs, title: Models and providers, url: https://kilo.ai/docs/gateway/models-and-providers}
      - {type: official_docs, title: Usage and billing, url: https://kilo.ai/docs/gateway/usage-and-billing}
      - {type: live_catalog, title: Gateway models API, url: https://api.kilo.ai/api/gateway/models}

  - id: scaleway_generative_apis
    name: Scaleway Generative APIs
    website: https://www.scaleway.com/en/generative-apis/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      compatibility: openai_compatible
      authentication: Scaleway credentials
    eligibility:
      account_required: true
      valid_payment_method_required_for_base_limits: true
    limits:
      free_tokens: 1000000
      transcription_minutes: 60
      recurrence: not_explicitly_documented
      caveat: Recorded as an account-level trial allowance, not a recurring free tier.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Free allowance is pooled across the current serverless catalog.
      discovery_url: https://www.scaleway.com/en/docs/generative-apis/reference-content/supported-models/
      representative_current_model_ids:
        - glm-5.2
        - deepseek-v4-flash-0731
        - gpt-oss-120b
        - mistral-small-3.2-24b-instruct-2506
        - pixtral-12b-2409
        - qwen3.6-35b-a3b
    history:
      - {event: Earlier beta pages advertised fully free access; current token-metered pricing supersedes that claim.}
    sources:
      - {type: official_pricing, title: Model as a Service pricing, url: https://www.scaleway.com/en/pricing/model-as-a-service/}
      - {type: official_docs, title: Supported models, url: https://www.scaleway.com/en/docs/generative-apis/reference-content/supported-models/}
      - {type: official_docs, title: Rate limits, url: https://www.scaleway.com/en/docs/generative-apis/reference-content/rate-limits/}

  - id: qwen_cloud
    name: Alibaba Cloud Model Studio / Qwen API
    website: https://www.alibabacloud.com/en/product/modelstudio
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      compatibility: [dashscope_native, openai_compatible]
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
      new_user_only: true
    limits:
      duration_days: 90
      basis: Each eligible model has its own token quota; existing models start at account activation and newly released models start at release.
      recurrence: false
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Models and per-model token amounts in the current free-quota table.
      discovery_url: https://docs.qwencloud.com/resources/free-quota
    history:
      - {date: "2026-04-15", event: Separate Qwen OAuth free access for Qwen Code was retired; this does not retire Model Studio new-user quotas.}
    sources:
      - {type: official_docs, title: Free quota, url: https://docs.qwencloud.com/resources/free-quota}
      - {type: official_pricing, title: Pricing overview, url: https://docs.qwencloud.com/developer-guides/getting-started/pricing}
      - {type: official_docs, title: Qwen Code authentication changes, url: https://qwenlm.github.io/qwen-code-docs/en/users/configuration/auth/}

  - id: sea_lion_api
    name: SEA-LION API
    operator: AI Singapore
    website: https://sea-lion.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    nus_relationship: AI Singapore is a national program hosted by the National University of Singapore. This is independent of the requested Nous Research provider.
    api:
      base_url: https://api.sea-lion.ai/v1
      compatibility: [openai_chat_completions, openai_embeddings]
      authentication: Bearer trial API key
    eligibility:
      account_required: true
      signup_method: Google account through SEA-LION Playground
      keys_per_user: 1
      payment_method_required: not_documented
    limits:
      requests_per_minute_per_user: 10
      effective_date: "2026-06-04"
      total_tokens_or_duration: not_documented
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://api.sea-lion.ai/v1/models
      documented_current_model_ids:
        - aisingapore/Qwen-SEA-LION-v4.5-27B-IT
        - aisingapore/Llama-SEA-LION-v3.5-70B-R
        - aisingapore/SEA-Guard
        - aisingapore/SEA-LION-ModernBERT-Embedding-600M
      other_recently_documented_model_ids:
        - aisingapore/Gemma-SEA-LION-v4-27B-IT
    caveat: The credential is explicitly called a trial key and no permanent total allowance is published.
    sources:
      - {type: official_docs, title: SEA-LION API inference guide, url: https://docs.sea-lion.ai/guides/inferencing/api}
      - {type: official_catalog, title: SEA-LION models, url: https://sea-lion.ai/models/}

  - id: ndif
    name: NSF National Deep Inference Fabric
    operator: Northeastern University with NCSA/UIUC
    website: https://ndif.us/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: development_only
    audience: research
    production_use_allowed: false
    api:
      compatibility: NNsight Python remote-execution API with access to and intervention on model internals; not OpenAI-compatible.
      authentication: Free NDIF API key
    eligibility:
      account_required: true
      payment_method_required: false
      allowed_use: research
    limits:
      published_numeric_limits: false
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://ndif.us/status/
      featured_model_ids:
        - meta-llama/Llama-3.1-70B
        - meta-llama/Llama-3.1-8B
        - meta-llama/Llama-3.1-405B
        - openai/gpt-oss-120b
      caveat: Live deployment status, not this static featured list, determines availability.
    sustainability: NSF Award 2408455 with compute from Delta at NCSA.
    sources:
      - {type: official_product, title: NDIF, url: https://ndif.us/}
      - {type: official_docs, title: Get started, url: https://ndif.us/get-started/}
      - {type: live_status, title: Model deployment status, url: https://ndif.us/status/}

  - id: ai_horde
    name: AI Horde
    website: https://aihorde.net/
    status: current
    qualifies: true
    headline_ongoing_free: conditional
    confidence: high
    free_kind: community_capacity
    api:
      base_url: https://aihorde.net/api/v2
      compatibility: Native asynchronous REST; submit text generation then poll status. Not OpenAI-compatible.
      authentication: Anonymous key 0000000000 or free registered key
    eligibility:
      account_required: false
      payment_method_required: false
    limits:
      mechanism: Nonmonetary kudos and queue priority rather than a fixed request quota.
      anonymous_priority: lowest
      load_shedding: Anonymous use can be restricted during load.
      monetary_purchase: Kudos cannot be bought or sold.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://aihorde.net/api/v2/status/models?type=text
      caveat: Availability, speed, queue, and context depend on volunteer workers.
      live_text_model_ids:
        - aphrodite/TheDrummer/Cydonia-24B-v4.3
        - aphrodite/TheDrummer/Skyfall-31B-v4.2
        - coder3101/gemma-4-E4B-it-qat-q4_0-unquantized-heretic
        - google/gemma-4-31b
        - koboldcpp/Angelic_Eclipse-12B
        - koboldcpp/Cydonia-24B-v4.3
        - koboldcpp/digo-prayudha/unsloth-llama-3.2-1b-gguf
        - koboldcpp/Gemma-3-1B
        - koboldcpp/gemma-4-31B-it-heretic
        - koboldcpp/gemma-4-E2B-it-qat-UD-Q4_K_XL.gguf
        - koboldcpp/Gemma-4-E4B-it-Ultra-Uncensored-Heretic
        - koboldcpp/Gemma-4-E4B-Uncensored-HauhauCS-Aggressive
        - koboldcpp/khrystan-worker
        - koboldcpp/L3-8B-Stheno-v3.2-Q5_K_M
        - koboldcpp/L3-Super-Nova-RP-8B
        - koboldcpp/Llama-3.2-1B-Instruct
        - koboldcpp/Llama-3.2-3B
        - koboldcpp/llama-3.2-3b-instruct-q4_k_m
        - koboldcpp/Meta-Llama-3-2-3B-Instruct.Q4_K_M
        - koboldcpp/mini-magnum-12b-v1.1
        - koboldcpp/MN-12B-Mag-Mell-R1.Q5_K_M
        - koboldcpp/mradermacher/Cerebras-GPT-111M-instruction-GGUF
        - koboldcpp/mradermacher/pythia-70m-deduped.f16.gguf
        - koboldcpp/Qwen_Qwen3-0.6B-IQ4_XS
        - koboldcpp/Qwen/Qwen3.5-0.8B
        - koboldcpp/Rocinante-X-12B
    sustainability: Distributed volunteer inference capacity.
    sources:
      - {type: official_repository, title: AI Horde, url: https://github.com/Haidra-Org/AI-Horde}
      - {type: official_integration_docs, title: Integration guide, url: https://github.com/Haidra-Org/AI-Horde/blob/main/README_integration.md}
      - {type: live_catalog, title: Live text models, url: "https://aihorde.net/api/v2/status/models?type=text"}

  - id: pollinations
    name: Pollinations.ai
    website: https://pollinations.ai/
    status: current_community
    qualifies: true
    headline_ongoing_free: conditional
    confidence: medium
    free_kind: rotating_zero_price
    api:
      base_url: https://gen.pollinations.ai
      compatibility: [openai_chat_completions, images, embeddings, audio]
      authentication: Bearer key from enter.pollinations.ai
    limits:
      mechanism: Endpoint-specific rate limits plus Pollen credits and BYOP.
      caveat: The recurring free Pollen refill amount is not clearly documented.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://gen.pollinations.ai/text/models
      openai_discovery_url: https://gen.pollinations.ai/v1/models
      live_catalog_count: 188
      observed_zero_price_community_models:
        - {id: Spit-fires/muse-glimmer, rpm: null}
        - {id: chigwell/llm7-fast, rpm: 250}
        - {id: "YoannDev90/muse-glimmer-30b:free", rpm: 10}
        - {id: chirag-gamer/gpt-oss-120b, rpm: 12}
        - {id: "YoannDev90/laguna-s-2.1:free", rpm: 30}
        - {id: "vendouple/laguna-s-2.1:free", rpm: 5}
        - {id: MarcosFRG/glm-4.6v-flash, rpm: 2}
        - {id: "YoannDev90/diffusiongemma-26b-a4b-it:free", rpm: 10}
        - {id: "vendouple/muse-glimmer-30b:free", rpm: 5}
      caveat: These are independently operated alpha/community endpoints, not a permanence guarantee; native Pollinations models now have positive Pollen prices.
    sources:
      - {type: official_repository, title: Pollinations repository, url: https://github.com/pollinations/pollinations}
      - {type: official_api_docs, title: API docs, url: https://github.com/pollinations/pollinations/blob/main/APIDOCS.md}
      - {type: live_catalog, title: Rich text model catalog, url: https://gen.pollinations.ai/text/models}

  - id: puter_js
    name: Puter.js AI
    website: https://puter.com/
    status: current_user_pays
    qualifies: true
    headline_ongoing_free: conditional
    confidence: high
    free_kind: user_pays
    api:
      compatibility: Browser and Node JavaScript puter.ai.chat() plus Puter Workers; not a conventional shared server API key.
      authentication: Each end user signs into Puter.
    eligibility:
      developer_provider_key_required: false
      end_user_account_required: true
    limits:
      free_monthly_end_user_allowance: amount_not_documented
      exhausted_behavior: End user is prompted to upgrade.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.puter.com/puterai/chat/models/details
      live_catalog_count: 879
      exact_zero_price_model_count: 0
      coverage: The user allowance can fund supported models; “free” means developer-free/user-pays, not zero model prices.
    sources:
      - {type: official_docs, title: User-pays model, url: https://docs.puter.com/user-pays-model/}
      - {type: official_docs, title: Puter AI, url: https://docs.puter.com/AI/}
      - {type: official_tutorial, title: Free LLM API tutorial, url: https://developer.puter.com/tutorials/free-llm-api/}
      - {type: live_catalog, title: Model details API, url: https://api.puter.com/puterai/chat/models/details}

  - id: public_ai
    name: Public AI Inference Utility
    operator: Public AI nonprofit
    website: https://platform.publicai.co/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.publicai.co/v1
      compatibility: [openai_compatible, huggingface_inference_provider]
      authentication: Bearer API key plus User-Agent header
    eligibility:
      account_required: true
      payment_method_required: not_documented
    limits:
      free_tier_requests_per_minute: 100
      starter_credit_amount: not_documented
      overage: Positive token prices are deducted from the wallet after starter credit.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://platform.publicai.co/models
      current_model_ids:
        - swiss-ai/apertus-v1.5-8b
        - swiss-ai/apertus-v1.5-8b-thinking
        - swiss-ai/apertus-v1.5-70b
        - swiss-ai/apertus-v1.5-70b-thinking
        - swiss-ai/apertus-8b-instruct
        - swiss-ai/apertus-70b-instruct
        - aisingapore/Gemma-SEA-LION-v4-27B-IT
        - aisingapore/Qwen-SEA-LION-v4-32B-IT
        - allenai/Olmo-3-7B-Instruct
        - speakleash/Bielik-11B-v3.0-Instruct
        - utter-project/EuroLLM-22B-Instruct-2512
    history:
      - {date: "2025-09-17", event: A Hugging Face launch article said Public AI was free at the time; current starter-credit and wallet pricing supersedes that statement.}
    sources:
      - {type: official_docs, title: Public AI API docs, url: https://platform.publicai.co/docs}
      - {type: official_pricing, title: Plans, url: https://platform.publicai.co/plans}
      - {type: official_catalog, title: Models, url: https://platform.publicai.co/models}
      - {type: historical_announcement, title: Public AI joins Inference Providers, url: https://huggingface.co/blog/inference-providers-publicai, published_at: "2025-09-17"}

  - id: lightning_ai_model_apis
    name: Lightning AI Model APIs / litAI
    website: https://lightning.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: recurring_credit
    api:
      base_url: https://lightning.ai/api/v1
      compatibility: [openai_chat_completions, litai]
      authentication: Lightning API key
    eligibility:
      account_required: true
      payment_method_required: false
      phone_verification_required: true
      country_availability_applies: true
      one_free_account_per_person: true
    limits:
      advertised_model_api_tokens_per_month: 30000000
      requests_per_minute: 15
      tokens_per_minute: 120000
      separate_platform_credits_usd_per_month: 15
      reset: Free platform balance is topped up to $15 on the first of each month and does not accumulate.
      caveat: Public pages do not fully reconcile the 30M-token promise with token-priced litAI calls and the separate $15 balance; the signed-in billing page is authoritative.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://api.lightning.ai/models
      coverage: Dynamic routed catalog, plus custom models deployed with Lightning credits.
      representative_current_ids: [lightning-ai/nvidia-nemotron-3-ultra-550b-a55b, lightning-ai/deepseek-v4-pro, lightning-ai/gemma-4-31B-it, lightning-ai/gpt-oss-120b, anthropic/claude-opus-4-8, google/gemini-3.5-flash, openai/gpt-5.5-2026-04-23]
    sources:
      - {type: official_docs, title: Model APIs, url: https://lightning.ai/docs/overview/model-apis}
      - {type: official_catalog, title: Model API catalog, url: https://api.lightning.ai/models}
      - {type: official_pricing, title: Lightning pricing, url: https://lightning.ai/pricing}
      - {type: official_docs, title: Account creation and free-credit eligibility, url: https://lightning.ai/docs/platform/overview/faq/create-account}
      - {type: official_docs, title: Billing, url: https://lightning.ai/docs/overview/faq/billing}

  - id: modal
    name: Modal
    website: https://modal.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: recurring_credit
    api:
      deployment_base_url_pattern: https://{workspace}--{endpoint}.{region}.modal.direct/v1
      shared_base_url: https://inference.us-west.modal.direct/v1
      compatibility: [openai_chat_completions, openai_responses, custom_http]
      authentication: Modal proxy token
    eligibility:
      account_required: true
      payment_method_required: true
    limits:
      plan: Starter
      plan_price_usd_per_month: 0
      included_credit_usd_per_month: 30
      containers: 100
      gpu_concurrency: 10
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Public Hugging Face models and custom/private fine-tunes deployed by the user; there is no fixed free-model list.
      exact_official_examples: [Qwen/Qwen3.5-4B, Qwen/Qwen3.6-27B, aisingapore/Qwen-SEA-LION-v4.5-27B-IT]
    sources:
      - {type: official_pricing, title: Pricing, url: https://modal.com/pricing}
      - {type: official_docs, title: Billing, url: https://modal.com/docs/guide/billing}
      - {type: official_docs, title: Endpoints, url: https://modal.com/docs/guide/endpoints}
      - {type: official_docs, title: Shared endpoint integrations, url: https://modal.com/docs/guide/endpoint-integrations}

  - id: beam_cloud
    name: Beam
    website: https://www.beam.cloud/
    status: current
    qualifies: true
    confidence: high
    free_kind: recurring_credit
    api:
      base_url_pattern: https://{app}-{deployment}-v{version}.app.beam.cloud
      compatibility: [custom_rest, openai_compatible_via_vllm_or_sglang]
      authentication: Bearer Beam token
    eligibility:
      account_required: true
      payment_method_required: true
      plan_selection_required: Developer pay-as-you-go
    limits:
      plan_price_usd_per_month: 0
      included_credit_usd_per_month: 30
      reset: monthly
      gpu_concurrency: 5
      cpu_concurrency: 30
      api_requests: unlimited
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Arbitrary public, private, or custom user-deployed models; no fixed hosted catalog.
      exact_official_examples: [OpenGVLab/InternVL3-8B-AWQ, 01-ai/Yi-Coder-9B-Chat, Qwen/Qwen2.5-7B-Instruct]
    sources:
      - {type: official_pricing, title: Pricing, url: https://www.beam.cloud/pricing}
      - {type: official_docs, title: FAQ, url: https://docs.beam.cloud/v2/resources/faq}
      - {type: official_docs, title: Endpoint overview, url: https://docs.beam.cloud/v2/endpoint/overview}
      - {type: official_example, title: vLLM OpenAI-compatible server, url: https://docs.beam.cloud/v2/examples/vllm}
      - {type: official_announcement, title: Monthly credit confirmation, url: https://www.beam.cloud/blog/serverless-gpu-reinforcement-learning, published_at: "2026-07-02"}

  - id: sail_research
    name: Sail Research
    website: https://www.sailresearch.com/
    status: current
    qualifies: true
    confidence: high
    free_kind: recurring_credit
    api:
      base_url: https://api.sailresearch.com/v1
      compatibility: [openai_responses, openai_chat_completions, anthropic_messages]
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required_for_monthly_refresh: true
    limits:
      included_credit_usd_per_month: 5
      reset: monthly
      rate_limits: No strict public limits; service priority depends on completion window.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://docs.sailresearch.com/
      representative_current_models: [GLM-5.2, DeepSeek V4 Pro, Kimi-K2.6, gpt-oss-120b, Nemotron 3 Super 120B, Gemma 4 31B IT]
    sources:
      - {type: official_pricing, title: Sail Research pricing and API overview, url: https://www.sailresearch.com/}
      - {type: official_docs, title: Sail Research documentation, url: https://docs.sailresearch.com/}

  - id: cartesia
    name: Cartesia
    website: https://www.cartesia.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    modalities: [text_to_speech, speech_to_text]
    api:
      base_url: https://api.cartesia.ai
      compatibility: cartesia_rest_and_websocket
      authentication: Cartesia API key
    eligibility:
      account_required: true
      payment_method_required: not_explicitly_documented
      commercial_use_on_free_plan: false
    limits:
      plan_price_usd_per_month: 0
      credits_per_month: 20000
      approximate_tts_minutes_per_month: 27
      approximate_stt_hours_per_month: 1.85
      tts_concurrent_requests: 2
      stt_concurrent_requests: 8
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      current_families: [Sonic-3.5, Ink-2]
    sources:
      - {type: official_pricing, title: Pricing, url: https://www.cartesia.ai/pricing}
      - {type: official_docs, title: API documentation, url: https://docs.cartesia.ai/}

  - id: elevenlabs_api
    name: ElevenLabs API
    website: https://elevenlabs.io/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    modalities: [text_to_speech, speech_to_text, sound_generation]
    api:
      base_url: https://api.elevenlabs.io/v1
      compatibility: elevenlabs_rest_and_websocket
      authentication: ElevenLabs API key
    eligibility:
      account_required: true
      payment_method_required: false
      free_plan_license: noncommercial_with_attribution
    limits:
      plan_price_usd_per_month: 0
      credits_reset: monthly
      flash_or_turbo_tts_characters_per_month: 20000
      multilingual_tts_characters_per_month: 10000
      caveat: Most, but not every, API endpoint is available on Free; each request consumes the account's credits.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      representative_current_models: [eleven_v3, eleven_flash_v2_5]
    sources:
      - {type: official_pricing, title: ElevenAPI pricing, url: https://elevenlabs.io/pricing/api}
      - {type: official_docs, title: API availability on Free, url: https://elevenlabs.io/docs/help-center/technical/how-much-does-it-cost-to-use-the-api}
      - {type: official_docs, title: Billing and Free plan, url: https://elevenlabs.io/docs/overview/administration/billing}

  - id: ovhcloud_ai_endpoints
    name: OVHcloud AI Endpoints
    website: https://www.ovhcloud.com/en/public-cloud/ai-endpoints/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    headline_general_llm_free: false
    scope_note: Qualifies for modality-agnostic hosted inference; it does not currently provide a free general-purpose chat model.
    api:
      unified_openai_base_url: https://oai.endpoints.kepler.ai.cloud.ovh.net/v1
      compatibility: [openai_images, model_specific_rest, grpc_tts, openai_style_guard]
      authentication: Anonymous calls are allowed on documented zero-price endpoints; an OVHcloud access key raises limits.
    limits:
      anonymous_requests_per_minute_per_ip_per_model: 2
      authenticated_requests_per_minute_per_project_per_model: 400
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      exact_zero_price_ids: [Qwen3Guard-Gen-8B, Qwen3Guard-Gen-0.6B, stable-diffusion-xl-base-v10, nvr-tts-de-de, nvr-tts-en-us, nvr-tts-es-es, nvr-tts-it-it]
      discovery_url: https://www.ovhcloud.com/en/public-cloud/ai-endpoints/catalog/
    separate_trial:
      credit_usd: 200
      duration: one month
      payment_method_required: true
      caveat: This finite Public Cloud trial is separate from the explicitly zero-price models.
    sources:
      - {type: official_catalog, title: AI Endpoints catalog, url: https://www.ovhcloud.com/en/public-cloud/ai-endpoints/catalog/}
      - {type: official_model_page, title: SDXL endpoint and rate limits, url: https://www.ovhcloud.com/en-gb/public-cloud/ai-endpoints/catalog/stable-diffusion-xl/}
      - {type: official_model_page, title: Qwen Guard endpoint, url: https://www.ovhcloud.com/de/public-cloud/ai-endpoints/catalog/qwen-guard-gen-8b/}
      - {type: official_trial, title: Public Cloud free trial, url: https://www.ovhcloud.com/en/public-cloud/free-trial/}

  - id: requesty
    name: Requesty LLM Gateway
    website: https://www.requesty.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://router.requesty.ai/v1
      eu_base_url: https://router.eu.requesty.ai/v1
      compatibility: [openai_compatible, anthropic_compatible]
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      requests_per_day: 200
      reset: daily
      caveat: The Free plan is restricted to zero-priced models; PAYG and BYOK catalogs are separate.
    models:
      snapshot_at: "2026-08-21T23:28:44-05:00"
      verification: live_catalog
      discovery_url: https://router.requesty.ai/v1/models
      live_total_models: 668
      zero_price_model_ids: [nvidia/nemotron-3-super-120b-a12b, nvidia/nemotron-3-nano-omni-30b-a3b-reasoning, nvidia/nemotron-3-nano-30b-a3b, nvidia/nemotron-3.5-content-safety, nvidia/nemotron-3-ultra-550b-a55b, nvidia/muse-glimmer-30b, nvidia/nemotron-3.5-lightning-30b-a3b, google/gemma-4-31b-it, poolside/laguna-xs.2, poolside/laguna-m.1, mistral/leanstral-1-5, novita/inclusionai/ling-3.0-tiny]
      caveat: Two old Poolside entries and nemotron-3-nano-30b-a3b reported max_output_tokens=0; runtime callability needs authenticated verification.
    sources:
      - {type: official_pricing, title: Pricing, url: https://www.requesty.ai/pricing}
      - {type: official_docs, title: Documentation index, url: https://docs.requesty.ai/llms.txt}
      - {type: live_catalog, title: Models API, url: https://router.requesty.ai/v1/models}

  - id: inception_platform
    name: Inception Platform
    website: https://platform.inceptionlabs.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.inceptionlabs.ai/v1
      compatibility: [openai_chat_completions, fim_completions, edit_completions]
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
      new_account_only: true
    limits:
      free_tokens: 100000000
      recurrence: false
      expiry: not_published
      requests_per_minute: 1000
      input_tokens_per_minute: 1000000
      output_tokens_per_minute: 100000
      caveat: The July 2026 announcement raised the older 10M-token documentation value to 100M; after depletion billing information is required.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      free_coverage: Shared initial allowance across both production models.
      model_ids: [mercury-2, mercury-edit-2]
    sources:
      - {type: official_docs, title: Platform quickstart, url: https://docs.inceptionlabs.ai/get-started/get-started}
      - {type: official_docs, title: Models and pricing, url: https://docs.inceptionlabs.ai/get-started/models}
      - {type: official_docs, title: Rate limits, url: https://docs.inceptionlabs.ai/get-started/rate-limits}
      - {type: official_announcement, title: Mercury 2 free-token increase, url: https://www.inceptionlabs.ai/blog/mercury-2-10x-free-tokens, published_at: "2026-07-29"}

  - id: poolside_direct_api
    name: Poolside direct API
    website: https://poolside.ai/models
    status: current_promotional
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: promotional
    api:
      base_url: https://inference.poolside.ai/v1
      compatibility: openai_compatible
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: not_documented
    limits:
      numeric_rate_limits: not_published
      expiry: not_published
      official_duration: Free for a limited time.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      current_names: [Laguna S 2.1, Laguna XS 2.1]
      confirmed_api_id: poolside/laguna-s-2.1
      likely_second_api_id: poolside/laguna-xs-2.1
      caveat: The second exact ID matches current router catalogs, but Poolside's direct models endpoint requires authentication.
    sources:
      - {type: official_product, title: Poolside models, url: https://poolside.ai/models}
      - {type: official_docs, title: Supported models, url: https://docs.poolside.ai/get-started/supported-models}
      - {type: official_announcement, title: Laguna XS.2 and M.1, url: https://poolside.ai/blog/introducing-laguna-xs2-m1, published_at: "2026-04-28"}

  - id: voyage_ai
    name: Voyage AI
    website: https://www.voyageai.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    modalities: [embeddings, multimodal_embeddings, reranking]
    api:
      base_url: https://api.voyageai.com/v1
      compatibility: voyage_rest
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required_for_initial_quota: false
    limits:
      recurrence: false
      free_text_tokens_by_model:
        200000000: [voyage-4-large, voyage-4, voyage-4-lite, voyage-context-4, voyage-code-3]
        50000000: [voyage-multilingual-2, voyage-finance-2, voyage-law-2, voyage-code-2]
      free_multimodal_text_tokens: 200000000
      free_multimodal_pixels: 150000000000
      free_rerank_tokens: 200000000
      batch_api_included: false
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      current_rerank_ids: [rerank-2.5, rerank-2.5-lite, rerank-2, rerank-2-lite]
      caveat: Current pricing text elsewhere mentions voyage-code-4 where the free-token table says voyage-code-3; preserve the table value until corrected.
    sources:
      - {type: official_pricing, title: Pricing and free tokens, url: https://docs.voyageai.com/docs/pricing}
      - {type: official_docs, title: Rate limits, url: https://docs.voyageai.com/docs/rate-limits}

  - id: ai21_studio
    name: AI21 Studio
    website: https://www.ai21.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.ai21.com/studio/v1
      compatibility: ai21_rest
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required_initially: false
    limits:
      signup_credit_usd: 10
      expires_after_months: 3
      recurrence: false
      coverage: API, SDK, and playground usage.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://docs.ai21.com/docs/models
    sources:
      - {type: official_pricing, title: Usage and cost, url: https://docs.ai21.com/docs/usage-cost}
      - {type: official_docs, title: Models, url: https://docs.ai21.com/docs/models}

  - id: deepgram
    name: Deepgram
    website: https://deepgram.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    modalities: [speech_to_text, text_to_speech, voice_agents]
    api:
      base_url: https://api.deepgram.com/v1
      compatibility: deepgram_rest_and_websocket
      authentication: Deepgram API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      signup_credit_usd: 200
      expires_after_years: 1
      recurrence: false
      coverage: All public model endpoints.
    sources:
      - {type: official_pricing, title: Pricing, url: https://deepgram.com/pricing}
      - {type: official_docs, title: Promotional credit expiry, url: https://developers.deepgram.com/guides/deep-dives/managing-projects}

  - id: jina_ai_search_foundation
    name: Jina Search Foundation API
    website: https://jina.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    modalities: [embeddings, multimodal_embeddings, reranking, classification, vision_language, grounded_research]
    api:
      base_url: https://api.jina.ai/v1
      compatibility: openai_embeddings_and_jina_rest
      authentication: Bearer Jina API key
    eligibility:
      account_required: true
      payment_method_required: false
      signup_token_use: noncommercial_only
      signup_token_license: CC-BY-NC
    limits:
      signup_tokens: 10000000
      recurrence: false
      free_token_expiry: not_documented
      conservative_embedding_and_rerank_requests_per_minute: 100
      conservative_embedding_and_rerank_tokens_per_minute: 100000
      caveat: The product table gives the conservative limits above while the versioned OpenAPI page publishes higher Free-tier limits; use the lower figures until response headers or the dashboard resolve the conflict.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.jina.ai/v1/models
      current_catalog_count: 29
      coverage: Current embedding, reranker, classifier, vision-language, and DeepSearch APIs; catalog entries have positive prices and draw down the one-time balance.
    sources:
      - {type: official_api_reference, title: Jina Search Foundation API, url: https://api.jina.ai/docs}
      - {type: official_product, title: Reader and Search API limits, url: https://jina.ai/reader/}

  - id: jina_ai_reader
    name: Jina Reader and anonymous utility APIs
    website: https://jina.ai/reader/
    status: current
    qualifies: true
    confidence: high
    free_kind: ongoing_free
    modalities: [webpage_to_markdown, pdf_extraction, image_captioning, tokenization, segmentation]
    api:
      reader_base_url: https://r.jina.ai
      segmenter_url: https://api.jina.ai/v1/segment
      compatibility: jina_rest
      authentication: None at anonymous limits; an optional Bearer key raises limits.
    eligibility:
      account_required: false
      payment_method_required: false
    limits:
      reader_anonymous_requests_per_minute_per_ip: 20
      segmenter_anonymous_requests_per_minute_per_ip: 20
      segmenter_tokens_charged: 0
      reader_with_free_key_requests_per_minute: 500
      caveat: Supplying a key to Reader charges its token balance; anonymous basic Reader calls remain free.
    models:
      snapshot_at: "2026-08-21"
      verification: live_call
      coverage: Service-level APIs rather than caller-selected model routes.
      implementation_models_named_by_jina: [ReaderLM-v2, jina-vlm]
      observed_calls: Anonymous Reader and Segmenter requests both returned non-empty HTTP 200 responses; Segmenter reported zero tokens used.
    sources:
      - {type: official_product, title: Reader API and current limits, url: https://jina.ai/reader/}
      - {type: official_terms, title: Legal information, url: https://jina.ai/legal/, updated_at: "2026-05-04"}

  - id: mancer_ai
    name: Mancer AI
    operator: Sunlit Software, Inc.
    website: https://mancer.tech/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://neuro.mancer.tech/oai/v1
      compatibility: [openai_chat_completions, openai_completions, openai_models]
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      numeric_quota: not_published
      official_claim: Free models continue working even with a negative credit balance.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://neuro.mancer.tech/oai/v1/models
      free_model_ids: [mytholite]
      model_details: {name: MythoLite, base: LLaMA 2 13B, context_tokens: 2560, max_completion_tokens: 150}
      caveat: A stale page fragment calls Rocinante free; the live catalog and current model table identify only mytholite at zero price.
    sources:
      - {type: official_product, title: Mancer AI, url: https://mancer.tech/}
      - {type: official_pricing, title: Pricing FAQ, url: https://mancer.tech/pricing}
      - {type: live_catalog, title: Models API, url: https://neuro.mancer.tech/oai/v1/models}
      - {type: official_api_reference, title: OpenAPI specification, url: https://mancer.tech/resources/api-docs-webui.yml}

  - id: mara_inference_cloud
    name: MARA Inference Cloud
    website: https://cloud.mara.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.cloud.mara.com/v1
      compatibility: openai_compatible
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      signup_credit_usd: 5
      expires_after_days: 30
      recurrence: false
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.cloud.mara.com/v1/models
      current_ids: [DeepSeek-V3.1, DeepSeek-V3.2, MiniMax-M2.7, gemma-4-31B-it, gpt-oss-120b]
      caveat: Plans distinguish production from preview access, so do not assume every catalog entry is available to trial accounts without an authenticated check.
    sources:
      - {type: official_pricing, title: Plans, url: https://cloud.mara.com/plans}
      - {type: official_product, title: Dashboard and quickstart, url: https://cloud.mara.com/}
      - {type: live_catalog, title: Models API, url: https://api.cloud.mara.com/v1/models}

  - id: assemblyai
    name: AssemblyAI
    website: https://www.assemblyai.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    modalities: [speech_to_text, speech_understanding, audio_guardrails]
    api:
      base_url: https://api.assemblyai.com
      streaming_url: wss://streaming.assemblyai.com/v3/ws
      compatibility: assemblyai_rest_and_websocket
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      signup_credit_usd: 50
      expiry: none
      recurrence: false
      exclusions: [LLM Gateway]
    sources:
      - {type: official_help, title: Free signup, url: https://support.assemblyai.com/articles/5370767329-can-i-sign-up-for-free}
      - {type: official_docs, title: Account billing and free credits, url: https://www.assemblyai.com/docs/faq/how-to-get-your-api-key}

  - id: speechmatics
    name: Speechmatics
    website: https://www.speechmatics.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    modalities: [speech_to_text, text_to_speech]
    api:
      compatibility: speechmatics_rest_and_realtime_websocket
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      signup_credit_usd: 100
      recurrence: false
      realtime_concurrent_sessions: 2
      languages: 55+
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Batch and real-time STT plus TTS under the shared credit balance.
    sources:
      - {type: official_pricing, title: Speech API pricing, url: https://www.speechmatics.com/pricing}
      - {type: official_announcement, title: Credit-based billing, url: https://www.speechmatics.com/company/articles-and-news/moving-to-credit-based-billing}

  - id: aws_bedrock
    name: Amazon Bedrock
    website: https://aws.amazon.com/bedrock/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      openai_base_url_pattern: https://bedrock-mantle.{region}.api.aws/openai/v1
      compatibility: [openai_responses, openai_chat_completions, anthropic_messages, bedrock_converse, bedrock_invoke]
      authentication: AWS credentials, SigV4, or Bedrock API key depending on endpoint
    eligibility:
      new_aws_customer_only: true
      payment_method_required: true
    limits:
      signup_credit_usd: 100
      additional_earnable_credit_usd: 100
      free_plan_duration_months: 6
      credit_expiry_months_from_account_creation: 12
      recurrence: false
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Dynamic regional multi-provider catalog; models have positive unit prices and draw down general AWS trial credit.
    sources:
      - {type: official_announcement, title: AWS Free Tier credits and six-month plan, url: https://aws.amazon.com/about-aws/whats-new/2025/07/aws-free-tier-credits-month-free-plan/, published_at: "2025-07-16"}
      - {type: official_docs, title: Free Tier FAQ, url: https://docs.aws.amazon.com/awsaccountbilling/latest/aboutv2/free-tier-FAQ.html}
      - {type: official_pricing, title: Bedrock pricing, url: https://aws.amazon.com/bedrock/pricing/}
      - {type: official_docs, title: OpenAI-compatible inference, url: https://docs.aws.amazon.com/bedrock/latest/userguide/inference-chat-completions-mantle.html}

  - id: azure_ai_foundry
    name: Microsoft Foundry Models
    website: https://azure.microsoft.com/products/ai-foundry/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      foundry_resource_base_url: https://{resource}.services.ai.azure.com/api/
      openai_base_url: https://{resource}.openai.azure.com/openai/v1/
      compatibility: [openai_v1, foundry_native]
    eligibility:
      new_customer_only: true
      phone_required: true
      non_prepaid_payment_card_required: true
      automatic_charging: false
    limits:
      general_cloud_credit_usd: 200
      duration_days: 30
      recurrence: false
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Dynamic regional Foundry catalog; positively priced serverless or provisioned inference draws down Azure trial credit.
    sources:
      - {type: official_trial, title: Azure free account, url: https://azure.microsoft.com/free/}
      - {type: official_pricing, title: Foundry Models pricing, url: https://azure.microsoft.com/en-us/pricing/details/ai-foundry-models/microsoft/}
      - {type: official_docs, title: Foundry application integration, url: https://learn.microsoft.com/en-us/azure/foundry/how-to/integrate-with-other-apps}

  - id: google_vertex_ai
    name: Google Cloud Vertex AI
    website: https://cloud.google.com/vertex-ai
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      native_url_pattern: https://{location}-aiplatform.googleapis.com/v1/projects/{project}/locations/{location}/publishers/google/models/{model}:generateContent
      openai_base_url_pattern: https://{location}-aiplatform.googleapis.com/v1/projects/{project}/locations/{location}/endpoints/openapi
      authentication: Google Cloud OAuth/IAM
    eligibility:
      new_customer_only: true
      payment_method_required_for_verification: true
    limits:
      general_cloud_credit_usd: 300
      duration_days: 90
      recurrence: false
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      coverage: Eligible first-party Google-managed Vertex services and models only.
      exclusions:
        - Gemini Developer API / AI Studio usage; its independent free tier is a separate record.
        - Generative AI partner models offered as managed API services.
    sources:
      - {type: official_trial, title: Google Cloud free program, url: https://docs.cloud.google.com/free/docs/free-cloud-features}
      - {type: official_pricing, title: Vertex generative AI pricing, url: https://cloud.google.com/vertex-ai/generative-ai/pricing}
      - {type: official_docs, title: OpenAI-compatible Gemini call on Vertex, url: https://docs.cloud.google.com/vertex-ai/generative-ai/docs/samples/generativeaionvertexai-gemini-chat-completions-non-streaming}

  - id: oracle_oci_generative_ai
    name: OCI Generative AI
    website: https://www.oracle.com/artificial-intelligence/generative-ai/generative-ai-service/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      openai_base_url_pattern: https://inference.generativeai.{region}.oci.oraclecloud.com/openai/v1
      compatibility: [openai_responses, openai_chat_completions, oci_native_inference]
      authentication: OCI Generative AI API key or OCI IAM
    eligibility:
      new_customer_only: true
      mobile_phone_required_for_most_users: true
      payment_card_required_for_most_users: true
    limits:
      general_cloud_credit_usd: 300
      duration_days: 30
      recurrence: false
    caveat: Generative AI is not an OCI Always Free service; only the finite general cloud trial offsets eligible usage.
    sources:
      - {type: official_trial, title: OCI Free Tier, url: https://docs.oracle.com/en-us/iaas/Content/FreeTier/freetier.htm}
      - {type: official_trial, title: Oracle Cloud Free, url: https://www.oracle.com/cloud/free/}
      - {type: official_docs, title: OpenAI-compatible API, url: https://docs.oracle.com/en-us/iaas/Content/generative-ai/openai-compatible-api.htm}

  - id: replicate
    name: Replicate
    website: https://replicate.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.replicate.com/v1
      compatibility: replicate_native_rest
      authentication: Bearer API token
    eligibility:
      account_required: true
      payment_method_required_initially: false
    limits:
      coverage: Dynamic subset of models can be run free for a small unquantified allowance.
      recurrence: none_documented
      no_card_granted_credit_requests_per_second: 1
      no_card_granted_credit_requests_per_minute: 6
      caveat: Billing setup is required after the model-specific initial allowance is exhausted.
    sources:
      - {type: official_docs, title: Billing, url: https://replicate.com/docs/topics/billing}
      - {type: official_docs, title: Prepaid credit, url: https://replicate.com/docs/topics/billing/prepaid-credit}
      - {type: official_docs, title: Prediction rate limits, url: https://replicate.com/docs/topics/predictions/rate-limits}

  - id: cerebras_inference
    name: Cerebras Inference
    website: https://www.cerebras.ai/inference
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.cerebras.ai/v1
      compatibility: openai_compatible
      authentication: API key
    eligibility:
      account_required: true
      verified_payment_method_required: true
    limits:
      signup_credit_usd: 5
      expires_after_days: 30
      recurrence: false
      caveat: Current documentation explicitly says no permanently free tier exists.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      trial_model_limits:
        - {id: gpt-oss-120b, rpm: 5, tpm: 30000, tpd: 1000000}
        - {id: gemma-4-31b, rpm: 5, tpm: 30000, tpd: 1000000}
    history:
      - {date: "2024-08-27", event: Launched with one million free tokens per day; that recurring offer ended.}
    sources:
      - {type: official_docs, title: Current trial and no-free-tier statement, url: https://inference-docs.cerebras.ai/support/rate-limits}
      - {type: official_announcement, title: Inference launch, url: https://www.cerebras.ai/blog/introducing-cerebras-inference-ai-at-instant-speed, published_at: "2024-08-27"}

  - id: clarifai
    name: Clarifai
    website: https://www.clarifai.com/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      native_base_url: https://api.clarifai.com
      openai_base_url: https://api.clarifai.com/v2/ext/openai/v1
      compatibility: [clarifai_native, openai_compatible]
    eligibility:
      phone_verification_required: true
      payment_method_required_initially: false
      payment_method_required_to_recharge: true
    limits:
      signup_credit_usd: 5
      expires_after_days: 30
      recurrence: false
      maximum_welcome_bonuses: 2
      default_requests_per_second: 15
    history:
      - {date: "2026-02-03", event: The recurring Community free plan retired and pay-as-you-go replaced it.}
    sources:
      - {type: official_docs, title: Account billing, url: https://docs.clarifai.com/control/account-billing/}
      - {type: official_docs, title: Inference, url: https://docs.clarifai.com/compute/inference/}
      - {type: official_docs, title: Rate limits, url: https://docs.clarifai.com/resources/api-overview/rate-limits/}
      - {type: official_changelog, title: Community plan retirement, url: https://docs.clarifai.com/product-updates/changelog/release121/, published_at: "2026-02-03"}

  - id: fireworks_ai
    name: Fireworks AI
    website: https://fireworks.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.fireworks.ai/inference/v1
      compatibility: openai_compatible
      authentication: API key
    limits:
      signup_credit_usd: 1
      recurrence: false
      caveat: This is a small promotional signup credit, not a recurring free tier.
    sources:
      - {type: official_pricing, title: Pricing, url: https://fireworks.ai/pricing}
      - {type: official_billing_faq, title: Billing and pricing FAQ, url: https://docs.fireworks.ai/faq-new/billing-pricing/how-much-does-fireworks-cost}

  - id: nebius_token_factory
    name: Nebius Token Factory
    website: https://nebius.com/token-factory
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    api:
      base_url: https://api.tokenfactory.nebius.com/v1
      compatibility: openai_compatible
      authentication: API key
    limits:
      signup_credit_usd: 1
      recurrence: false
    source: {type: official_pricing, title: Token Factory prices, url: https://nebius.com/token-factory/prices}

  - id: novita_ai
    name: Novita AI
    website: https://novita.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    limits:
      signup_voucher_amount: dashboard_only
      recurrence: false
      caveat: The finite new-user voucher is followed by prepaid top-ups; promotional vouchers can expire.
    sources:
      - {type: official_quickstart, title: Quickstart, url: https://novita.ai/docs/guides/quickstart}
      - {type: official_pricing, title: Pricing, url: https://novita.ai/pricing}

  - id: hyperbolic
    name: Hyperbolic
    website: https://www.hyperbolic.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    eligibility:
      phone_verification_required: true
    limits:
      signup_credit_usd: 1
      recurrence: false
    source: {type: official_billing_docs, title: Billing and payments, url: https://www.hyperbolic.ai/docs/general/billing-payments}

  - id: waterfall
    name: Waterfall
    website: https://www.getwaterfall.org/
    status: current_community
    qualifies: true
    headline_ongoing_free: conditional
    confidence: high
    free_kind: rotating_zero_price
    api:
      base_url: https://api.getwaterfall.org/v1
      compatibility: openai_compatible
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      numeric_limit: not_public
      policy: community_rate_limits
      routing: free_smart
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.getwaterfall.org/v1/models
      live_total_models: 439
      free_model_ids: [nemotron-30b-free, nemotron-9b-free, nemotron-3-super-120b-free, nemotron-3-nano-30b-free, nemotron-nano-12b-vl-free, nemotron-nano-9b-v2-free, gemma-4-31b-free, gemma-4-26b-free, nemotron-12b-video-free, gemma-4-26b-a4b-it-free, gemma-4-31b-it-free, nemotron-3-nano-30b-a3b-free, nemotron-3-nano-omni-30b-a3b-reasoning-free, nemotron-3-super-120b-a12b-free, nemotron-3-ultra-550b-a55b-free, nemotron-3.5-content-safety-free, nemotron-nano-12b-v2-vl-free, laguna-xs-2.1-free, north-mini-code-free, laguna-s-2.1-free, nemotron-3.5-lightning-free, lfm-2.5-2.6b-free, dots-3-note-preview-free, glm-5.2-free, inkling-free, inkling-small-free, lyria-3-clip-preview-free, lyria-3-pro-preview-free]
    sustainability: Bootstrapped single-maintainer service with no SLA or permanence guarantee.
    sources:
      - {type: official_product, title: Waterfall, url: https://www.getwaterfall.org/}
      - {type: official_pricing, title: Pricing, url: https://www.getwaterfall.org/pricing/}
      - {type: official_docs, title: Documentation, url: https://www.getwaterfall.org/docs/}
      - {type: live_catalog, title: Models API, url: https://api.getwaterfall.org/v1/models}

  - id: logfare
    name: Logfare
    website: https://logfare.ai/
    status: current_community
    qualifies: true
    headline_ongoing_free: conditional
    confidence: high
    free_kind: ongoing_free
    api:
      base_url: https://logfare.ai/v1
      compatibility: [openai_chat_completions, openai_responses, anthropic_messages, embeddings, images, audio]
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      numeric_limit: none_published
      policy: fair_use
      caveat: Excessive or automated traffic may be throttled or blocked.
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://logfare.ai/v1/models
      standard_model_ids: [gemma-4-26b]
      premium_training_opt_in_ids: [kiro-auto, minimax-m3, moondream3.1, deepseek-v4-pro, glm-5.2, qwen-3.8-2.4t-a95b, kimi-k3, deepseek-v4-flash-0731, qwen-3.8-27b, deepseek-v4-pro-0813]
      other_live_models: [sdxl-lightning, whisper-large-v3-turbo, phoenix-1.0, qwen3-embedding-8b, flux-2-pro, aura-2-en, nova-3, lucid-origin, text-embedding-3-small]
    privacy:
      logging: Request and response bodies, IP, headers, and metadata are logged.
      standard_tier: May be used for internal evaluation after best-effort PII scrubbing.
      premium_tier: Requires prospective opt-in to model-training use.
    sustainability: Independent, self-funded, donation-supported service with no SLA.
    sources:
      - {type: official_product, title: Logfare, url: https://logfare.ai/}
      - {type: official_docs, title: API docs, url: https://logfare.ai/docs}
      - {type: official_terms, title: Terms, url: https://logfare.ai/tos}
      - {type: live_catalog, title: Models API, url: https://logfare.ai/v1/models}

  - id: bazaarlink
    name: BazaarLink
    operator: 集聯科技有限公司
    website: https://bazaarlink.ai/
    status: current
    qualifies: true
    confidence: high
    free_kind: rotating_zero_price
    geography: Taiwan
    api:
      base_url: https://api.bazaarlink.ai/v1
      compatibility: openai_compatible
      authentication: Bearer API key; programmatic agent registration is also supported.
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      requests_per_minute: 10
      requests_per_day: 50
      requests_per_day_after_any_topup: 100
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.bazaarlink.ai/v1/models
      free_model_ids: ["qwen/qwen3.7-flash:free", "auto:free"]
    sources:
      - {type: official_product, title: Free models, url: https://bazaarlink.ai/free}
      - {type: official_docs, title: API docs, url: https://bazaarlink.ai/en/docs/api}
      - {type: official_terms, title: Terms, url: https://bazaarlink.ai/en/terms}
      - {type: live_catalog, title: Models API, url: https://api.bazaarlink.ai/v1/models}

  - id: dreamprompting
    name: DreamPrompting
    website: https://dreamprompting.com/
    status: current_community
    qualifies: true
    headline_ongoing_free: conditional
    confidence: medium_high
    free_kind: ongoing_free
    api:
      base_url: https://dreamprompting.com/api/v1
      compatibility: openai_chat_completions
      authentication: Google sign-in and Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      requests_per_minute_per_ip: 100
      tokens_per_rolling_24_hours: 500000
      requests_per_rolling_24_hours: 5000
      max_input_tokens: 32000
      max_output_tokens: 8192
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      routing: Aggregates upstream free tiers, trials, and the keyless ch.at service.
      representative_ids: [google/gemini-3.1-flash-lite, groq/llama-3.3-70b-versatile, nvidia/meta/llama-3.3-70b-instruct, "openrouter/google/gemma-4-31b-it:free", mistral/mistral-small-latest, cohere/command-a-03-2025, chat/ch.at]
    caveat: Young upstream-dependent service; terms prohibit some unapproved automation while API docs promote agent use.
    sources:
      - {type: official_product, title: DreamPrompting, url: "https://dreamprompting.com/?lang=en"}
      - {type: official_docs, title: API docs, url: https://dreamprompting.com/api-docs}
      - {type: official_catalog, title: Models, url: https://dreamprompting.com/models}
      - {type: official_terms, title: Terms, url: https://dreamprompting.com/terms}

  - id: ch_at
    name: ch.at
    website: https://ch.at/
    status: current_community
    qualifies: true
    headline_ongoing_free: conditional
    confidence: high
    free_kind: community_capacity
    api:
      base_url: https://ch.at/v1
      compatibility: openai_chat_completions
      authentication: none
    eligibility:
      account_required: false
      payment_method_required: false
    limits:
      requests_per_minute_per_ip: 100
      burst: 10
      history_cap: 64KB
    models:
      snapshot_at: "2026-08-21"
      verification: authenticated_call
      operator_selected: true
      caveat: A live unauthenticated generation returned HTTP 200, but the response model is blank and the server ignores the caller's model field; do not publish a stable model ID.
    sustainability: Community service with no SLA.
    sources:
      - {type: official_repository, title: ch.at repository, url: https://github.com/Deep-ai-inc/ch.at}
      - {type: live_endpoint, title: Chat completions endpoint, url: https://ch.at/v1/chat/completions}

  - id: opentyphoon
    name: OpenTyphoon API
    operator: SCB 10X
    website: https://opentyphoon.ai/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: development_only
    geography: Thailand
    api:
      base_url: https://api.opentyphoon.ai/v1
      compatibility: openai_compatible
      authentication: Playground account and Bearer API key
    eligibility:
      account_required: true
      payment_method_required: not_stated
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      current_limits:
        - {id: typhoon-v2.5-30b-a3b-instruct, context_tokens: 128000, rps: 5, rpm: 200}
        - {id: typhoon-v2.1-12b-instruct, context_tokens: 56000, rps: 5, rpm: 200}
        - {id: typhoon-ocr, rpm: 20}
        - {id: typhoon-asr-realtime, rpm: 100}
    caveat: Beta research showcase provided as-is without formal support.
    sources:
      - {type: official_docs, title: Documentation, url: https://docs.opentyphoon.ai/en/}
      - {type: official_docs, title: FAQ, url: https://docs.opentyphoon.ai/en/faq/}
      - {type: official_catalog, title: Models and limits, url: https://docs.opentyphoon.ai/en/models/}
      - {type: official_api_reference, title: API reference, url: https://docs.opentyphoon.ai/en/api-reference/}

  - id: alcf_inference_endpoints
    name: Argonne Leadership Computing Facility Inference Endpoints
    website: https://www.alcf.anl.gov/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: development_only
    geography: United States
    audience: approved_research
    api:
      service_root: https://inference-api.alcf.anl.gov
      compatibility: [openai_chat, openai_responses, anthropic_messages, completions, embeddings, batches]
      authentication: ALCF account and Globus OAuth
    eligibility:
      approved_project_required: true
      programs: [Directors Discretionary, INCITE, ALCC, NAIRR]
      open_research_cost: generally_no_cost_compute_allocation
      proprietary_research: cost_recovery
    limits:
      numeric_user_limit: allocation_specific
      batch_max_requests: 150000
      cold_start_minutes: 10-15
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      discovery_url: https://inference-api.alcf.anl.gov/resource_server/list-endpoints
      representative_ids: [meta-llama/Meta-Llama-3.1-405B-Instruct, meta-llama/Llama-4-Maverick-17B-128E-Instruct, mistralai/Mistral-Large-Instruct-2407, openai/gpt-oss-120b, argonne/AuroraGPT-IT-v4-0125, google/gemma-4-31B-it, nvidia/nemotron-3-super-120b]
    sources:
      - {type: official_docs, title: Inference endpoints, url: https://docs.alcf.anl.gov/services/inference-endpoints/}
      - {type: official_docs, title: Allocation management, url: https://docs.alcf.anl.gov/account-project-management/allocation-management/}
      - {type: official_program, title: Discretionary allocations, url: https://www.alcf.anl.gov/science/directors-discretionary-allocation-program}
      - {type: official_repository, title: Inference endpoints repository, url: https://github.com/argonne-lcf/inference-endpoints}

  - id: fikra_api
    name: Fikra API
    website: https://fikraapi.co.ke/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    geography: Kenya
    api:
      base_url: https://api.fikraapi.co.ke/v1
      compatibility: openai_chat_completions
      authentication: Bearer API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      signup_tokens: 300000
      recurrence: false
      requests_per_minute: 30
      exhausted_behavior: HTTP 402 until top-up.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      model_ids: [fikra-fast-8b, fikra-pro-20b, fikra-pro-120b]
    sources:
      - {type: official_product, title: Fikra API, url: https://fikraapi.co.ke/}
      - {type: official_docs, title: Documentation, url: https://docs.fikraapi.co.ke/}

  - id: sarvam_ai
    name: Sarvam AI
    website: https://www.sarvam.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    geography: India
    api:
      base_url: https://api.sarvam.ai/v2
      compatibility: sarvam_rest_with_openai_style_chat_parameters
      authentication: API subscription key
    eligibility:
      account_required: true
      payment_method_required: not_explicitly_stated
    limits:
      signup_credit_inr: 100
      expiry: none
      recurrence: false
      starter_requests_per_minute: 60
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.sarvam.ai/v2/models
      public_ids: [glm5.2, gemma4, sarvam-105b]
      caveat: glm5.2 and gemma4 require per-key beta whitelisting.
    sources:
      - {type: official_docs, title: Rate limits, url: https://docs.sarvam.ai/api/getting-started/ratelimits}
      - {type: official_catalog, title: Open-source models, url: https://docs.sarvam.ai/api/getting-started/models/open-source}
      - {type: official_pricing, title: API pricing, url: https://web.sarvam.dev/api-pricing}
      - {type: live_catalog, title: Models API, url: https://api.sarvam.ai/v2/models}

  - id: byteplus_modelark
    name: BytePlus ModelArk
    website: https://www.byteplus.com/en/product/modelark
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    geography: Singapore and supported global regions
    api:
      base_url: https://ark.ap-southeast.bytepluses.com/api/v3
      compatibility: [openai_chat_completions, openai_responses]
      authentication: ARK API key
    eligibility:
      enterprise_information_submission_required: true
      payment_method_required: not_explicitly_stated
    limits:
      typical_tokens_per_eligible_model: 500000
      recurrence: once_per_account
      exact_models_and_expiry: Dynamic in Model Activation and Billing Center.
      exclusions: [plugins, knowledge_bases, batch_inference, cache_storage]
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      representative_ids: [seed-1.6, seed-1.6-flash, seed-2-1-turbo, seed-translation, skylark-pro, skylark-vision]
    sources:
      - {type: official_docs, title: Free token package, url: https://docs.byteplus.com/en/docs/modelark/1399514}
      - {type: official_docs, title: API overview, url: https://docs.byteplus.com/api/docs/modelark/1465347}
      - {type: official_terms, title: Free-token campaign terms, url: https://docs.byteplus.com/en/docs/legal/termsandconditions_modelark_free-token_campaign}

  - id: tencent_hunyuan
    name: Tencent Hunyuan
    website: https://cloud.tencent.com/product/hunyuan
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    geography: China
    api:
      base_url: https://api.hunyuan.cloud.tencent.com/v1
      migration_base_url: https://tokenhub.tencentmaas.com/v1
      compatibility: openai_compatible
      authentication: Bearer API key
    eligibility:
      account_required: true
      real_name_verification_required: true
      payment_method_required: false_if_postpaid_not_enabled
    limits:
      activation_tokens: 1000000
      separate_embedding_tokens: 1000000
      expires_after_years: 1
      recurrence: false
      default_concurrency: 5
      exhausted_behavior: Does not auto-switch to paid unless postpaid is explicitly enabled.
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      representative_ids: [Hunyuan-a13b, Hunyuan-role-latest, Hunyuan-translation, Hunyuan-translation-lite, Tencent-HY-Vision-1.5-Instruct, Hunyuan-embedding]
    sources:
      - {type: official_docs, title: Free resource package, url: https://cloud.tencent.com/document/product/1729/97731}
      - {type: official_docs, title: API access, url: https://cloud.tencent.com/document/product/1729/111007}
      - {type: official_docs, title: Migration and current aliases, url: https://cloud.tencent.com/document/product/1729/131925}

  - id: upstage
    name: Upstage
    website: https://www.upstage.ai/
    status: current_trial_and_restricted_program
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    geography: South Korea and supported regions
    api:
      base_url: https://api.upstage.ai/v1
      compatibility: openai_compatible
      authentication: Bearer API key
    general_signup:
      credit_usd: 10
      recurrence: false
      expiry: not_public
      current_model_example: solar-pro3
    institutional_program:
      eligibility: [K-12 schools, universities, university hospitals, nonprofits, NGOs]
      coverage: [solar-pro2, solar-pro3, document-parse]
      duration: up to one year
      final_access_date: "2027-03-31T23:00:00+09:00"
      approval_required: true
    sources:
      - {type: official_guide, title: Console and API guide, url: https://www.upstage.ai/blog/en/guide-1-upstage-console-api}
      - {type: official_pricing, title: API pricing, url: https://www.upstage.ai/pricing/api}
      - {type: official_program, title: AI Initiative 2026, url: https://www.upstage.ai/events/ai-initiative-2026-en}

  - id: maritaca_academic_credits
    name: Maritaca AI Academic Credits
    website: https://www.maritaca.ai/
    status: current_restricted
    qualifies: true
    confidence: high
    free_kind: development_only
    geography: Brazil
    audience: research_and_teaching
    api:
      base_url: https://chat.maritaca.ai/api
      compatibility: openai_compatible
      authentication: Approved account and API key
    eligibility:
      application_required: true
      applicants: [students, faculty, researchers]
      project_types: [teaching, scientific_research]
    limits:
      credit_amount: not_public
      duration: not_public
      tier_0_requests_per_minute: 60
      tier_0_input_tokens_per_minute: 128000
      tier_0_output_tokens_per_minute: 10000
    models:
      snapshot_at: "2026-08-21"
      verification: official_current
      current_ids: [sabia-4-thinking, sabia-4-thinking-br-sp, sabia-4, sabia-4-2026-01-06, sabia-4-br-sp, sabiazinho-4, sabiazinho-4-2026-01-06, sabia-4-small, sabiazim-4, sabiazinho-4-br-sp]
    sources:
      - {type: official_program, title: Academic credits, url: https://www.maritaca.ai/research}
      - {type: official_catalog, title: Models, url: https://docs.maritaca.ai/pt/modelos}
      - {type: official_docs, title: Rate limits, url: https://docs.maritaca.ai/pt/rate-limits}
      - {type: official_pricing, title: Pricing, url: https://docs.maritaca.ai/pt/precos}

  - id: wavespeedai
    name: WaveSpeedAI
    website: https://wavespeed.ai/
    status: current_trial
    qualifies: true
    headline_ongoing_free: false
    confidence: high
    free_kind: trial
    modalities: [language_models, image, video, audio]
    api:
      compatibility: wavespeed_native_rest
      authentication: API key
    eligibility:
      account_required: true
      payment_method_required: false
    limits:
      signup_credit_usd: 1
      recurrence: false
      expiry: not_published
      caveat: Some premium models are unavailable to trial balances; eligible catalog entries have positive prices and draw down the credit.
    sources:
      - {type: official_pricing, title: Pricing and signup credit, url: https://wavespeed.ai/pricing}

conditional_or_needs_verification:
  - id: anyrouter
    name: AnyRouter
    status: conditional
    confidence: high
    api_base_url: https://anyrouter.dev/api/v1
    model_id: anyrouter/free
    zero_token_price: true
    unlock_requirement:
      - Pay $1/month for the Go plan.
      - Or donate a working upstream provider API key, which unlocks Go without a card.
    limits:
      requests_per_day: 1000
      daily_reset: "00:00 UTC"
      go_requests_per_minute: 60
      go_rolling_five_hour_requests: 3000
    current_router_models: [cohere/north-mini-code, deepseek/DeepSeek-V3.1]
    reason_not_headline_free: Requires money or contribution of an upstream credential/capacity.
    sources:
      - {type: official_product, title: Free model, url: https://anyrouter.dev/free}
      - {type: official_model_page, title: anyrouter/free, url: https://anyrouter.dev/model/anyrouter/free}
      - {type: official_docs, title: Rate limits, url: https://docs.anyrouter.dev/features/rate-limits}

  - id: nexusrouter
    name: NexusRouter
    status: needs_authenticated_verification
    confidence: low
    api_base_url: https://api.nexusrouter.net/v1
    compatibility: [openai_compatible, anthropic_compatible]
    public_claims:
      fair_use_tokens_per_day: 1000000
      requests_per_minute: 10
      reset: "07:00 Asia/Jakarta"
      free_model_included: true
    conflict: The current client also contains “Free — 100K tokens/month,” and no exact free model is publicly named.
    recommendation: Do not include in headline counts until the authenticated dashboard and models endpoint resolve the plan and model.
    sources:
      - {type: official_product, title: NexusRouter, url: https://nexusrouter.net/}
      - {type: official_terms, title: Fair use, url: https://nexusrouter.net/fair-use}
      - {type: official_pricing, title: Pricing, url: https://nexusrouter.net/pricing}

  - id: siliconflow
    name: SiliconFlow / SiliconCloud
    status: conflicting_current_sources
    confidence: medium
    api_base_url: https://api.siliconflow.com/v1
    positive_evidence:
      - Generic docs say free models have fixed limits and paid variants use Pro/ prefixes.
      - Chinese docs say real-name-verified users can use models currently labeled free at zero cost.
    negative_evidence:
      - The current global public catalog gives positive prices to formerly cited DeepSeek-V3 and Qwen3-8B models.
      - The older exact free list contains dated Qwen2-era IDs.
      - No current exact zero-price ID was publicly verifiable without authentication.
    china_endpoint_note: The China service may retain a separate real-name-verified free catalog from the global service, but it could not be enumerated anonymously.
    recommendation: Require an authenticated model response or billing-free call before treating an ID as current.
    sources:
      - {type: official_docs, title: Global rate limits, url: https://docs.siliconflow.com/en/userguide/rate-limits/rate-limit-and-upgradation}
      - {type: official_docs, title: China rate limits, url: https://docs.siliconflow.cn/cn/userguide/rate-limits/rate-limit-and-upgradation}
      - {type: official_pricing, title: Current global pricing, url: https://www.siliconflow.com/pricing}
      - {type: official_api_reference, title: Models API, url: https://docs.siliconflow.cn/cn/api-reference/models/get-model-list}

  - id: freemodel_dev
    name: FreeModel.dev
    status: insufficiently_documented_trial
    confidence: low
    public_claim: Signup credits with no card.
    missing: Credit amount, recurrence, exact current catalog, and durable quota are not published clearly enough to validate ongoing free inference.
    source: {type: official_product, url: https://www.freemodel.dev/}

  - id: arouter
    name: ARouter
    status: insufficiently_documented
    confidence: medium
    api_base_url: https://api.arouter.ai/v1
    finding: Documentation describes generic :free routing and possible promotional credits, but exposes no exact public free model, stable quota, or guaranteed starter amount.
    source: {type: official_docs, url: https://docs.arouter.ai/en/faq}

  - id: baseten
    name: Baseten
    website: https://www.baseten.co/
    status: signup_credit_amount_unpublished
    confidence: high
    headline_ongoing_free: false
    free_kind: trial
    api:
      base_url: https://inference.baseten.co/v1
      compatibility: [openai_chat_completions, anthropic_messages]
      authentication: API key
    positive_evidence: Current pricing says new accounts receive credits to experiment for free.
    unresolved: [credit_amount, expiry, card_requirement, recurrence]
    models:
      coverage: Dynamic Model API catalog plus arbitrary self-deployed models; all published token prices are positive.
      representative_ids: [zai-org/GLM-5, deepseek-ai/DeepSeek-V4-Pro, moonshotai/Kimi-K2.6]
    sources:
      - {type: official_pricing, title: Pricing, url: https://www.baseten.co/pricing/}
      - {type: official_docs, title: Inference API overview, url: https://docs.baseten.co/reference/inference-api/overview}

  - id: nscale_serverless_inference
    name: Nscale Serverless Inference
    website: https://www.nscale.com/
    status: promotional_credits_only
    confidence: high
    headline_ongoing_free: false
    free_kind: trial
    api:
      compatibility: openai_compatible
      authentication: Service token
    current_requirement: Current overview says to add at least $5 credit before using the service.
    possible_exception: Current quickstart says early users may be eligible for unspecified promotional credits.
    history: The 2025 launch offered every new user $5, but the current docs no longer make that universal promise.
    recommendation: Do not count without an account-specific promotion.
    sources:
      - {type: official_docs, title: Current overview and minimum credit, url: https://docs.nscale.com/docs/getting-started/overview}
      - {type: official_docs, title: Current quickstart and possible promotions, url: https://docs.nscale.com/docs/getting-started/quickstart}
      - {type: historical_announcement, title: 2025 serverless launch, url: https://www.nscale.com/blog/introducing-nscale-serverless-inference-scalable-ai-without-infrastructure-hassles, published_at: "2025-04-02"}

  - id: bentoml_cloud
    name: BentoML Cloud / Bento Inference Platform
    website: https://www.bentoml.com/
    status: signup_credit_amount_unpublished
    confidence: medium_high
    headline_ongoing_free: false
    free_kind: trial
    positive_evidence: Current pricing promises one-time free compute credit.
    unresolved: [credit_amount, expiry]
    eligibility:
      payment_method_required_for_initial_trial: false
      payment_method_required_for_starter_upgrade: true
    api:
      base_url: deployment_specific
      compatibility: custom_rest_or_user_deployed_openai
    conflict: Older BentoML material mentioned $10, but current pricing no longer confirms that amount.
    sources:
      - {type: official_pricing, title: Pricing, url: https://www.bentoml.com/pricing}
      - {type: historical_product_post, title: Deployment example with older credit claim, url: https://www.bentoml.com/blog/deploying-a-text-to-speech-application-with-bentoml}

  - id: friendli_ai
    name: FriendliAI
    website: https://friendli.ai/
    status: signup_credit_amount_unpublished
    confidence: medium_high
    headline_ongoing_free: false
    free_kind: trial
    api:
      base_url: https://api.friendli.ai
      compatibility: [friendli_rest, openai_compatible]
      authentication: API key
    positive_evidence: Current billing docs confirm promotional signup credit.
    unresolved: [current_amount, expiry, card_requirement]
    exhausted_behavior: Purchased credit with a $10 minimum is required after the promotion.
    sources:
      - {type: official_docs, title: Credits, url: https://friendli.ai/docs/guides/suite/credits}
      - {type: official_pricing, title: Pricing, url: https://friendli.ai/pricing}
      - {type: official_api_reference, title: API introduction, url: https://friendli.ai/docs/openapi/introduction}

  - id: akashml
    name: AkashML
    website: https://akashml.com/
    status: signup_credit_amount_unpublished
    confidence: medium_high
    headline_ongoing_free: false
    free_kind: trial
    api:
      base_url: https://api.akashml.com/v1
      compatibility: [openai_compatible, anthropic_compatible]
      authentication: API key
    positive_evidence: Current FAQ says new users automatically receive free credits.
    unresolved: [credit_amount, expiry, card_requirement, recurrence, exact_trial_model_ids]
    models:
      representative_current_names: [QWEN3.8 27B, DeepSeek V4 Flash 0731, QWEN3.6 35B A3B, Llama 3.3 70B, GPT OSS 120B, GPT OSS 20B]
      caveat: All public model prices are positive and the models endpoint requires authentication.
    sources:
      - {type: official_product, title: AkashML and FAQ, url: https://akashml.com/}
      - {type: official_docs, title: Documentation, url: https://akashml.com/docs}

  - id: eden_ai
    name: Eden AI
    website: https://www.edenai.co/
    status: trial_plus_zero_price_routes_need_runtime_verification
    confidence: medium_high
    headline_ongoing_free: false
    free_kind: trial
    api:
      base_url: https://api.edenai.run/v3
      compatibility: [edenai_unified, openai_compatible_llm]
      authentication: Bearer API token
    trial:
      signup_credit_usd: 10
      recurrence: false
      unresolved: [expiry, payment_method_requirement]
    models:
      snapshot_at: "2026-08-21"
      verification: live_catalog
      discovery_url: https://api.edenai.run/v3/models
      live_total_models: 928
      zero_price_entries: [cloudflare/@cf/google/gemma-2b-it-lora, cloudflare/@cf/google/gemma-7b-it-lora, cloudflare/@cf/meta-llama/llama-2-7b-chat-hf-lora, cloudflare/@cf/mistral/mistral-7b-instruct-v0.2-lora, google/gemma-4-26b-a4b-it, google/gemma-4-31b-it]
      caveat: The catalog reports zero prices, but billing docs do not confirm those routes work at a zero wallet balance; sandbox responses are simulated and do not qualify.
    sources:
      - {type: official_docs, title: Eden AI overview, url: https://www.edenai.co/docs/index.md}
      - {type: official_docs, title: LLM quickstart, url: https://www.edenai.co/docs/v3/quickstart/first-llm-call.md}
      - {type: official_pricing, title: Pricing and signup credit, url: https://www.edenai.co/pricing}
      - {type: live_catalog, title: Models API, url: https://api.edenai.run/v3/models}
      - {type: official_docs, title: Sandbox is simulated, url: https://www.edenai.co/docs/v3/general/sandbox.md}

  - id: sambanova_cloud
    name: SambaNova Cloud
    website: https://cloud.sambanova.ai/
    status: conflicting_current_sources
    confidence: medium_high
    headline_ongoing_free: conditional
    free_kind: ongoing_free
    api:
      base_url: https://api.sambanova.ai/v1
      compatibility: openai_compatible
      authentication: API key
    positive_evidence:
      - Current rate-limit docs define a no-payment-method Free Tier.
      - Official staff posts in 2026 describe accounts without a card as Free Tier accounts.
    negative_evidence:
      - The current plans page says Free users must add a payment method and purchase credits before the first request.
    advertised_free_limits:
      production_model_ids: [DeepSeek-V3.1, Meta-Llama-3.3-70B-Instruct, gpt-oss-120b]
      preview_evaluation_only_ids: [DeepSeek-V3.2, gemma-4-31B-it]
      per_model: {requests_per_minute: 20, requests_per_day: 20, tokens_per_day: 200000}
    recommendation: Require a zero-balance authenticated generation before counting this as an ongoing tier.
    sources:
      - {type: official_docs, title: Model rate limits and Free Tier, url: https://docs.sambanova.ai/docs/en/models/rate-limits}
      - {type: official_plans, title: Contradictory plans page, url: https://cloud.sambanova.ai/plans}
      - {type: official_community_staff, title: No-card Free Tier confirmation, url: https://community.sambanova.ai/t/minimax-2-5-is-now-available-on-the-sambacloud-developer-tier/1601/10, published_at: "2026-03-18"}

  - id: fal_ai_builder_grant
    name: fal Builder Grant
    website: https://fal.ai/builder-grant
    status: application_required
    confidence: high
    headline_ongoing_free: false
    free_kind: selective_program
    eligibility:
      regions: [Europe, Asia]
      intended_use: Direct generative-media applications
      approval_guaranteed: false
      application_frequency: once_per_team_per_3_months
    grant_packs:
      - {name: Starter, credit_usd: 25}
      - {name: Plus, credit_usd: 100, unlock_monthly_usage_usd: 50}
      - {name: Launch, credit_usd: 250, unlock_monthly_usage_usd: 100}
    caveat: Sandbox coupons and free credits explicitly cannot be used through the API or Workflows; only an approved Builder Grant funds API usage.
    sources:
      - {type: official_program, title: Builder Grant, url: https://fal.ai/builder-grant}
      - {type: official_docs, title: Sandbox limitations, url: https://fal.ai/docs/documentation/model-apis/sandbox}
      - {type: official_pricing, title: API pricing, url: https://fal.ai/docs/documentation/model-apis/pricing}

  - id: openai_data_sharing_tokens
    name: OpenAI API complimentary data-sharing tokens
    website: https://platform.openai.com/
    status: conditional_eligibility_and_positive_balance_required
    confidence: high
    headline_ongoing_free: conditional
    free_kind: contribution_based
    eligibility:
      organization_must_be_selected_by_openai: true
      input_output_data_sharing_opt_in_required: true
      positive_account_balance_required: true
      unavailable_to: [Enterprise, Zero Data Retention organizations]
    limits:
      reset: "00:00 UTC daily"
      tiers_1_and_2: {large_model_group_tokens_per_day: 250000, small_model_group_tokens_per_day: 2500000}
      tiers_3_to_5: {large_model_group_tokens_per_day: 1000000, small_model_group_tokens_per_day: 10000000}
      caveat: A request that would cross the quota is billed in full; fine-tuning, evals, and tool use are excluded.
    models:
      representative_large_group: [gpt-5.5-2026-04-23, gpt-5.4-2026-03-05, gpt-4.1-2025-04-14, o3-2025-04-16]
      representative_small_group: [gpt-5.4-mini-2026-03-17, gpt-5.4-nano-2026-03-17, gpt-4.1-mini-2025-04-14, o4-mini-2025-04-16]
    sources:
      - {type: official_help, title: Complimentary tokens for shared API traffic, url: https://help.openai.com/en/articles/10306912-sharing-feedback-and-api-inputs-and-outputs-with-openai}
      - {type: official_docs, title: API data controls, url: https://platform.openai.com/docs/models/default-usage-policies-by-endpoint}

  - id: llm_gateway
    name: LLM Gateway
    website: https://llmgateway.io/
    status: conflicting_current_sources
    confidence: medium
    api_base_url: https://api.llmgateway.io/v1
    public_claim: Pricing says three zero-cost models at 20 requests per minute with no card.
    conflict:
      - Current models page renders zero Free Models.
      - The unauthenticated models endpoint returned 256 models with none marked free.
    recommendation: Do not count until an authenticated catalog or generation confirms a free route.
    sources:
      - {type: official_pricing, url: https://llmgateway.io/pricing}
      - {type: official_catalog, url: https://llmgateway.io/models}
      - {type: live_catalog, url: https://api.llmgateway.io/v1/models}

  - id: baidu_qianfan
    name: Baidu Qianfan
    status: current_trial_terms_incomplete
    confidence: medium
    geography: China
    public_claim: New customers can receive more than one million trial tokens.
    unresolved: [exact_models, per_model_quota, duration, payment_requirement]
    historical_conflict: A 2024 notice promised long-term-free ERNIE Speed/Lite/Tiny routes, but current pricing shows positive prices.
    sources:
      - {type: official_product, url: https://cloud.baidu.com/product/qianfan.html}
      - {type: official_docs, url: https://cloud.baidu.com/doc/qianfan/s/wmh4sv6ya}
      - {type: official_notice, url: https://cloud.baidu.com/news/notice_c7a145b9-5c99-4870-939f-e09ba506aab3}

  - id: tera_promotional_credits
    name: Tera
    website: https://www.tera.gw/
    status: application_required
    confidence: high
    headline_ongoing_free: false
    free_kind: selective_program
    api_base_url: https://api.tera.gw/v1
    eligibility:
      application_required: true
      founder_call_required: true
      grants_per_company: 1
    credits:
      solo_developers_usd: 150
      vc_backed_startups_usd: 250
      expires_after_days: 45
    models:
      representative_ids: [gpt-oss-20b, gpt-oss-120b, DeepSeek-V4-Flash, DeepSeek-V3.2, Qwen3-Coder-480B, kimi-k2-thinking, GLM-5, Kimi-K2.6]
    sources:
      - {type: official_program, title: Credits, url: https://www.tera.gw/credits}
      - {type: official_pricing, title: Pricing, url: https://www.tera.gw/pricing}

  - id: runpod_startup_program
    name: Runpod startup credits
    website: https://www.runpod.io/startup-program
    status: application_required
    confidence: high
    headline_ongoing_free: false
    free_kind: selective_program
    offer:
      startup_credit_usd: 1000
      approval_required: true
    caveat: Normal serverless/public-endpoint inference is prepaid and requires a positive balance; this is not an automatic free tier.
    sources:
      - {type: official_program, title: Startup program, url: https://www.runpod.io/startup-program}
      - {type: official_docs, title: Billing, url: https://docs.runpod.io/accounts-billing/billing}

  - id: jina_ai_llm_serp
    name: Jina LLM-as-SERP API
    status: advertised_free_but_callability_unresolved
    confidence: medium
    api_base_url: https://llm-serp.jina.ai
    positive_evidence: Current product page says the API is free and optional keys raise limits without being charged.
    negative_evidence: Two anonymous calls returned HTTP 200 with empty results and zero token usage; no numeric public limit is stated.
    recommendation: Require a non-empty successful response before moving to current providers.
    sources:
      - {type: official_product, url: https://jina.ai/api-dashboard/llm-serp/}
      - {type: company_press_release, url: https://jina.ai/news/llm-as-serp-search-engine-result-pages-from-large-language-models/, published_at: "2025-02-27"}

  - id: neurlap
    name: Neurlap
    website: https://neurlap.ai/
    status: early_access_waitlist
    confidence: low
    free_kind: contributor_credits
    condition: Users share GPU capacity to earn credits.
    missing: [public_signup, stable_catalog, quotas, complete_terms]
    source: {type: official_product, url: https://neurlap.ai/}

historical_or_excluded:
  - id: aimlapi
    name: AI/ML API
    status: free_tier_paused
    confidence: high
    reason: The authoritative current FAQ says the Free Tier is paused; live aliases containing :free do not override that statement.
    source: {type: official_docs, title: Free Tier FAQ, url: https://docs.aimlapi.com/faq/free-tier.md}

  - id: runpod
    name: Runpod Serverless and Public Endpoints
    status: prepaid_only
    confidence: high
    reason: Normal inference requires a positive balance; selective startup credits are recorded separately as conditional.
    sources:
      - {type: official_docs, title: Billing, url: https://docs.runpod.io/accounts-billing/billing}
      - {type: official_docs, title: Public endpoints quickstart, url: https://docs.runpod.io/public-endpoints/quickstart}

  - id: digitalocean_gradient_inference
    name: DigitalOcean Gradient serverless inference
    status: prepaid_only
    confidence: high
    reason: A separate prepaid inference balance is mandatory; the free router preview does not make routed model inference free.
    sources:
      - {type: official_pricing, url: https://docs.digitalocean.com/products/gradient-platform/details/pricing/}
      - {type: official_api_reference, url: https://docs.digitalocean.com/reference/api/reference/serverless-inference/}

  - id: fal_ai_general_api
    name: fal Model APIs
    status: prepaid_only
    confidence: high
    reason: Sandbox coupons cannot fund API or Workflow calls; only the application-based Builder Grant is potentially free.
    sources:
      - {type: official_pricing, url: https://fal.ai/docs/documentation/model-apis/pricing}
      - {type: official_docs, url: https://fal.ai/docs/documentation/model-apis/sandbox}

  - id: featherless_ai
    name: Featherless AI
    status: paid_subscription_only
    confidence: high
    reason: The lowest current plan is paid and monthly credits belong to paid subscriptions.
    sources:
      - {type: official_plans, url: https://featherless.ai/docs/plans}
      - {type: official_pricing, url: https://featherless.ai/docs/request-pricing-and-credits}

  - id: inference_net
    name: Inference.net
    status: free_gateway_not_free_inference
    confidence: high
    reason: The free plan covers gateway requests and tracing; all live model entries have positive inference prices.
    sources:
      - {type: official_pricing, url: https://inference.net/pricing/}
      - {type: live_catalog, url: https://api.inference.net/v1/models}

  - id: venice_api
    name: Venice API
    status: free_web_chat_but_paid_api
    confidence: high
    reason: API calls require spendable DIEM, bundled, or USD balance; the free web-chat plan does not fund API inference.
    sources:
      - {type: official_pricing, url: https://venice.ai/pricing}
      - {type: official_docs, title: Generating an API key, url: https://docs.venice.ai/guides/getting-started/generating-api-key.md}

  - id: infermatic
    name: Infermatic
    status: free_ui_only
    confidence: high
    reason: The $0 plan explicitly has no API access; API access starts on a paid plan.
    source: {type: official_pricing, url: https://infermatic.ai/pricing/}

  - id: modular_model_api
    name: Modular shared Model API
    status: paid_hosted_api
    confidence: high
    reason: Free-forever language refers to self-hosted MAX; hosted shared endpoints have positive token prices.
    source: {type: official_pricing, url: https://www.modular.com/pricing}

  - id: lambda_inference_api
    name: Lambda hosted Inference API
    status: winding_down_paid_service
    confidence: high
    reason: The hosted API is winding down and the current alternative is paid GPU instances.
    source: {type: official_product, url: https://lambda.ai/inference}

  - id: avian_api
    name: Avian API
    status: prepaid_only
    confidence: high
    reason: Current models have positive prices and credits must be prepaid; Start Free is signup wording, not a compute allowance.
    source: {type: official_pricing, url: https://api.avian.io/pricing/}

  - id: nextbit
    name: NextBit
    status: prepaid_only
    confidence: medium_high
    reason: Current docs and terms describe prepaid positive-price API usage with no quantified automatic free allowance.
    source: {type: official_docs, url: https://www.nextbit256.com/docs}

  - id: darkbloom
    name: Darkbloom
    status: paid_public_alpha
    confidence: medium_high
    reason: Public alpha access is evaluation-only, but every displayed model has a positive token price.
    source: {type: official_product, url: https://www.darkbloom.dev/}

  - id: wafer
    name: Wafer
    status: paid_subscription_or_inference
    confidence: high
    reason: Current terms describe paid Wafer Pass and inference; third-party claims of a free route lack first-party support.
    source: {type: official_terms, url: https://www.wafer.ai/terms}

  - id: liquid_ai_hosted_api
    name: Liquid AI hosted API
    status: self_hosted_weights_only
    confidence: high
    reason: Current free offer covers downloadable/on-device models and tooling, not operator-funded remote inference.
    sources:
      - {type: official_pricing, url: https://www.liquid.ai/pricing}
      - {type: official_pricing, title: LEAP, url: https://leap.liquid.ai/pricing}

  - id: petals_public_chat
    name: Petals public chat endpoint
    status: currently_nonfunctional
    confidence: high
    reason: A live call returned MissingBlocksError because no peers held the required model blocks.
    source: {type: official_repository, url: https://github.com/petals-infra/chat.petals.dev}

  - id: webinfer_resource_pool
    name: WebInfer resource pool
    status: insufficiently_documented
    confidence: low
    reason: No stable gateway base URL, quota, catalog endpoint, signup flow, or complete terms are published.
    sources:
      - {type: official_product, url: https://webllm.org/providers/resource-pool}
      - {type: official_product, url: https://webinfer.com/providers}

  - id: cscs_inference
    name: CSCS inference service
    status: project_accounted_not_free
    confidence: high
    reason: Non-SwissAI use consumes project credits converted from allocated node-hours at a published CHF rate.
    source: {type: official_docs, url: https://docs.cscs.ch/services/inference/api/}

  - id: isambard_ai_inference
    name: Isambard-AI inference service
    status: not_yet_available
    confidence: high
    reason: Research allocations exist, but the July 2026 update says the shared inference service is still being prepared.
    sources:
      - {type: official_announcement, url: https://www.bristol.ac.uk/research/centres/bristol-supercomputing/articles/2026/one-year-of-isambard-ai.html}

  - id: alia_spain
    name: ALIA Spain
    status: open_models_no_public_api
    confidence: high
    reason: Publishes open models and resources, but no stable public developer inference endpoint with current auth and quotas.
    source: {type: official_product, url: https://alia.gob.es/eng}

  - id: falcon_tii
    name: Falcon / Technology Innovation Institute
    status: weights_and_playground_only
    confidence: high
    reason: Free weights and a playground are available, but no stable documented public inference API with quotas and auth.
    source: {type: official_product, url: https://falconllm.tii.ae/index.html}

  - id: ai2_playground
    name: Ai2 Playground
    status: web_playground_only
    confidence: high
    reason: Ai2 directs API users to external paid routes; its playground is not a documented general developer API.
    source: {type: official_docs, url: https://docs.allenai.org/quick_start/apis}

  - id: llm_jp
    name: LLM-jp
    status: weights_and_tooling_only
    confidence: high
    reason: Publishes weights, corpora, and tooling rather than a stable hosted inference API.
    source: {type: official_product, url: https://llm-jp.nii.ac.jp/release/}

  - id: opengradient
    name: OpenGradient
    status: paid_x402_inference
    confidence: high
    reason: Hosted inference uses x402 or OPG payment rather than a free tier.
    source: {type: official_docs, url: https://docs.opengradient.ai/about/}

  - id: swarmllm
    name: SwarmLLM
    status: self_hosted_only
    confidence: high
    reason: Peer-to-peer software without an operator-hosted public endpoint.
    source: {type: official_repository, url: https://github.com/enapt/SwarmLLM}

  - id: freellmapi_co
    name: FreeLLMAPI.co
    status: self_hosted_byok_router
    confidence: high
    reason: Users self-host the router and supply upstream provider keys.
    source: {type: official_product, url: https://freellmapi.co/}

  - id: krutrim_cloud
    name: Krutrim Cloud AI Studio
    status: prepaid_only
    confidence: high
    reason: Current AI Studio requires purchased credits and has no verified automatic free allowance.
    source: {type: official_docs, url: https://docs.cloud.olakrutrim.com/basics/ai-studio/billing-for-ai-studio}

  - id: duckk_africa
    name: Duckk Africa
    status: waitlist_without_verified_free_terms
    confidence: low
    reason: OpenAI-compatible marketing and a waitlist exist, but no validated free quota or public production endpoint.
    source: {type: official_product, url: https://www.duckk.org/}

  - id: national_university_of_singapore
    name: National University of Singapore generic AI services (unrelated name match)
    status: not_a_public_general_inference_api
    requested_name_resolution: The user confirmed the requested provider was Nous Research; this record remains only to document the unrelated name-match check.
    current_findings:
      - NUS ChatGPT Edu is institution-restricted to students, faculty, and staff beginning 2026-08-31.
      - AI Sense Maker is an NUS web application, not a general public inference API.
    sources:
      - {type: official_announcement, title: NUS and OpenAI strategic collaboration, url: https://news.nus.edu.sg/nus-powers-education-research-and-administration-to-new-heights-with-ai-through-a-strategic-collaboration-with-openai/, published_at: "2026-08-11"}
      - {type: official_research_site, title: NUS AI Institute, url: https://ai.nus.edu.sg/research/}

  - id: github_models
    name: GitHub Models
    status: retired
    confidence: high
    retired_on: "2026-07-30"
    retired_components: [playground, model_catalog, inference_api, byok]
    historical_api_base_url: https://models.github.ai/inference
    history:
      - {date: "2024-10-29", event: Public preview with rate-limited free model inference}
      - {date: "2025-06-24", event: Paid usage beyond free limits added}
      - {date: "2026-06-16", event: Closed to new customers}
      - {date: "2026-07-30", event: Fully retired}
    sources:
      - {type: official_retirement, title: Full retirement announcement, url: https://github.blog/changelog/2026-07-01-github-models-is-being-fully-retired-on-july-30-2026/}
      - {type: official_announcement, title: Public preview, url: https://github.blog/changelog/2024-10-29-github-models-is-now-available-in-public-preview/}

  - id: chutes
    name: Chutes
    status: paid_only
    confidence: high
    api_base_url: https://llm.chutes.ai/v1
    historical_offer:
      requests_per_day: 200
      tee_access_removed: "2026-02-27"
      non_tee_free_access_ended: "2026-03-15"
      replacement: One month of Base or $5 credit.
    sustainability_explanation: Chutes reported removing about 20B daily tokens of sponsored OpenRouter traffic and 10B daily tokens from its direct free quota; estimated direct cost was about $6/user/month across roughly 11,000 users.
    sources:
      - {type: official_pricing, title: Current pricing, url: https://chutes.ai/pricing}
      - {type: official_announcement, title: February community announcement, url: https://chutes.ai/news/community-announcement-february}
      - {type: official_postmortem, title: Building a sustainable inference platform, url: https://chutes.ai/news/from-volume-to-value-building-a-sustainable-ai-inference-platform-2}

  - id: together_ai
    name: Together AI
    status: prepaid_only
    confidence: high
    reason: Official billing docs say there are no free trials; API access requires at least a $5 credit purchase.
    sources:
      - {type: official_docs, title: Billing and credits, url: https://docs.together.ai/docs/billing-credits}
      - {type: official_pricing, title: Inference pricing, url: https://docs.together.ai/docs/inference/pricing}

  - id: deepinfra
    name: DeepInfra
    status: paid_only
    confidence: medium_high
    reason: The current public catalog has positive per-token or per-execution prices and no permanent recurring free quota.
    sources:
      - {type: official_pricing, url: https://deepinfra.com/pricing}
      - {type: official_api_reference, url: https://docs.deepinfra.com/api-reference/introduction}

  - id: qwen_code_oauth
    name: Qwen OAuth free access for Qwen Code
    status: retired
    retired_on: "2026-04-15"
    note: Model Studio/Qwen Cloud new-user quotas remain separately recorded as a current trial.
    source: {type: official_docs, url: https://qwenlm.github.io/qwen-code-docs/en/users/configuration/auth/}

  - id: public_ai_old_blanket_free
    name: Public AI old blanket-free offer
    status: superseded
    old_claim_date: "2025-09-17"
    old_claim: Free of charge at the time of writing.
    current_reality: Starter credit followed by positive per-token wallet billing.
    source: {type: historical_announcement, url: https://huggingface.co/blog/inference-providers-publicai}

  - id: huggingface_spaces_as_a_class
    name: Hugging Face Spaces / Gradio demos
    status: excluded_provider_class
    reason: Many Spaces expose callable Gradio endpoints, but there is no shared quota, catalog guarantee, authentication policy, or availability commitment; individual Spaces may sleep or disappear.
    inclusion_rule: Include a Space only when its owner documents a public API, quota, access policy, and continuing availability.

  - id: self_hosted_open_weights
    name: Self-hosted open-weight models
    status: out_of_scope
    reason: The model weights may be free, but remote hosted compute is not being provided.
