{
  "slug": "fal-ai",
  "name": "Fal.ai",
  "url": "https://fal.ai",
  "cat": "ai",
  "group": "ai",
  "tagline": "Generative media inference host",
  "price": "Per output or per second for media models: images ~$0.02-$0.04 at 1MP, video $0.05-$0.40 per second. Serverless GPU rates $1.10-$8.50/hr with committed-use discounts.",
  "tier": "paid",
  "pick": false,
  "what": "An <b>inference host</b> specialised in generative media — diffusion image models, video models such as Veo, Kling and Wan, plus audio. It runs its own optimised inference stack, so the same open model is usually faster here than on a general-purpose GPU cloud.",
  "why": "If your app generates images or video, fal is the media equivalent of Groq — low latency, streaming previews, and a clean JS and Python client. Per-image pricing makes unit economics easy to reason about.",
  "warn": "Video generation gets expensive fast — at $0.40 per second a five-second clip is $2, so an unthrottled public demo can burn hundreds of dollars in an afternoon.",
  "install": "npm i @fal-ai/client (or pip install fal-client)",
  "category": {
    "key": "ai",
    "name": "Model APIs, routers & observability",
    "desc": "Providers, gateways, inference hosts and eval layers — constantly conflated, listed separately here.",
    "group": "ai",
    "groupName": "Put AI inside it"
  },
  "freshness": null,
  "health": {
    "ok": true,
    "status_code": 200,
    "final_url": "https://fal.ai/",
    "checked_at": "2026-08-18T18:09:22.039+00:00"
  },
  "corrected_at": null,
  "upstream_changed_at": null
}