{
  "slug": "braintrust",
  "name": "Braintrust",
  "url": "https://www.braintrust.dev",
  "cat": "ai",
  "group": "ai",
  "tagline": "LLM eval and experiment platform",
  "price": "Starter free ($10/mo model credits, 1GB data, 10k scores, 14-day retention), Pro $249/mo ($249 credits, 5GB, 50k scores), Enterprise custom. Overage $3-4/GB and $1.50-$2.50 per 1k scores.",
  "tier": "mixed",
  "pick": false,
  "what": "An eval-first <b>evaluation</b> platform: datasets, scoring functions, experiment comparison, a prompt playground and production logging in one place. It aims to answer 'did this prompt change make things better or worse' with numbers, rather than only showing traces.",
  "why": "Turns prompt tweaking from vibes into a regression suite — pin a dataset, run the eval, see which model or prompt scores higher before you ship. Included model credits mean the graders themselves are covered.",
  "warn": "Billing on processed data volume and score counts punishes verbose traces — logging whole documents in every span pushes you past 1GB fast. Free retention is only 14 days.",
  "install": "npm i braintrust (or pip install braintrust)",
  "category": {
    "key": "ai",
    "name": "Model APIs, routers & observability",
    "desc": "Providers, gateways, inference hosts and eval layers — constantly conflated, listed separately here.",
    "group": "ai",
    "groupName": "Put AI inside it"
  },
  "freshness": null,
  "health": {
    "ok": true,
    "status_code": 200,
    "final_url": "https://www.braintrust.dev/",
    "checked_at": "2026-08-18T18:09:22.039+00:00"
  },
  "corrected_at": null,
  "upstream_changed_at": null
}