{
  "id": "anthropic-sitemap:research:bloom",
  "type": "article",
  "title": "Introducing Bloom: an open source tool for automated behavioral evaluations",
  "abstract": "_We're releasing Bloom, an open source agentic framework for generating behavioral evaluations of frontier AI models. Bloom takes a researcher-specified behavior and quantifies its frequency and severity across automatically generated scenarios. Bloom's evaluations correlate strongly with our hand-labeled judgments and we find they reliably separate baseline models from intentionally misaligned ones. As examples of this, we release benchmark results for four alignment relevant behaviors on 16 models. Bloom is available here._",
  "issued": {
    "date-parts": [
      [
        2025,
        12,
        19
      ]
    ]
  },
  "URL": "https://www.anthropic.com/research/bloom",
  "publisher": "Anthropic",
  "source": "vendor/anthropic-sitemap/research/bloom.md"
}
