{
  "schema_version": "1.0",
  "id": "s8:https://www.youtube.com/watch?v=O72p-rBb2bA",
  "slug": "evals-driven-development-for-a-mental-health-ai-coach-akele-reed-dave-re-1l6vc33",
  "url": "https://feed7.dev/p/evals-driven-development-for-a-mental-health-ai-coach-akele-reed-dave-re-1l6vc33",
  "title": "Evals-Driven Development for a Mental Health AI Coach — Akele Reed & Dave Revere, SonderMind",
  "why_included": "SonderMind turns clinician-reviewed failures into release-gating evals, keeping mental-health guardrails modular and testing false positives, false negatives, category, and timing.",
  "summary": "SonderMind built **separate input and output guardrail LLMs** around its coaching model instead of relying on general-purpose moderation. Flagged conversations are reviewed by clinicians, normalized into **typed evals**, and committed as **CI release gates**.",
  "practical_implication": "For high-stakes agents, trace ambiguous failures and make domain experts responsible for defining acceptable behavior. Evaluate whether the correct category fired at the correct conversational moment, including both false negatives and harmful over-triggering.",
  "agent_context": "SonderMind built **separate input and output guardrail LLMs** around its coaching model instead of relying on general-purpose moderation. Flagged conversations are reviewed by clinicians, normalized into **typed evals**, and committed as **CI release gates**.\n\nFor high-stakes agents, trace ambiguous failures and make domain experts responsible for defining acceptable behavior. Evaluate whether the correct category fired at the correct conversational moment, including both false negatives and harmful over-triggering.\n\nThe talk provides an engineering process, not measured safety rates or a claim of perfect detection. Clinical language can remain ambiguous, separate judge calls add latency and cost, and benchmark optimization can drift away from the people the tests are meant to protect.",
  "source": {
    "name": "AI Engineer",
    "url": "https://www.youtube.com/watch?v=O72p-rBb2bA",
    "published_at": "2026-07-25T23:00:36.000Z"
  },
  "source_class": "video",
  "content_type": "Video",
  "layer": "benchmark",
  "domains": [],
  "topics": [
    "agent-evals",
    "agent-reliability"
  ],
  "verification": {
    "status": "source_linked",
    "label": "Source Linked",
    "method": "source_feed",
    "verified_at": null
  },
  "uncertainty": [
    "The talk provides an engineering process, not measured safety rates or a claim of perfect detection. Clinical language can remain ambiguous, separate judge calls add latency and cost, and benchmark optimization can drift away from the people the tests are meant to protect."
  ],
  "lifecycle": "Current",
  "published_at": "2026-07-25T23:00:36.000Z",
  "modified_at": "2026-07-25T23:00:36.000Z",
  "supersedes": [],
  "expires_at": null,
  "formats": {
    "html": "https://feed7.dev/p/evals-driven-development-for-a-mental-health-ai-coach-akele-reed-dave-re-1l6vc33",
    "json": "https://feed7.dev/p/evals-driven-development-for-a-mental-health-ai-coach-akele-reed-dave-re-1l6vc33.json",
    "markdown": "https://feed7.dev/p/evals-driven-development-for-a-mental-health-ai-coach-akele-reed-dave-re-1l6vc33.md"
  }
}