{"slug":"evalview","title":"evalview","description":"Behavior regression testing for AI agents. EvalView detects when your\nagent's behavior drifts e.g changed tool calls, different outputs, or degraded quality even when traditional tests still pass. Snapshot baselines, check for regressions, and auto-heal flaky results, all from inside Claude Code.","developer":null,"category":null,"reported_installs":null,"active":true,"first_seen_at":"2026-10-03T22:13:33.942Z","last_seen_at":"2026-10-08T18:00:00.447Z","last_changed_at":"2026-10-03T22:13:33.942Z","platform":"claude","sources":["Community"],"works_with":["Claude Code"],"marketplace_url":null,"repository_url":"https://github.com/hidai25/eval-view.git","web_data":{},"github_data":{},"catalog_sources":{"community":[{"name":"evalview","source":{"sha":"afbc7f6c715d7cd1ab1ea71aefd5e26a9bbcd767","url":"https://github.com/hidai25/eval-view.git","source":"url"},"homepage":"https://evalview.com","description":"Behavior regression testing for AI agents. EvalView detects when your\nagent's behavior drifts e.g changed tool calls, different outputs, or degraded quality even when traditional tests still pass. Snapshot baselines, check for regressions, and auto-heal flaky results, all from inside Claude Code."}]}}