{
 "license": "CC-BY-4.0",
 "attribution": "Fide AI, Agentic Cyber Explorer",
 "url": "https://agentic-cyber-explorer.pages.dev/findings/trusted-editing-beats-monitoring-in-toy-setting/",
 "asOf": "2026-09-26",
 "id": "trusted-editing-beats-monitoring-in-toy-setting",
 "claim": "In an early AI-control experiment, trusted editing gave a better safety-usefulness tradeoff than untrusted monitoring.",
 "evidenceKind": "measured",
 "scope": "A 2023 code-backdoor setting with models of that time.",
 "topics": [
  "monitoring-and-control"
 ],
 "atlas": [
  "monitor"
 ],
 "evidence": [
  {
   "event": "redwood-ai-control-paper-2023"
  }
 ],
 "relations": [],
 "statusHistory": [
  {
   "status": "reported",
   "on": "2023-12-12",
   "why": "Redwood Research AI control paper.",
   "event": "redwood-ai-control-paper-2023",
   "kind": "evidence"
  }
 ],
 "halfLifeDays": 730,
 "wouldChange": "A replication with current models and agentic tasks.",
 "fideQuestions": [],
 "methods": [
  "ai-monitoring",
  "monitor-evasion"
 ],
 "review": "assistant-drafted",
 "addedOn": "2026-09-25"
}