{
  "id": 3413,
  "url": "https://arxiv.org/abs/2606.00341v1",
  "title": "ROGUE: Misaligned Agent Behavior Arising from Ordinary Computer Use",
  "summary": "As AI agents are increasingly deployed in real personal and corporate settings (email accounts, development workflows, company databases, etc.), safety considerations surrounding these agents become paramount. Although much work has focused on agent safety in the presence of an adversary, we show that agents can exhibit misaligned behavior even in benign settings, taking unsafe actions when those actions are instrumental to task completion. We study this failure mode through the lens of corrigib",
  "authors": "Jeremy Tien, Abishek Anand, Yu-Rou Tuan, Yuchen Shen, J. Zico Kolter, Aran Nayebi",
  "category": "research",
  "topics": "safety-alignment,agents-autonomy",
  "orgs": null,
  "regions": null,
  "published_at": "2026-05-29T20:29:35.000Z",
  "fetched_at": "2026-07-14T16:30:14.369Z",
  "source_slug": "arxiv-ethics",
  "source_name": "arXiv",
  "source_homepage": "https://arxiv.org",
  "ethics_ai_record_url": "https://ethics.ai/record/3413",
  "original_url": "https://arxiv.org/abs/2606.00341v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}