{
  "id": 14800,
  "url": "https://arxiv.org/abs/2607.27191v1",
  "title": "Can AI agents conduct open-ended AI research? Early evidence from two case studies",
  "summary": "Forecasts of explosive AI progress hinge on AI agents automating AI research. But evidence on whether agents can carry out open-ended AI research is thin. Current evaluations either test agents on narrow, verifiable tasks, which excludes open-ended research, or submit AI-generated papers to blind peer review, which is overstretched, stochastic, and suffers from poor review quality. We introduce a third way to measure progress towards AI R\\&D automation. An agent takes on the central, open-ended",
  "authors": "Peter Kirgis, Sayash Kapoor, Andrew Schwartz, Stephan Rabanser, David Africa, Konstantinos Voudouris, Viet Nguyen, Toby Pilditch, Magda Dubois, Harry Coppock, Cozmin Ududec, Nitya Nadgir, Matilda Orona, Tilman Bayer, Derrick Chan-Sew, Yue Ling, Abhishek Shetty, Helen Toner, Gillian Hadfield, Seth Lazar, Steve Newman, Shoshannah Tekofsky, Rishi Bommasani, Arvind Narayanan",
  "category": "research",
  "topics": "jobs-economy,agents-autonomy",
  "orgs": null,
  "regions": null,
  "published_at": "2026-07-29T17:57:19.000Z",
  "fetched_at": "2026-07-30T05:10:24.387Z",
  "source_slug": "x-arxiv-cs-ai",
  "source_name": "arXiv cs.AI",
  "source_homepage": "https://arxiv.org/list/cs.AI/recent",
  "ethics_ai_record_url": "https://ethics.ai/record/14800",
  "original_url": "https://arxiv.org/abs/2607.27191v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}