{
  "id": 14428,
  "url": "https://arxiv.org/abs/2607.25904v1",
  "title": "Interactive Reward Agent: GUI Task Evaluation via Environment-State Verification",
  "summary": "Graphical user interface task evaluation aims to determine whether a GUI agent has successfully completed a user instruction. Automated GUI task evaluation has received increasing attention because the evaluation results can serve as reward signals for both test-time scaling and post-training. However, reliable GUI task evaluation remains challenging because the judgments often require access to environment states, such as system configurations, file data, and application settings, beyond the sc",
  "authors": "Chenrui Shi, Yuwei Wu, Yang Liu, Ruining Feng, Zirui Shang, Zhi Gao, Lifeng Fan, Che Sun",
  "category": "research",
  "topics": "agents-autonomy,environment",
  "orgs": null,
  "regions": null,
  "published_at": "2026-07-28T16:01:38.000Z",
  "fetched_at": "2026-07-29T05:10:12.205Z",
  "source_slug": "x-arxiv-cs-ai",
  "source_name": "arXiv cs.AI",
  "source_homepage": "https://arxiv.org/list/cs.AI/recent",
  "ethics_ai_record_url": "https://ethics.ai/record/14428",
  "original_url": "https://arxiv.org/abs/2607.25904v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}