{
  "id": 10915,
  "url": "https://arxiv.org/abs/2607.13621v1",
  "title": "UESF-Bench: Benchmarking and Probing for Unified Embodied Seeking and Following",
  "summary": "Language-guided human following is an important capability for embodied agents, but existing benchmarks typically assume that the target person is visible at the start of an episode. This setting simplifies the problem and overlooks a more realistic requirement: an agent often needs to first find a language-described target and then persistently follow that target in a dynamic environment. While recent work has started to study human search, existing settings are typically evaluated in task-spec",
  "authors": "Kun Yu, Jianhua Yang, Yixiang Chen, Changwei Wang, Hongyuan Yu, Yan Huang, Fushuo Huo, Ya Jing, Zhumin Chen, Keji He",
  "category": "research",
  "topics": "agents-autonomy,environment",
  "orgs": null,
  "regions": null,
  "published_at": "2026-07-15T09:09:25.000Z",
  "fetched_at": "2026-07-16T05:10:56.605Z",
  "source_slug": "x-arxiv-cs-ai",
  "source_name": "arXiv cs.AI",
  "source_homepage": "https://arxiv.org/list/cs.AI/recent",
  "ethics_ai_record_url": "https://ethics.ai/record/10915",
  "original_url": "https://arxiv.org/abs/2607.13621v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}