{
  "id": 6460,
  "url": "https://arxiv.org/abs/2604.02389v1",
  "title": "Audio Spatially-Guided Fusion for Audio-Visual Navigation",
  "summary": "Audio-visual Navigation refers to an agent utilizing visual and auditory information in complex 3D environments to accomplish target localization and path planning, thereby achieving autonomous navigation. The core challenge of this task lies in the following: how the agent can break free from the dependence on training data and achieve autonomous navigation with good generalization performance when facing changes in environments and sound sources. To address this challenge, we propose an Audio ",
  "authors": "Xinyu Zhou, Yinfeng Yu",
  "category": "research",
  "topics": "agents-autonomy,transparency,environment",
  "orgs": null,
  "regions": null,
  "published_at": "2026-04-02T07:15:17.000Z",
  "fetched_at": "2026-07-14T16:32:28.613Z",
  "source_slug": "arxiv-ethics",
  "source_name": "arXiv",
  "source_homepage": "https://arxiv.org",
  "ethics_ai_record_url": "https://ethics.ai/record/6460",
  "original_url": "https://arxiv.org/abs/2604.02389v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}