{
  "id": 1519,
  "url": "https://arxiv.org/abs/2607.11849",
  "title": "AdvancedMathBench: A Benchmark Suite for Advanced Mathematical Proof Generation and Verification",
  "summary": "Large language models (LLMs) have achieved remarkable performance on high-school and olympiad-style mathematics, yet their capabilities on advanced mathematics remain poorly understood. Existing benchmarks, however, fall short in both scope and evaluation granularity: they provide limited disciplinary coverage and often rely on final-answer correctness or coarse judgments, leaving the validity of the reasoning process inadequately assessed. To bridge this gap, we introduce AdvancedMathBench, a b",
  "authors": "Lingkai Kong, Zijian Wu, Yuzhe Gu, Haiteng Zhao, Wenyong Huang, Shuang Sun",
  "category": "research",
  "topics": "children-education",
  "orgs": null,
  "regions": null,
  "published_at": "2026-07-13T13:38:22.000Z",
  "fetched_at": "2026-07-14T16:04:12.223Z",
  "source_slug": "hf-daily",
  "source_name": "HuggingFace Daily Papers",
  "source_homepage": "https://huggingface.co/papers",
  "ethics_ai_record_url": "https://ethics.ai/record/1519",
  "original_url": "https://arxiv.org/abs/2607.11849",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}