{
  "id": 15254,
  "url": "https://arxiv.org/abs/2607.27968v1",
  "title": "Beyond Binary Rewards: A Comparative Study of Reward Design for Reinforcement Unlearning",
  "summary": "Machine unlearning seeks to selectively remove specific knowledge from trained language models without full retraining, a growing necessity under privacy regulations such as GDPR and the EU AI Act. Recent work has reformulated unlearning as a Reinforcement Learning with Verifiable Rewards (RLVR) problem, where models are optimized against verifiable rewards computed directly from their outputs. However, existing methods rely on sparse binary rewards that provide minimal learning signal, indicati",
  "authors": "Efstratios Zaradoukas, Davide Gabrielli, Bardh Prenkaj, Gjergji Kasneci",
  "category": "research",
  "topics": "regulation,privacy-surveillance",
  "orgs": null,
  "regions": "eu",
  "published_at": "2026-07-30T10:15:59.000Z",
  "fetched_at": "2026-07-31T05:10:57.675Z",
  "source_slug": "arxiv-cslg",
  "source_name": "arXiv cs.LG",
  "source_homepage": "https://arxiv.org/list/cs.LG/recent",
  "ethics_ai_record_url": "https://ethics.ai/record/15254",
  "original_url": "https://arxiv.org/abs/2607.27968v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}