{
  "id": 17427,
  "url": "https://arxiv.org/abs/2608.05138",
  "title": "Teaching Nemotron Greek: Mining a Corpus, Adapting Retrieval, and Grounding Generation for Modern Greek across Specialist Domains",
  "summary": "Modern Greek is absent from NVIDIA's Nemotron retrieval models and from major multilingual retrieval benchmarks, despite being important for retrieval-augmented generation (RAG) in legal, energy, financial, and medical applications. We present an end-to-end adaptation of the Nemotron retrieval stack for Modern Greek, including corpus mining, synthetic supervision, retrieval model training, reranker adaptation, reader fine-tuning, and a new benchmark called HERA. Our study shows that a parameter-",
  "authors": "Ayoub Kirouane, Christos Petrocheilos",
  "category": "research",
  "topics": "healthcare,environment",
  "orgs": "nvidia",
  "regions": null,
  "published_at": "2026-08-04T20:00:00.000Z",
  "fetched_at": "2026-08-08T05:10:34.355Z",
  "source_slug": "hf-daily",
  "source_name": "HuggingFace Daily Papers",
  "source_homepage": "https://huggingface.co/papers",
  "ethics_ai_record_url": "https://ethics.ai/record/17427",
  "original_url": "https://arxiv.org/abs/2608.05138",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}