{
  "id": 6263,
  "url": "https://arxiv.org/abs/2604.05591v1",
  "title": "AI-Driven Modular Services for Accessible Multilingual Education in Immersive Extended Reality Settings: Integrating Speech Processing, Translation, and Sign Language Rendering",
  "summary": "This work introduces a modular platform that brings together six AI services, automatic speech recognition via OpenAI Whisper, multilingual translation through Meta NLLB, speech synthesis using AWS Polly, emotion classification with RoBERTa, dialogue summarisation via flan t5 base samsum, and International Sign (IS) rendering through Google MediaPipe. A corpus of IS gesture recordings was processed to derive hand landmark coordinates, which were subsequently mapped onto three dimensional avatar ",
  "authors": "N. D. Tantaroudas, A. J. McCracken, I. Karachalios, E. Papatheou",
  "category": "research",
  "topics": "children-education",
  "orgs": "openai,google,meta,amazon",
  "regions": null,
  "published_at": "2026-04-07T08:35:53.000Z",
  "fetched_at": "2026-07-14T16:32:24.287Z",
  "source_slug": "arxiv-ethics",
  "source_name": "arXiv",
  "source_homepage": "https://arxiv.org",
  "ethics_ai_record_url": "https://ethics.ai/record/6263",
  "original_url": "https://arxiv.org/abs/2604.05591v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}