{
  "count": 50,
  "items": [
    {
      "id": 19474,
      "url": "https://arxiv.org/abs/2608.13190v1",
      "title": "ProME: Prototype-Margin Environments with Repair-Aware Selection for Group-Robust Learning",
      "summary": "Group-robust learning is crucial for maintaining accuracy on rare subpopulations when training-group labels are unavailable. However, existing methods often infer environments from a separate reference model and select representations before fitting the classifier used at deployment, leaving both decisions misaligned with the deployed predictor. In this work, we formulate group robustness without training-group labels as the endogenous environments with repair-aware selection (ERAS) problem, and",
      "authors": "Qianqian Wang, Yunshan Li, Dawei Huang, Wenwu Gong, Lili Yang",
      "category": "research",
      "topics": "safety-alignment,environment",
      "published_at": "2026-08-13T12:57:41.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19474"
    },
    {
      "id": 19475,
      "url": "https://arxiv.org/abs/2608.12957v1",
      "title": "I-SDPO: Instance-Level Adaptive Self-Distillation Policy Optimization",
      "summary": "Group Relative Policy Optimization (GRPO) learns from reward differences within a rollout group, but receives no useful relative signal when every sampled response is incorrect. Privileged self-distillation can fill this gap with dense token supervision, yet applying it throughout training creates a different failure mode: the teacher is a biased, low-variance surrogate for the reward objective, so persistent imitation can oppose reward-improving updates after the policy becomes capable of produ",
      "authors": "Yubo Zhang, Xinhong Ma, Zezhong Tan, Ziqiang Dong",
      "category": "research",
      "topics": "bias-fairness,regulation",
      "published_at": "2026-08-13T08:37:24.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19475"
    },
    {
      "id": 19476,
      "url": "https://arxiv.org/abs/2608.12489v1",
      "title": "When Can You Trust Offline Evaluation of Equal-Cost Top-k Allocation? A Controlled, Reproducible Benchmark and Practitioner's Guide",
      "summary": "Organizations decide whom to treat under a budget and want to know what a targeting rule would have earned before deploying it. Off-policy evaluation promises this from logged data, but the deployable rule is a deterministic top-k policy: it removes all averaging over actions, so weak overlap hits the estimate directly. We benchmark six estimators across five datasets and two known-effect sweeps, and validate the mechanisms against a non-simulated paired reference. First, weak overlap is governe",
      "authors": "Binshuang Li",
      "category": "research",
      "topics": "regulation",
      "published_at": "2026-08-12T18:10:10.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19476"
    },
    {
      "id": 19477,
      "url": "https://arxiv.org/abs/2608.12477v1",
      "title": "Learning Under Treatment-Induced Label Indeterminacy with Expert Annotations of Counterfactual Outcomes: A Case Study in Neurological Prognostication",
      "summary": "Clinical prediction models are often developed as if the outcome of interest were cleanly observed for every patient. This assumption fails when treatment decisions make the clinically relevant outcome permanently unobservable. As a case study of this problem, we consider post-cardiac-arrest neurological prognostication using a cohort of 2,497 patients, including 1,429 patients whose outcomes were rendered indeterminate by treatment decisions. These patients with indeterminate outcomes were revi",
      "authors": "Xiaobin Shen, Chloe Y. H. Huang, Jonathan Elmer, George H. Chen",
      "category": "research",
      "topics": "healthcare",
      "published_at": "2026-08-12T18:01:13.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19477"
    },
    {
      "id": 19478,
      "url": "https://arxiv.org/abs/2608.12447v1",
      "title": "Geometric and Behavioral Stratification in Transformer Residual Streams",
      "summary": "Trained transformer models develop privileged bases: coordinate axes whose statistics differ from the rest of the residual stream. But what kind of direction does such a basis select? We investigate the prediction direction, the unembedding direction of the token a model currently predicts, and find that it functions as a content-defined privileged anchor. Measured with respect to this anchor, residual-stream variation is geometrically and behaviorally stratified by proximity to the prediction.",
      "authors": "Nelson Guda",
      "category": "research",
      "topics": "finance-investment",
      "published_at": "2026-08-12T17:42:20.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19478"
    },
    {
      "id": 19479,
      "url": "https://arxiv.org/abs/2608.12444v1",
      "title": "Non-Degenerate Risk Certification for Automated Security Decisions: A Decision-Contract Theory with ATT\\&CK-Aligned Triage as a Worked Instance",
      "summary": "An unconditional risk bound on automated decisions can be satisfied without automating anything, since a selector that never acts drives the bound to zero. We show this is structural: any risk certificate is defined over a decision contract, the inputs a system acts on plus the semantic relation under which an output counts correct, and weakening either hides base-classifier error. We develop a decision-contract theory: an error-conservation law showing error is only reassigned among harmful aut",
      "authors": "Zhenpeng Li",
      "category": "research",
      "topics": "regulation",
      "published_at": "2026-08-12T16:59:19.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19479"
    },
    {
      "id": 19480,
      "url": "https://arxiv.org/abs/2608.12441v1",
      "title": "Dual Spatial-Temporal Attribution: Architecture-Aligned Post-Hoc Explainability for Recurrent Graph Anomaly Detection",
      "summary": "Deep learning detectors for anomalies in dynamic graphs have reached strong accuracy, yet they remain opaque: when an edge is flagged, the analyst receives a score but no reason. This opacity is untenable in the cooperative, regulated information systems where such detectors are deployed, where automated decisions must be auditable and trustworthy. We address this gap for AddGraph, the foundational GCN+GRU framework for edge-level anomaly detection in dynamic graphs, which to our knowledge has n",
      "authors": "Iyad Assaad Nekka, Hamida Seba, Khaled Walid Hidouci, Karima Amrouche",
      "category": "research",
      "topics": "regulation,transparency",
      "published_at": "2026-08-12T15:58:27.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19480"
    },
    {
      "id": 19066,
      "url": "https://arxiv.org/abs/2608.12086v1",
      "title": "Look What the Probes Dragged In! Real-World Chest X-ray Shortcuts in MedCLIP",
      "summary": "Vision-language models, such as contrastive language-image pre-training (CLIP)-based approaches, have reached state-of-the-art (SOTA) results in medical artificial intelligence. However, recent work reveals that CLIP-based models remain vulnerable to shortcuts. We investigate how real-world shortcuts manifest across different layers of the medical CLIP-based model, MedCLIP, and its vision encoder, a frozen ResNet-50. We attach 17 linear classification probes to the intermediate layers of the Res",
      "authors": "Nikolette Pedersen, Regitze Sydendal, Veronika Cheplygina, Théo Sourget",
      "category": "research",
      "topics": "healthcare,finance-investment",
      "published_at": "2026-08-12T14:08:52.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19066"
    },
    {
      "id": 19067,
      "url": "https://arxiv.org/abs/2608.12084v1",
      "title": "NAE: Normalizing AutoEncoder",
      "summary": "We consider the setting of Normalizing flows with approximate inverses, an established paradigm spanning both full-dimensional ($d=D$) and bottleneck ($d<D$) settings, and group these models under the term flow autoencoders. We present a theoretical investigation into their training dynamics and prove that the proposed loss used by existing approaches is suboptimal; specifically, both encoder and decoder surrogates must be optimized in alignment with reconstruction loss. Guided by these insights",
      "authors": "Muhammad Abdur Rafae, Niels Landwehr",
      "category": "research",
      "topics": "safety-alignment,finance-investment",
      "published_at": "2026-08-12T14:07:09.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19067"
    },
    {
      "id": 19481,
      "url": "https://arxiv.org/abs/2608.11698v2",
      "title": "REOPD: Reliability-Adaptive Reward Extrapolation for On-Policy Distillation",
      "summary": "On-policy distillation (OPD) trains a student on its own trajectories under dense token-level supervision from a teacher. Reward-extrapolation methods such as ExOPD amplify the teacher-reference log-likelihood ratio to move beyond direct imitation, but apply a single global coefficient $λ$ to every token. This can drive the student to fit extreme peaks in the implicit reward, causing reward hacking and unstable training, and the optimal $λ$ varies across domains, requiring costly sweeps. We prop",
      "authors": "Yang Sun, Lichao Ma, Houyuan Qin, Yuxin Liu, Hanyang Lu, Yao Zhu et al.",
      "category": "research",
      "topics": "regulation,children-education",
      "published_at": "2026-08-12T06:15:33.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19481"
    },
    {
      "id": 19068,
      "url": "https://arxiv.org/abs/2608.11656v1",
      "title": "Continuous-Latent Predictive Modeling with Semantic Alignment for EEG-Language Foundation Models",
      "summary": "Recent advances in EEG foundation models have demonstrated the potential of large-scale pretraining to enable generalizable neural decoding across subjects, recording environments, and datasets. However, dominant pretraining paradigms face key challenges: masked autoencoding tends to prioritize low-level signal reconstruction over task-relevant semantics, while autoregressive modeling creates a mismatch between continuous neural dynamics and discrete token spaces. To address these challenges, ne",
      "authors": "Myeong-Ju Cho, Hye-Bin Shin, Seo-Hyun Lee, Seong-Whan Lee",
      "category": "research",
      "topics": "safety-alignment,environment",
      "published_at": "2026-08-12T04:54:43.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19068"
    },
    {
      "id": 19069,
      "url": "https://arxiv.org/abs/2608.11638v1",
      "title": "Transferable Above-Ground Biomass (AGB) Estimation Model from Multi-Sensor Data with Sparse Field Calibration",
      "summary": "Spatially continuous quantification of forest above-ground biomass (AGB) is what makes carbon accounting credible and mitigation strategies actionable. While field inventories provide high localized accuracy, they are spatially sparse; conversely, spaceborne LiDAR from the Global Ecosystem Dynamics Investigation (GEDI) offers broad biomass samples but lacks spatial continuity and systematic underestimation of high-biomass forests. This paper presents an operational framework centered on a single",
      "authors": "Pann Thinzar Seint, Bryan Atwood, Subas Chhatkuli",
      "category": "research",
      "topics": "environment,finance-investment",
      "published_at": "2026-08-12T04:36:26.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19069"
    },
    {
      "id": 19070,
      "url": "https://arxiv.org/abs/2608.11560v1",
      "title": "When Offline Evaluation Misleads: A Diagnostic Protocol for Reward and Policy Selection in Delayed-Feedback Contextual Bandits",
      "summary": "Personalizing marketing messages with contextual multi-armed bandits (CMABs) drives real business value, yet the objective that ultimately matters - a downstream conversion - is observed only weeks later, too late to drive online learning. Teams therefore train the bandit on a fast proxy reward, and separately must judge whether a contextual bandit is worth its complexity over sending one best message. Settling both decisions with the usual offline checks - a batch off-policy estimate, a margina",
      "authors": "Sang Su Lee, Vineeth Loganathan, Shishir Dash, Vijay Raghavan",
      "category": "research",
      "topics": "regulation,healthcare",
      "published_at": "2026-08-12T01:53:21.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19070"
    },
    {
      "id": 18693,
      "url": "https://arxiv.org/abs/2608.11167v1",
      "title": "MultiModal Code-Switching: Interleaving Visual Objects into Language for Explicit Object-Level Alignment",
      "summary": "Existing Multimodal Large Language Models (MLLMs) predominantly rely on image-text pairs for modality alignment pretraining, mapping global image representations to long textual descriptions. However, this image-level alignment suffers from referential ambiguity: models struggle to infer the correspondences between multiple visual objects and textual entities from the global representation, leading to data inefficiency and suboptimal semantic grounding. To address this, we propose MultiModal Cod",
      "authors": "Changhao Xiang, Shangyu Xing, Zhen Wu, Jianbing Zhang, Xinyu Dai",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-11T17:28:52.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18693"
    },
    {
      "id": 18694,
      "url": "https://arxiv.org/abs/2608.11034v1",
      "title": "SCOUT: Symmetric Consensus Outlier Detection for Failure Localization in LLM Pre-Training",
      "summary": "In LLM pre-training, synchronization propagates rank-local stalls, slowdowns, and numerical errors into job-wide symptoms, obscuring their origin. Existing diagnosis often relies on in-process monitors that cannot report after the trainer blocks or terminates, or on post-mortem logs that preserve only synchronized symptoms; offline health tests lose the workload and operating conditions that triggered the failure. We present SCOUT, a unified runtime failure-localization framework built on one de",
      "authors": "Zhuang Wang",
      "category": "research",
      "topics": "jobs-economy,healthcare",
      "published_at": "2026-08-11T15:12:14.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18694"
    },
    {
      "id": 18695,
      "url": "https://arxiv.org/abs/2608.10897v1",
      "title": "Partially Observable Learning for Multi-Platform Dispatch Optimization",
      "summary": "Instant delivery platforms have become a critical component of urban logistics, increasingly relying on crowdsourced couriers to fulfill highly dynamic orders. In real-world systems, couriers are not exclusive to a single platform and may concurrently serve multiple platforms, while each platform can only observe its own orders and couriers' interactions due to privacy and operational constraints. This results in a multi-platform dispatch environment with inherent partial observability. However,",
      "authors": "Fengming Yao, Man Luo",
      "category": "research",
      "topics": "privacy-surveillance,environment",
      "published_at": "2026-08-11T13:20:12.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18695"
    },
    {
      "id": 18696,
      "url": "https://arxiv.org/abs/2608.10634v1",
      "title": "IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning",
      "summary": "Model-based reinforcement learning (MBRL), which learns environment dynamics to generate synthetic experience, is a promising approach to sample-efficient decision making. Numerous methods have been developed to improve dynamics prediction and policy optimization for MBRL through uncertainty estimation, model regularization, and conservative value learning. However, these methods typically treat the transition model and critic as monolithic predictors, overlooking the policy-induced data bias. C",
      "authors": "Zefeng Liang, Jie Qiao, Ruichu Cai, Weilin Chen, Zhifeng Hao",
      "category": "research",
      "topics": "bias-fairness,regulation,environment",
      "published_at": "2026-08-11T08:20:02.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18696"
    },
    {
      "id": 18697,
      "url": "https://arxiv.org/abs/2608.10619v1",
      "title": "Pair-Centric Graph Rewiring for Over-Squashing via Optimal Transport-Guided Communication Alignment",
      "summary": "Message-passing neural networks (MPNNs) often struggle when task-relevant information is distributed across distant regions of a graph, since local propagation must compress remote signals through limited structural interfaces. Graph rewiring provides a structural response to over-squashing. Most existing methods rely on edge-level bottleneck scores or graph-level connectivity surrogates. With a limited rewiring budget, the key question is which pairwise communications most need structural suppo",
      "authors": "Yan Wang, Chuan-Xian Ren",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-11T08:05:32.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18697"
    },
    {
      "id": 18698,
      "url": "https://arxiv.org/abs/2608.10406v1",
      "title": "Post-Calibration Reliability Reranking of Relevance Decisions via Label-wise Monotone Projection",
      "summary": "Web search, product search, and question-answering retrieval systems often assign a relevance label and confidence score to each query-candidate pair. The relevance label describes how well a page, product, or passage matches the query, while the confidence often guides downstream use or fallback decisions. Post-hoc calibration is therefore needed because misaligned confidence can make systems over-trust wrong predictions or unnecessarily defer correct ones. However, calibration mainly aligns co",
      "authors": "Inwoo Tae, Yongjae Lee",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-11T02:51:35.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18698"
    },
    {
      "id": 18699,
      "url": "https://arxiv.org/abs/2608.10316v1",
      "title": "UniMod: Enhancing Multi-Modal Medical Diagnosis through Cross-Modality and Within-Modality Alignment",
      "summary": "Multi-modal learning combining medical images and clinical text is promising for disease diagnosis. However, standard multi-modal training leads to shortcut learning: models exploit the easier modality (e.g., diagnostic cues in text) while neglecting harder-to-learn features (e.g., subtle visual patterns). We propose UniMod, a framework that mitigates shortcut learning by requiring each modality to predict the diagnosis on its own. It supervises image-only, text-only, and multi-modal classificat",
      "authors": "Zijian Gu, Weikai Lin, Shuang Zhou, Zihan Chen, Song Wang",
      "category": "research",
      "topics": "safety-alignment,healthcare",
      "published_at": "2026-08-10T23:39:49.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18699"
    },
    {
      "id": 19071,
      "url": "https://arxiv.org/abs/2608.10300v2",
      "title": "Logit-Boundary Geometric Belief Interfaces and Sparse Sheaf-Enclave Protocols: A Self-Contained Substrate for Secure Network Electronic Health Record (EHR) Interoperability",
      "summary": "Electronic health-record interoperability is a boundary problem: legacy systems, generative models, terminology services, identity systems, and human reviewers may each expose rich internal states, while operational exchange requires a narrow shared interface of typed claims, bounded uncertainty, provenance, and explicit admission or abstention. This paper details a mathematical and engineering architecture for that interface. The organizing idea is the logit boundary: a discovery model may prop",
      "authors": "Alvin Spivey, Yu Huang",
      "category": "research",
      "topics": "healthcare",
      "published_at": "2026-08-10T23:10:53.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/19071"
    },
    {
      "id": 18321,
      "url": "https://arxiv.org/abs/2608.09332v1",
      "title": "Hallucinations and Constraints : Regulating surgical workflow recognition beyond accuracy",
      "summary": "Hallucinations are a major concern for the integration of artificial intelligence into medicine, although less explored in the realm of medical image processing. Unlike problems in natural text understanding and reasoning therewith, determining whether or not predictions derived from biomedical images and signals is less intuitively clear. This article suggests that topological errors could constitute hallucinations in a way that can be more readily measured and thus regulated. Certain of these",
      "authors": "John S. H. Baxter, Pierre Jannin",
      "category": "research",
      "topics": "regulation,healthcare",
      "published_at": "2026-08-10T09:11:48.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18321"
    },
    {
      "id": 18322,
      "url": "https://arxiv.org/abs/2608.09193v1",
      "title": "CPDA: Class-Conditional Path Distribution Alignment for Unsupervised Time-Series Domain Adaptation",
      "summary": "Unsupervised time-series domain adaptation (DA) addresses the challenge of transferring a classifier from a labeled source domain to an unlabeled target domain under distribution shifts induced by different users, sensors, devices, acquisition conditions, or temporal dynamics. Existing methods typically mitigate this shift by aligning marginal feature distributions through adversarial training, optimal transport, or moment-based discrepancies. In this paper, we propose Class-Conditional Path Dis",
      "authors": "Felix Ott, Christopher Mutschler",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-10T07:05:58.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18322"
    },
    {
      "id": 18323,
      "url": "https://arxiv.org/abs/2608.09053v1",
      "title": "Diagnosing as Cardiologists Do: ECG Agents with Doctor-Grounded Priors for Clinical Reasoning Across Diseases and Populations",
      "summary": "Cardiologists interpret electrocardiograms by localizing waveform components, measuring rhythm and interval patterns, and translating these structured observations into diagnostic evidence. Whether this expert reading process can serve as an effective prior for ECG agents remains unclear. To address this question, we introduce LuminaECG, a clinically structured ECG reasoning framework that reformulates ECG interpretation as measurement-grounded visual reading. ECG signals are rendered on standar",
      "authors": "Hongxiang Gao, He-yang Xu, Yuwen Li, Minghui Zhao, Zhipeng Cai, Xingyao Wang et al.",
      "category": "research",
      "topics": "healthcare,agents-autonomy",
      "published_at": "2026-08-10T03:00:35.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18323"
    },
    {
      "id": 18324,
      "url": "https://arxiv.org/abs/2608.08976v1",
      "title": "Label-Free Parkinson's Disease Screening from Face and Voice through Mechanistic Interpretability",
      "summary": "Parkinson's disease (PD) is the second most common neurodegenerative disorder. Typical machine learning screening methods require PD labels, but the available data is limited by privacy concerns and the need for expert annotation. We propose a label-free face-plus-voice PD screen built entirely on frozen pretrained encoders--a face-expression Vision Transformer and HuBERT--in which no PD label touches any fit; the reference is training controls only. The voice modality uses a synthetic-dysarthri",
      "authors": "Jiaheng Su, Yu Sun",
      "category": "research",
      "topics": "safety-alignment,privacy-surveillance",
      "published_at": "2026-08-10T00:43:57.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18324"
    },
    {
      "id": 18325,
      "url": "https://arxiv.org/abs/2608.08926v1",
      "title": "Decoding Phenotypes: A Framework for Fusing Genomic Language Models and Neuroimaging",
      "summary": "Neuroimaging and genetic testing are two important clinical references for nervous system diseases, offering complementary diagnostic information. However, integrating genomic and neuroimaging data for precise disease diagnosis is challenging due to cross-modality heterogeneity. Existing imaging-genetics approaches mainly encode genetic information as hard-coded labels, which lose the local sequence context around disease-associated variants. To address this limitation, we propose GeneFuse, a mu",
      "authors": "Tianli Tao, Ziyang Wang, Emma Robinson, Rachel Sparks, Le Zhang",
      "category": "research",
      "topics": "healthcare,biotech",
      "published_at": "2026-08-09T21:40:17.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18325"
    },
    {
      "id": 18326,
      "url": "https://arxiv.org/abs/2608.08815v1",
      "title": "Distilling Vision-Language Models for Robust Traffic Sign Perception in Autonomous Vehicles",
      "summary": "Traffic sign recognition (TSR) models based on deep neural networks achieve strong clean-data performance but remain vulnerable to physically realizable adversarial attacks, including shadow perturbations, natural-light interference, and printed patches. Existing defenses often improve robustness against one attack type while degrading performance on others, and can reduce clean accuracy. We propose LAMDA (Language-Anchored Model for Direction Alignment), a training framework that transfers lang",
      "authors": "Pedram MohajerAnsari, Amir Salarpour, Mert D. Pesé",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-09T17:05:00.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18326"
    },
    {
      "id": 18327,
      "url": "https://arxiv.org/abs/2608.08764v1",
      "title": "Learning from Consensus and Disagreement: Unsupervised On-Policy Self-Distillation with Minority-Trajectory Contrast",
      "summary": "On-policy self-distillation improves language-model reasoning by querying a teacher on states actually visited by the student. Recent methods create a powerful information asymmetry by exposing the teacher to privileged context, yet they fundamentally rely on external supervision---such as gold solutions or verifiers---to construct this advantage. We introduce CoDA (Consensus and Disagreement Alignment), a fully unsupervised framework that creates reliable privileged information entirely from th",
      "authors": "Jiaxin Guo, Yanwei Yue, Xuanbo Fan, Chunyu Yang, Yan Zhang",
      "category": "research",
      "topics": "regulation,safety-alignment,children-education",
      "published_at": "2026-08-09T15:23:25.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18327"
    },
    {
      "id": 18328,
      "url": "https://arxiv.org/abs/2608.08553v1",
      "title": "MotionCraft: Latent World Modeling with Sparse Attention for Visual Upscaling",
      "summary": "Video super-resolution (VSR) aims to recover high-fidelity high-resolution videos from low-resolution inputs and is central to applications ranging from mobile capture to streaming and archival restoration. Existing approaches trade off among local-detail fidelity, long-range spatio-temporal modeling, perceptual realism, and efficiency: convolutional alignment techniques preserve local structure but suffer when motion is large or degradations are complex; transformer-based methods capture long-r",
      "authors": "Rong Fu, Chunlei Meng, Yangchen Zeng, Xiaowen Ma, Yongtai Liu, Wangyu Wu et al.",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-09T07:54:43.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18328"
    },
    {
      "id": 18329,
      "url": "https://arxiv.org/abs/2608.08440v1",
      "title": "MGMCL: Multi-Granularity Manifold Contrastive Learning With Neural ODEs for Cross-Subject EEG Emotion Recognition",
      "summary": "Cross-subject electroencephalogram (EEG)-based emotion recognition remains challenging due to substantial inter-individual variability and discrete formulation that overlooks affective continuity. Existing methods operate in Euclidean space and focus on marginal distribution alignment, failing to preserve the semantic structure of emotions across subjects. This article proposes MGMCL, reconceptualizing emotion recognition as learning continuous representations on symmetric positive definite (SPD",
      "authors": "Xiang Xie",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-09T03:13:51.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18329"
    },
    {
      "id": 18330,
      "url": "https://arxiv.org/abs/2608.08309v1",
      "title": "Three Necessary Principles for Self-Supervised Visual Representation Learning",
      "summary": "We argue that learning visual representations without labels requires a training signal jointly complete across three non-overlapping objectives: semantic invariance across augmented views, patch-level spatial prediction, and representational non-degeneracy. We formalize these as the observation, prediction, and regularization principles and prove (i) that combining observation and prediction without regularization admits the constant encoder as a global minimizer under negative-free alignment;",
      "authors": "Nikos Giakoumoglou, Paschalis Giakoumoglou, Tania Stathaki",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-08T19:37:05.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18330"
    },
    {
      "id": 18331,
      "url": "https://arxiv.org/abs/2608.08294v1",
      "title": "A Controlled Study of Feature-Based Knowledge Distillation Across Student Designs",
      "summary": "Knowledge distillation trains a smaller student to match the outputs of a larger teacher. Feature-based methods also align intermediate representations, but this extra constraint may affect students differently. We study this question on CIFAR-100 using a ResNet-50 teacher, a width-controlled CustomResNet family and MobileNetV2 as a cross-design comparison. For each student, we evaluate each feature method against a matched logit-KD run using the same teacher, optimizer settings, training schedu",
      "authors": "Abhinand Balachandran, Praveen Prashant",
      "category": "research",
      "topics": "children-education",
      "published_at": "2026-08-08T19:10:17.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18331"
    },
    {
      "id": 18332,
      "url": "https://arxiv.org/abs/2608.08182v1",
      "title": "Biologically Informed Representation Learning for Robust Cross-Center Generalization of MALDI-TOF Mass Spectrometry",
      "summary": "Machine learning models for MALDI-TOF mass spectrometry have shown considerable promise for clinical microbiology tasks such as microbial identification and antimicrobial resistance prediction. However, their deployment across institutions remains limited by domain shift, as acquisition-specific variability often leads models to capture technical artifacts rather than transferable biological information. Existing representation learning approaches primarily address this problem through statistic",
      "authors": "Alejandro L. García-Navarro, Carlos Sevilla-Salcedo, Belén Rodríguez-Sánchez, Vanessa Gómez-Verdejo",
      "category": "research",
      "topics": "healthcare,biotech",
      "published_at": "2026-08-08T15:19:53.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18332"
    },
    {
      "id": 18333,
      "url": "https://arxiv.org/abs/2608.08148v1",
      "title": "DoGMA: A Central-Dogma-Guided Foundation Model for Multi-Omics Alignment and Multi-Task Learning in Oncology",
      "summary": "Attention mechanisms have been widely utilized in modern deep learning, and many existing multi-omics models inherit their conventional use to allow unrestricted bidirectional interactions. However, the fundamental logic of life is directional. Existing designs often overlook the directionality suggested by the central dogma, potentially limiting transfer across heterogeneous cancers, downstream tasks, and incomplete modality settings.In this work, we present DoGMA, a central-dogma-guided founda",
      "authors": "Junfei Ling, Bangzheng Pu, Bingsen Xue, Tianle Li, Ruying Hu, Cheng Jin",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-08T14:15:43.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18333"
    },
    {
      "id": 18334,
      "url": "https://arxiv.org/abs/2608.08138v1",
      "title": "EFFEKT: Efficient Federated Knowledge Transfer to Foundation Models",
      "summary": "Recent data protection laws have accelerated the adoption of Federated Learning (FL) for privacy-preserving decentralized training. Nevertheless, increasing model sizes impose substantial computational demands on client devices, limiting FL applicability in resource-constrained settings. We introduce a novel multi-domain federated learning framework in which lightweight client-side proxy models collaborate with a server-side Foundation Model (FM) to learn new concepts without sharing private dat",
      "authors": "Matteo Caligiuri, Francesco Barbato, Pietro Zanuttigh, Francesco Restuccia",
      "category": "research",
      "topics": "privacy-surveillance",
      "published_at": "2026-08-08T13:54:48.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18334"
    },
    {
      "id": 18335,
      "url": "https://arxiv.org/abs/2608.08111v1",
      "title": "Hierarchical Multi-Task Federated Learning in VANETs",
      "summary": "Vehicular Ad hoc Networks (VANETs) increasingly rely on federated learning (FL) to enable collaborative intelligence without sharing raw sensory data. However, most existing vehicular FL frameworks assume that all vehicles train a single global model for a common task, which limits their applicability in practical vehicular environments where vehicles may perform heterogeneous learning tasks under non-independent and identically distributed (non-IID) data, intermittent connectivity, and high mob",
      "authors": "M. Saeid HaghighiFard, Sinem Coleri",
      "category": "research",
      "topics": "environment",
      "published_at": "2026-08-08T12:46:04.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18335"
    },
    {
      "id": 18336,
      "url": "https://arxiv.org/abs/2608.07935v1",
      "title": "Adaptive Supervised Anchoring for On-Policy Self-Distillation",
      "summary": "On-policy self-distillation (OPSD) adapts a language model by distilling guidance from a frozen teacher on trajectories sampled from the student. Its effectiveness, however, depends critically on the quality of those trajectories. We show that when student rollouts drift from target trajectories, conditioning the teacher on off-target prefixes substantially weakens its task-relevant supervision. Controlled prefix-corruption experiments expose this failure mode, which we term rollout-conditioned",
      "authors": "Meilin Yang, Zixuan Ding, Jianhao Nie, Weite Zhang, Yuxin Zhang, Zhiming Shao et al.",
      "category": "research",
      "topics": "regulation,children-education",
      "published_at": "2026-08-08T05:46:32.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18336"
    },
    {
      "id": 18337,
      "url": "https://arxiv.org/abs/2608.07827v1",
      "title": "From token probabilities to calibrated confidence: An empirical study of mathematical question answering",
      "summary": "Confidence estimation for large language models (LLMs) aims to estimate the probability that a generated answer is correct, while calibration aligns these estimates with empirical accuracy. Prior work has shown that token probabilities are often overconfident, we investigate whether these readily available signals can nevertheless provide well-calibrated confidence estimation for mathematical question answering. We compare single-pass estimators, which reuse token probabilities from the original",
      "authors": "Avery Ma, Lorne Schell, Vin Bhaskara, Leila Pishdad",
      "category": "research",
      "topics": "finance-investment",
      "published_at": "2026-08-08T00:00:19.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18337"
    },
    {
      "id": 18338,
      "url": "https://arxiv.org/abs/2608.07786v1",
      "title": "Who Built This Model? Tracing LLM Lineage via Spectral Fingerprints in Weight Space",
      "summary": "Open-weight large language models (LLMs) are increasingly developed through complex, multi-stage pipelines, leading to intricate lineage relationships that reflect model origin, ownership, and evolution. Understanding these relationships is important for model provenance, governance, and supply-chain integrity. In this work, we investigate the notion of LLM \"biometrics\" (analogous to human biometrics) to ask whether LLMs exhibit intrinsic fingerprints in weight space alone, without access to inp",
      "authors": "Yiwei Chen, Bingqi Shang, Sijia Liu",
      "category": "research",
      "topics": "regulation,finance-investment",
      "published_at": "2026-08-07T22:18:02.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/18338"
    },
    {
      "id": 17937,
      "url": "https://arxiv.org/abs/2608.07419v1",
      "title": "Beyond Post-Hoc Temperature Scaling: Bilevel Optimization for LLM Calibration",
      "summary": "Preference alignment often makes large language models (LLMs) overconfident and poorly calibrated. Traditional post-hoc temperature scaling is inherently domain-dependent: a temperature fitted on one domain does not generalize across domains. This motivates us to modify model parameters during training to improve calibration. We propose maximizing the entropy of predictive distributions as the calibration objective, which directly targets overconfidence by discouraging overly concentrated predic",
      "authors": "Ruochen Jin, Zhanliang Wang, Zongyu Dai, Jiancong Xiao, Bojian Hou",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-07T17:05:10.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17937"
    },
    {
      "id": 17938,
      "url": "https://arxiv.org/abs/2608.07393v1",
      "title": "FedDOSE: Federated Learning Framework Decomposing Site Effects for Modeling Brain Dynamic Functional Connectivity",
      "summary": "Functional Magnetic Resonance Imaging ( fMRI ) data are often pooled into collaborative multi-site consortia, as deep learning models for analyses require large datasets to generalize well. While Federated Learning (FL) offers a privacy-preserving paradigm for collaborative training, standard approaches continue to struggle with statistical heterogeneity. In particular, site differences pose a key challenge in multi-site data settings. Additionally, existing FL approaches for fMRI rely on static",
      "authors": "Deepank Girish, Yi Hao Chan, Yubin Zheng, Sukrit Gupta, Jagath C. Rajapakse",
      "category": "research",
      "topics": "privacy-surveillance",
      "published_at": "2026-08-07T16:37:30.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17938"
    },
    {
      "id": 17939,
      "url": "https://arxiv.org/abs/2608.07371v1",
      "title": "Trajectory-Relative Hindsight Distillation for Agentic Reinforcement Learning",
      "summary": "Recent agentic reinforcement learning methods use hindsight to complement sparse outcome rewards. However, a completed rollout can yield many such signals, leaving their appropriate allocation across turns unclear. We introduce TRIAL, a trajectory-relative hindsight distillation framework with a unified turn-aligned scoring protocol. For each decision turn, TRIAL extracts an outcome view of that decision's realized consequence and evaluates the same response under ordinary and hindsight-conditio",
      "authors": "Haoyu Zheng, Yun Zhu, Qing Wang, Wenqiao Zhang",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-07T16:12:58.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17939"
    },
    {
      "id": 17940,
      "url": "https://arxiv.org/abs/2608.07328v1",
      "title": "Learning Fault-Tolerant Locomotion with Adaptive Gait Timing",
      "summary": "Hardware failures require legged robots to rapidly reorganize coordination and gait timing to maintain stability and mobility. This is particularly challenging for larger quadrupeds, where increased mass and tighter actuation limits reduce the feasibility of aggressive, high-frequency compensation strategies often observed on smaller platforms. In this work, we propose a deep reinforcement learning approach for fault-tolerant locomotion under actuator power loss. The method employs an asymmetric",
      "authors": "Giovanbattista Gravina, Luca Rossini, Carlo Rizzardo, Arturo Laurenzi, Nikos Tsagarakis",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-07T15:27:03.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17940"
    },
    {
      "id": 17941,
      "url": "https://arxiv.org/abs/2608.07281v1",
      "title": "High-dimensional ridgeless least squares interpolation under spiked covariance structures",
      "summary": "This paper investigates the asymptotic behavior of the out-of-sample prediction risk of the high-dimensional ridgeless least-squares estimator when the feature dimension $p$ and the sample size $n$ grow proportionally. We consider a generalized spiked population covariance model with multiple latent factors, where the number of spiked eigenvalues may remain finite or increase with $n$, and the spiked eigenvalues may be bounded or diverge at arbitrary rates. Beyond characterizing the impact of co",
      "authors": "Zhijun Liu, Dandan Jiang",
      "category": "research",
      "topics": "finance-investment",
      "published_at": "2026-08-07T14:40:44.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17941"
    },
    {
      "id": 17942,
      "url": "https://arxiv.org/abs/2608.06880v1",
      "title": "SkillAligner: Treating Retrieved Skills as Adaptable Drafts at Execution Time",
      "summary": "General-purpose skills promise reusable procedural knowledge for language agents, yet semantic relevance does not guarantee execution utility: a retrieved skill may encode assumptions that conflict with the current task, execution environment, or other retrieved skills. We formalize this problem as the skill--execution misfit. To address it, we propose SkillAligner, a training-free execution-time skill adaptation framework that treats retrieved skills as adaptable drafts rather than fixed instru",
      "authors": "Qinfeng Li, Dalin He, Yuntai Bao, Ying Yang, Ruoxi Chen, Xinyan Yu et al.",
      "category": "research",
      "topics": "agents-autonomy,environment",
      "published_at": "2026-08-07T07:05:53.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17942"
    },
    {
      "id": 17943,
      "url": "https://arxiv.org/abs/2608.06772v1",
      "title": "ArchEGraph: A Large-Scale Graph Dataset for Geometry-Topology-Physics Aligned Building Energy Modeling",
      "summary": "Accurate estimation of building energy use is essential for achieving carbon neutral and sustainable buildings. To better understand the influence of design decisions on building energy use and calibrate machine learning models that can give architects and engineers rapid design feedback, large-scale datasets are needed that explicitly map building geometry to performance. We present ArchEGraph, a large-scale benchmark dataset that represents buildings as heterogeneous graphs with aligned geomet",
      "authors": "Yihui Li, Yihui Chen, Kaidi Zha, Xiaoyue Yan, Zhexuan Yu, Shiqi Dai et al.",
      "category": "research",
      "topics": "environment",
      "published_at": "2026-08-07T03:45:41.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17943"
    },
    {
      "id": 17944,
      "url": "https://arxiv.org/abs/2608.06595v1",
      "title": "Flowing Through States: Neural ODE Regularization for Reinforcement Learning",
      "summary": "Neural networks applied to sequential decision-making tasks typically rely on latent representations of environment states. While environment dynamics dictate how semantic states evolve, the corresponding latent transitions are usually left implicit, creating a potential misalignment between the two. We propose to model latent dynamics explicitly by drawing an analogy between Markov decision process (MDP) trajectories and ordinary differential equation (ODE) flows: in both cases, the current sta",
      "authors": "Mohamed Ghanem, Bernd Finkbeiner",
      "category": "research",
      "topics": "safety-alignment,environment",
      "published_at": "2026-08-06T21:11:16.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17944"
    },
    {
      "id": 17945,
      "url": "https://arxiv.org/abs/2608.06582v1",
      "title": "CrystalGRPO: Target-Aligned and Coverage-Preserving Reinforcement Learning for Flow-Based Crystal Structure Prediction",
      "summary": "Flow-based generative models can efficiently produce candidate structures for crystal structure prediction (CSP), but their pretrained objectives do not directly optimize downstream target recovery. Reinforcement-learning post-training offers a flexible solution, yet existing approaches rely primarily on energy rewards and coordinate-only stochastic policies. Predicted energy does not identify the reference polymorph, while reward-driven concentration can reduce the candidate coverage required f",
      "authors": "Kaixiang Su, Hongfei Xue, Qiang Zhu",
      "category": "research",
      "topics": "environment",
      "published_at": "2026-08-06T20:53:00.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17945"
    },
    {
      "id": 17391,
      "url": "https://arxiv.org/abs/2608.06246v1",
      "title": "A Six-Dimensional Taxonomy of Post-Training Adaptation Techniques with Applications in AI Governance",
      "summary": "Post-training adaptation has become central to modern machine learning practice and includes techniques such as retraining, fine-tuning, parameter-efficient adaptation, alignment, retrieval augmentation, model editing, unlearning, calibration, and Multimodal Instruction Tuning. However, the literature remains fragmented across technique families, model classes, and deployment contexts, making it difficult to compare methods or describe how a trained model has been modified. This survey synthesiz",
      "authors": "Fardin Afdideh, Fernando Seoane, Farhad Abtahi",
      "category": "research",
      "topics": "regulation,safety-alignment",
      "published_at": "2026-08-06T16:32:26.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17391"
    },
    {
      "id": 17392,
      "url": "https://arxiv.org/abs/2608.06179v1",
      "title": "SAGA: Score-Weighted Adaptive Generation Alignment for Low-Resource Nordic Language Models",
      "summary": "Preference optimisation has proven effective for improving large language models but typically relies on costly human preference annotations. Extending these methods to morphologically rich, low-resource languages remains challenging because such annotations are scarce. We present SAGA (Score-weighted Adaptive Generation Alignment), a parser-guided preference optimisation framework that replaces human labels with dependency-parser supervision. SAGA converts parser judgements into preference pair",
      "authors": "Hoda Fakharzadehjahromy, Emil Wiman, Andreas Bueff, Hafsteinn Einarsson, Fredrik Heintz",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-06T15:41:02.000Z",
      "source": "arXiv cs.LG",
      "ethics_ai_record_url": "https://ethics.ai/record/17392"
    }
  ],
  "attribution": "via ethics.ai"
}