{
  "count": 50,
  "items": [
    {
      "id": 19155,
      "url": "https://arxiv.org/abs/2608.13560v1",
      "title": "AutoDesign: Meta-Harness Optimization for Long-Horizon Agentic Design",
      "summary": "Transforming multimodal sources into condensed and structured media outputs can be fundamentally conceptualized as a long-horizon agentic process centered on a model-harness system. While an ideal harness system should align with human design priors and accumulate reusable experience through empirical exploration to drive recursive self-improvement, existing paradigms remain static and fall short of this capability. In this paper, we present AutoDesign, a framework that aligns with human design",
      "authors": "Yaxin Luo, Haobin Jiang, Jialv Zou, Xu Huang, Wenhao Yan, Haodong Li et al.",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-13T17:59:57.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19155"
    },
    {
      "id": 19156,
      "url": "https://arxiv.org/abs/2608.13555v1",
      "title": "HumanTracker: Towards Comprehensive and Human-Aligned Motion Tracking Benchmark",
      "summary": "Humanoid motion tracking is central to teleoperation and whole-body imitation, yet evaluation often disagrees with what people perceive in videos. Kinematic errors average per-frame pose differences but miss the physical artifacts that matter most, particularly unstable support and incorrect contacts such as foot skating and mistimed touch-downs. Meanwhile, widely used test suites are small and lack the diversity needed to stress contact-rich, long-horizon behaviors. We introduce HumanTracker to",
      "authors": "Dairu Liu, Zekun Qi, Jiayu Zeng, Ruixi Yu, Yu Guan, Yintianrun Zhang et al.",
      "category": "research",
      "topics": "privacy-surveillance",
      "published_at": "2026-08-13T17:59:40.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19156"
    },
    {
      "id": 19157,
      "url": "https://arxiv.org/abs/2608.13517v1",
      "title": "DFM Mimir v1: An Open HRM Delivering Frontier Performance at 1B Parameters Using Only Permissible Post-Training Data",
      "summary": "Current large language model development relies on massive, often non-permissible datasets, creating a high barrier for researchers committed to open-source and ethically sourced data. We introduce Mimir v1, a 1-billion-parameter language model based on the Hierarchical Reasoning Model (HRM) architecture, that is trained from scratch and delivers highly competitive performance for English and sets a new state of the art for Danish using only permissible post-training data. Trained on a mixture o",
      "authors": "Peter Schneider-Kamp, Jacob Nielsen, Gianluca Barmina, Kenneth Enevoldsen, Lukas Galke Poech",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T17:37:53.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19157"
    },
    {
      "id": 19158,
      "url": "https://arxiv.org/abs/2608.13492v1",
      "title": "AlayaWorld: Interactive Long-Horizon World Modeling - Full Technical Report (v1.1)",
      "summary": "This report presents an improved version of AlayaWorld. While the backbone architecture, chunk-wise autoregressive generation scheme, and training data remain unchanged from the previous release, we substantially revise how conditioning signals are represented and integrated into the model. The new design is guided by a simple principle: conditioning signals should match the generated content as closely as possible in both latent representation and temporal structure. To this end, we make two ma",
      "authors": "AlayaWorld Team, Kaipeng Zhang, Chuanhao Li, Yifan Zhan, Yongtao Ge, Yuanyang Yin et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T17:21:03.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19158"
    },
    {
      "id": 19159,
      "url": "https://arxiv.org/abs/2608.13484v1",
      "title": "Toward a Gricean Retreat: Probing LLMs for Knowledge Boundaries and Referent Specificity",
      "summary": "When asked about entities outside their knowledge boundary, LLMs routinely fabricate plausible-sounding details rather than backing off to safer, more general claims. We frame this failure through a Gricean lens: a cooperative speaker who is uncertain about a referent retreats up the specificity hierarchy, trading informativeness for truthfulness. We ask whether LLMs have the ingredients to perform this retreat. Using a T-REx-based benchmark that varies entity familiarity and referent specificit",
      "authors": "Dananjay Srinivas, Saksham Khatwani, Maria Pacheco",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T17:13:41.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19159"
    },
    {
      "id": 19160,
      "url": "https://arxiv.org/abs/2608.13482v1",
      "title": "Synthetic Persona Pretraining: Alignment from Token Zero",
      "summary": "As language-model-based AI is increasingly deployed in autonomous settings, aligning its goals and values with those of humans becomes critical. Today, alignment, and the assistant identity itself, are typically introduced only after pretraining, once behavioral priors are already established. This can make values a thin overlay, rather than deeply rooted, and facilitate subsequent misalignment. Pursuing a different paradigm, we introduce Synthetic Persona Pretraining (SPP), which installs the d",
      "authors": "Julian Minder, Viktor Moskvoretskii, Raghav Singhal, Difan Jiao, Andy Arditi, Shaobo Cui et al.",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-13T17:12:04.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19160"
    },
    {
      "id": 19161,
      "url": "https://arxiv.org/abs/2608.13456v1",
      "title": "A Unifying Perspective on Causal World Models: From Observations to Representations to Structure",
      "summary": "World Models (WM) are increasingly seen as a foundation for intelligent agents that can predict, plan, and act beyond their training distribution. In this paper, we study WMs from a causal perspective across multiple levels of abstraction, ranging from perceptual observations to building a conceptual representation of the structure governing the environment dynamics. We argue that useful WMs must go beyond generative capabilities alone: they should also capture entity properties, entity-to-entit",
      "authors": "Avinash Kori, Fabrizio Russo",
      "category": "research",
      "topics": "agents-autonomy,environment",
      "published_at": "2026-08-13T16:40:35.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19161"
    },
    {
      "id": 19162,
      "url": "https://arxiv.org/abs/2608.13453v1",
      "title": "UniTexture: Cross-Task Universal Adversarial Textures for Vision-Language-Action Models",
      "summary": "Vision-Language-Action (VLA) models have emerged as generalist robotic policies capable of following diverse language instructions and performing a wide range of manipulation tasks. However, their direct control over embodied agents also exposes them to adversarial interference that may cause unsafe physical behaviors. Existing attacks on robotic policies are typically optimized for a single task or instruction, leaving the cross-task vulnerabilities of multitask VLAs largely unexplored. We intr",
      "authors": "Yukun Dai, Mingzhe Dai, Tianshi Wang, Fengling Li, Jingjing Li, Lei Zhu",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-13T16:38:57.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19162"
    },
    {
      "id": 19163,
      "url": "https://arxiv.org/abs/2608.13447v1",
      "title": "Academic League of Artificial Intelligence - An Integrative Perspective of Teaching, Research, and Extension",
      "summary": "Academic leagues have become important mechanisms for promoting extracurricular education and strengthening the integration between universities and society. This paper presents the organizational framework adopted by the Academic League of Artificial Intelligence (LIA) at the Federal University of Santa Catarina (UFSC), designed to integrate teaching, research, and university extension through a student-centered, project-based approach. The framework combines democratic governance, collaborativ",
      "authors": "Alison R. Panisson, Maria Eduarda W. M. Vianna, Italo Firmino da Silva, Heitor Henrique da Silva, Rafaela Fernandes Savaris, Bernardo Pandolfi Costa et al.",
      "category": "research",
      "topics": "regulation,children-education",
      "published_at": "2026-08-13T16:32:44.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19163"
    },
    {
      "id": 19164,
      "url": "https://arxiv.org/abs/2608.13444v1",
      "title": "Algorithmic Gender Prediction Is Illegitimate, But Gender Imputation Can Yield Valid Measurements",
      "summary": "Machine learning ethics researchers and critical HCI scholars have argued that algorithmically predicting gender is wrong. At the same time, other researchers rely on predicted gender labels to study gender disparities and develop algorithmic fairness techniques. How do we reconcile these two seemingly contradictory intuitions? We differentiate two ways gender prediction may be wrong: being illegitimate, thereby contributing to harm; and being invalid, thereby producing unusable measurements. Ou",
      "authors": "Evan Dong, Angelina Wang",
      "category": "research",
      "topics": "bias-fairness",
      "published_at": "2026-08-13T16:30:47.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19164"
    },
    {
      "id": 19165,
      "url": "https://arxiv.org/abs/2608.13394v1",
      "title": "Heterogeneity-Aware Belief Synchronization for Semantic Communication in AI-Native 6G Networks",
      "summary": "6G networks will not be serving as communication infrastructures only; rather, they are expected to evolve into intelligent systems, where thousands of autonomous artificial intelligence (AI) agents are interconnected. The agents are deployed across a wide range of platforms including low Earth orbit (LEO) satellites, high-altitude platforms (HAPs), unmanned aerial vehicles (UAVs), edge servers, and terrestrial devices. These agents continuously observe their environment and exchange information",
      "authors": "Muhammad Hannan Akram, Muhammad Abubakar Rashid, Wassi Haider Kabir, Haejoon Jung, Kapal Dev, Syed Ali Hassan",
      "category": "research",
      "topics": "agents-autonomy,environment",
      "published_at": "2026-08-13T15:53:15.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19165"
    },
    {
      "id": 19166,
      "url": "https://arxiv.org/abs/2608.13389v1",
      "title": "TopoIntent: Compiling Security Intent into Executable, Compliance-Checked Network Topologies",
      "summary": "Enterprise security topology design requires translating business intent, regulatory requirements, and risk assumptions into zones, boundary devices, inter-zone paths, and access-control policies. Existing NetOps automation tools mainly operate after this design is fixed, providing limited support for generating structured security topologies from underspecified natural-language requirements. We present TopoIntent, a system that compiles security intent into executable, compliance-checked networ",
      "authors": "Xiaokang Qu, Jianliang Ma, Zao Fan, Tianshu Chu, Tianlong Fan, Linyuan Lü",
      "category": "research",
      "topics": "regulation,jobs-economy",
      "published_at": "2026-08-13T15:49:37.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19166"
    },
    {
      "id": 19167,
      "url": "https://arxiv.org/abs/2608.13345v1",
      "title": "Rules or Character? Scaling Laws for AI Safety Design",
      "summary": "Artificial Intelligence (AI) safety systems combine character shaping (e.g., Reinforcement Learning from Human Feedback [RLHF], Constitutional AI), which modifies behavioral distributions at training time, with rule enforcement (e.g., output filters, safety classifiers), which blocks harmful outputs at inference time, yet little formal analysis exists on how their optimal balance should change as deployment scales increase. We introduce a stylized comparative-statics model that parameterizes saf",
      "authors": "Satoshi Takahashi, Nobuji Kouno, Masaaki Komatsu, Ryuji Hamamoto",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-13T15:15:09.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19167"
    },
    {
      "id": 19168,
      "url": "https://arxiv.org/abs/2608.13344v1",
      "title": "LongEarth-R1: Benchmarking and Aligning Vision-Language Models for Long-Horizon Earth Observation Reasoning",
      "summary": "Long-horizon Earth observation reasoning requires models to organize multi-stage geographic evolution, localize spatial changes, detect temporal anomalies, and infer future from extended image sequences. However, existing remote sensing vision-language models mainly focus on isolated images, image pairs, or short sequences, limiting reliable grounding in the relevant frames and regions. We introduce LongEarth-Bench, a benchmark containing approximately 120k question-answering samples derived fro",
      "authors": "Yupan Ding, Jing Xiao, Zhenyuan Zhang, Chaofeng Chen, Liang Liao, Gui-Song Xia et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T15:14:59.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19168"
    },
    {
      "id": 19169,
      "url": "https://arxiv.org/abs/2608.13341v1",
      "title": "Simulation-to-real transfer learning for infrared spectroscopic chemical sensing and analysis from molecules to complex samples",
      "summary": "Infrared (IR) spectroscopy is widely used for chemical sensing, but extracting reliable chemical information from spectra remains challenging. Conventional interpretation is labor-intensive, relies on prior knowledge and reference spectra, and is difficult to scale, whereas most machine-learning methods are tailored to individual tasks or datasets, require large labeled training sets, and transfer poorly across analytical objectives and experimental datasets. Here we introduce UltraIR, a foundat",
      "authors": "Yusen Tan, Yixuan Chen, Zheng Fang, Pan Liu, Yifan Li, Qinyu Guo et al.",
      "category": "research",
      "topics": "jobs-economy",
      "published_at": "2026-08-13T15:11:50.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19169"
    },
    {
      "id": 19170,
      "url": "https://arxiv.org/abs/2608.13317v1",
      "title": "StateBridge: Training-free Hidden-state Alignment for Latent Communication in LLM Multi-Agent Systems",
      "summary": "Large language model based multi-agent systems usually communicate in text, i.e., using discrete tokens. However, text introduces a discrete bottleneck. Converting the sender's continuous hidden states into discrete tokens discards information that token identities alone cannot capture. Recent work proposes latent communication as an alternative, where agents transmit hidden representations directly without converting them to text. However, existing latent methods either inject working memory la",
      "authors": "Yanwen Peng, Delvin Ce Zhang, Xi Wang, Nikolaos Aletras",
      "category": "research",
      "topics": "safety-alignment,agents-autonomy",
      "published_at": "2026-08-13T14:40:59.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19170"
    },
    {
      "id": 19171,
      "url": "https://arxiv.org/abs/2608.13277v1",
      "title": "Mixture of Training: Recombining Small-Scale Scaffolded Pretraining Runs into a Larger Language Model",
      "summary": "We ask whether language-model pre-training can be decomposed into smaller, independently trainable jobs that can later be recomposed into a coherent larger model. We introduce Mixture of Training (MoT), a scaffolded modular pre-training procedure that partitions a target Transformer into contiguous layer blocks, trains each block inside a frozen pretrained aligner scaffold, and then recomposes the trained blocks with an optional short end-to-end adaptation pass. On a 1.3B-parameter Gemma-style m",
      "authors": "Mohammed Sabry, Sean Augenstein, Keith Rush, Lucio Dery",
      "category": "research",
      "topics": "jobs-economy",
      "published_at": "2026-08-13T14:13:46.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19171"
    },
    {
      "id": 19172,
      "url": "https://arxiv.org/abs/2608.13272v1",
      "title": "Sovereign by necessity? Frontier AI export controls, cyber security, and the limits of national AI capability",
      "summary": "A small number of firms based in two states produce the most capable frontier AI models. The governments of those states have shown both the legal power and the political will to decide which other countries may use these systems. In June 2026 the United States required a leading developer to obtain licences before releasing its most advanced models to any foreign person, including foreign nationals resident in the United States. The affected models were withdrawn worldwide at short notice, part",
      "authors": "Alan Woodward, Andrew Rogoyski",
      "category": "research",
      "topics": "military-security",
      "published_at": "2026-08-13T14:09:45.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19172"
    },
    {
      "id": 19173,
      "url": "https://arxiv.org/abs/2608.13262v1",
      "title": "Into the ORBIT for Time Series: Training Regimes for Foundation Models",
      "summary": "Time series foundation models (TSFMs) have advanced primarily through architectural innovation, while training regimes for large-scale heterogeneous corpora remain under-explored. As a result, pre-training distributions are often poorly controlled with respect to domain imbalance, context requirements, prediction horizons, and missingness. We introduce ORBIT (Omni-Range Bootstrap Incremental Training), a training paradigm that makes this distribution explicit and controllable. ORBIT combines Boo",
      "authors": "Hongjie Xia, Yiding Liu, Yifan Hu, Peiyuan Liu, Zewei Dong",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T14:00:39.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19173"
    },
    {
      "id": 19174,
      "url": "https://arxiv.org/abs/2608.13256v1",
      "title": "Novel Knowledge-Guided Generative Methods for Synthetic Transcriptomic Data",
      "summary": "As biomedical research increasingly relies on data-intensive tools, the quality and utility of datasets are critical. Challenges such as imbalances, biases, and ethical or legal constraints often limit access to high-quality data. Synthetic data generation can help overcome these limitations. Here, we present a comparative analysis of generative models for transcriptomic data, investigating strategies to incorporate prior biological knowledge via gene graphs. This ensures that synthetic data cap",
      "authors": "Francesca Pia Panaccione, Sofia Mongardi, Marco Masseroli, Pietro Pinoli",
      "category": "research",
      "topics": "finance-investment,biotech",
      "published_at": "2026-08-13T13:57:47.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19174"
    },
    {
      "id": 19175,
      "url": "https://arxiv.org/abs/2608.13255v1",
      "title": "GeoCache: Training-Free Acceleration of Multi-View Texture Diffusion via Geometric Delta Transport",
      "summary": "Geometry-conditioned multi-view diffusion enables high-quality 3D texture generation, but its repeated per-view denoiser evaluations introduce substantial computational cost. Existing training-free accelerators primarily exploit temporal redundancy by reusing computation across denoising steps. In multi-view texturing, however, skipping a step also removes the cross-view interaction that continually aligns different observations of the same surface, leading to rapidly degraded consistency and fi",
      "authors": "Haotang Li, Zhenyu Qi, Shaohan Henry Wang, Kebin Peng, Yutong Zhao, Zi Wang et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T13:57:35.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19175"
    },
    {
      "id": 19176,
      "url": "https://arxiv.org/abs/2608.13250v1",
      "title": "Follow the Norm: Accounting for Fine-Tuning and Prompt Effects on Model Rationales",
      "summary": "Normative datasets are often used to train and align AI systems, but the norms they contain can function as action-guiding patterns rather than neutral moral knowledge. We propose treating the AI system as a proxy actor and test whether dataset-level norms can shift it away from its baseline safety behavior when it faces high-conflict dilemmas. We make three contributions. First, we demonstrate in controlled experiments that norm-breaking fine-tuning yields norm-divergent actions justified by se",
      "authors": "Long Hoang Nguyen, Brice Valentin Kok-Shun, Guangyu Du, Ali Sunyaev",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T13:55:03.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19176"
    },
    {
      "id": 19177,
      "url": "https://arxiv.org/abs/2608.13228v1",
      "title": "Capability Sheaves for Compositional Agent-Harness Repair: Controlled Quotients and a Real-Repository Stress Test",
      "summary": "Agent harnesses combine retrieval, routing, state, provenance, and verification, but locally successful components may disagree on shared state. We model this failure with a finite \\emph{capability sheaf}: stalks encode typed behavior signatures, restriction maps retain shared fields, and accepted runs are useful global sections. An exact finite constraint-satisfaction problem (CSP) defines acceptance, while a linearized relative cohomology class provides a diagnostic and search feature. A contr",
      "authors": "Saveliy Batruin",
      "category": "research",
      "topics": "healthcare,agents-autonomy",
      "published_at": "2026-08-13T13:31:09.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19177"
    },
    {
      "id": 19178,
      "url": "https://arxiv.org/abs/2608.13160v1",
      "title": "Better Decomposition, Free Aggregation: A Synthesizer-Folding Framework for Multilingual Multi-Hop Question Answering",
      "summary": "Multilingual retrieval-augmented generation (mRAG) equips large language models with access to globally distributed external knowledge for complex multilingual question answering. Recent approaches either translate retrieved documents into English or the query language to bridge the cross-lingual semantic gap, or decompose a complex query into sub-questions and aggregate the intermediate reasoning process. However, both lines of work suffer from two limitations. First, one-size-fits-all translat",
      "authors": "Yilin Wang, Yuchun Fan, Weidong Bao, Zili Wei, Shi Feng, Tong Xiao et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T12:25:59.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19178"
    },
    {
      "id": 19179,
      "url": "https://arxiv.org/abs/2608.13136v1",
      "title": "LigBench: A Unified and Human-Aligned Benchmark for LLM-based Research Idea Generation",
      "summary": "With the rapid advancement of large language models (LLMs), research idea generation has attracted increasing attention. Existing approaches enable LLMs to retrieve relevant literature and propose novel ideas for research areas. However, current evaluation practices for idea generation remain fragmented and lack objective standards, often relying on direct LLM scoring, which limits their ability to provide unified and reliable assessments across a coherent distribution of generated ideas. To add",
      "authors": "Chenrun Wang, Mingxuan Zhu, Tiancheng Huang, Wenjie Li, Yujie Zhang, Zichen Zhu et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T12:11:23.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19179"
    },
    {
      "id": 19180,
      "url": "https://arxiv.org/abs/2608.13120v1",
      "title": "SkillEvo: Self-Renewing Evolution Gradients from Multi-Turn Interaction Feedback",
      "summary": "Agent Skills are today either hand-authored or produced in a single LLM generation pass, and consequently possess no closed loop through which they might improve from the interaction failures they actually cause. Recent work does close this loop, but derives its feedback from single-turn question-answering evaluation. The consequence is a sharp asymmetry: once the first round has patched the gaps that a single exchange can reveal, the evolution gradient decays, the defects that surface only acro",
      "authors": "Qianxi Yan, Chunrong Chen, Jiuzhou Zhao, Min Zhang, Yongzhou Xu, Xiaochuan Xu",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-13T11:49:02.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19180"
    },
    {
      "id": 19181,
      "url": "https://arxiv.org/abs/2608.13100v1",
      "title": "Multi-Layer Context Camouflaging: A Semantic Superposition and Contextual Lamination Framework for Malpractice-Resilient Online Assessment",
      "summary": "Contemporary online assessment systems rely primarily on browser lockdown, webcam monitoring, and behavioural analytics, yet remain vulnerable to attacks that extract the assessment content itself through screenshots, screen sharing, optical character recognition, and automated scraping. This paper extends the Multi-dimensional Spatio-Temporal Context Camouflaging Model (MSCCM) within the MARS (Multi-modal Assessment Resilience Suite) by introducing the Multi-Layer Context Camouflaging Theory (M",
      "authors": "Gupta Lovi Raj, Kaur Kamalpreet, Dama Sri Ram, Parani Prajithaa",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T11:25:12.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19181"
    },
    {
      "id": 19182,
      "url": "https://arxiv.org/abs/2608.13072v1",
      "title": "EEG-PRIME: Prototype-Aligned Representation Learning with Multi-Level Conditioning for EEG Decoding",
      "summary": "Electroencephalography (EEG) decoding models often generalize poorly across datasets and subjects due to domain shifts in acquisition protocols and individual neurophysiology. We propose EEG-PRIME, a two-stage EEG foundation model for cross-dataset multi-task decoding. EEG-PRIME combines masked pretraining with prototype-aligned instruction tuning to enable instruction-aware and subject-invariant decoding across diverse BCI paradigms. During pretraining, an EEG encoder learns transferable repres",
      "authors": "Shuailei Zhang, Muyun Jiang, Wei Zhang, Jinbo Chen, Zhiwei Guo, Yong Li et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T10:35:51.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19182"
    },
    {
      "id": 19183,
      "url": "https://arxiv.org/abs/2608.13069v1",
      "title": "Behavioral Reprogramming of Open-Weights Models: Cognitive Plasticity and Alignment Bounds",
      "summary": "Large language models (LLMs) are predominantly aligned to function as passive, sycophantic assistants. We challenge this default paradigm by empirically evaluating the cognitive plasticity of open-weight architectures when subjected to rigorous behavioral reprogramming. Our objective is to induce a proactive, Socratic conversational framework, characterized by high-frequency question generation under strictly constrained high-performance computing (HPC) conditions. Through a massively paralleliz",
      "authors": "Lucia Malíčková",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-13T10:33:00.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19183"
    },
    {
      "id": 19184,
      "url": "https://arxiv.org/abs/2608.13043v1",
      "title": "From Local Mismatch to Global Impact: Optimizing Cache Reuse Policy for Efficient Diffusion",
      "summary": "Diffusion models have achieved dominant performance in visual generation but suffer from substantial inference overhead. While cache-based acceleration has emerged as a promising solution, existing policies rely on local similarity heuristics, which we identify as being significantly misaligned with final generation quality. This discrepancy stems from the non-uniform propagation and accumulation of errors along the denoising trajectory. To address this, we propose Global-Impact Cache (GCache).",
      "authors": "Xichen Ye, Yifan Wu, Zhikang Xie, Xiangyu Yue, Cheng Jin, Weizhong Zhang",
      "category": "research",
      "topics": "regulation,safety-alignment",
      "published_at": "2026-08-13T10:08:47.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19184"
    },
    {
      "id": 19185,
      "url": "https://arxiv.org/abs/2608.13022v1",
      "title": "Applied and Filtered: An End-to-End Algorithmic Fairness Audit of A Public Employment Agency",
      "summary": "Algorithmic fairness evaluation commonly assesses AI systems as bounded technical components, abstracting away the organizational context in which they operate. We present, to our knowledge, the first independent end-to-end fairness audit of a semi-automated hiring system operated by Barcelona Activa, a public employment agency using the third-party TalentClue platform for candidate search and shortlisting. We analyze approximately 497,000 candidate-vacancy pipeline entries from September 2017 t",
      "authors": "Gemma Galdón-Clavell",
      "category": "research",
      "topics": "bias-fairness,jobs-economy,transparency",
      "published_at": "2026-08-13T09:46:27.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19185"
    },
    {
      "id": 19186,
      "url": "https://arxiv.org/abs/2608.12935v1",
      "title": "Decomposition of Evidence, Contradiction, and Fragility in Perturbation Responses",
      "summary": "Perturbation methods explain model decisions by measuring prediction changes under altered inputs, but response magnitude tells us only how much a model reacts, not what that reaction means. The same magnitude can support the final factual-counterfactual difference, oppose it, or arise strongly along the perturbation path yet vanish at the endpoint. We therefore track how the contrast develops as paired inputs are progressively revealed, using the final contrast to interpret the trajectory. We i",
      "authors": "Lei You",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T08:13:22.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19186"
    },
    {
      "id": 19187,
      "url": "https://arxiv.org/abs/2608.12854v1",
      "title": "BrainWAM: Action-Space Coordination of Semantic Priors and Predictive Dynamics for Autonomous Driving",
      "summary": "Autonomous driving requires planning under both semantic constraints and predictive dynamics. Existing end-to-end driving approaches, however, typically emphasize only one side of this requirement: Vision-Language-Action (VLA) models exploit VLM priors for semantic reasoning, while World Action Models (WAMs) provide future-aware prediction through generative world modeling. This naturally motivates a unified planner that can leverage both semantic priors and predictive dynamics. However, we find",
      "authors": "Bing Zhan, Shuyao Shang, Jiahao Gu, Shuo Lu, Yuan Xu, Zhao Wang et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T05:56:17.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19187"
    },
    {
      "id": 19188,
      "url": "https://arxiv.org/abs/2608.12851v1",
      "title": "Practice Makes Unsafe: Skill Misevolution in Self-Improving LLM Agents",
      "summary": "Self-improving LLM agents convert successful trajectories into persistent cross-task state. An unsafe success can thereby become reusable policy after its triggering input disappears. Skill evolution makes this failure measurable by distilling operational trajectories into executable, transferable, and inspectable procedures. Because evolution optimizes task outcomes rather than procedure safety, compromised experience can cause skill misevolution. Existing benchmarks measure current behavior or",
      "authors": "Xutao Mao, Liangjie Zhao, Xiang Zheng, Cong Wang",
      "category": "research",
      "topics": "regulation,agents-autonomy",
      "published_at": "2026-08-13T05:47:43.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19188"
    },
    {
      "id": 19189,
      "url": "https://arxiv.org/abs/2608.12845v1",
      "title": "FSGR: Mitigating Token Frequency Bias for Fair SID-Based Generative Recommendation",
      "summary": "Semantic ID (SID)-based generative recommendation has recently achieved remarkable success. However, existing methods suffer from a previously overlooked fairness issue, which we term \\textbf{Token Frequency Bias}, where high-frequency SID tokens are systematically over-predicted while low-frequency SID tokens are under-predicted. This bias originates from the combined effects of imbalanced semantic codebooks during SID construction, and popularity bias together with the maximum likelihood estim",
      "authors": "Yuchen Zheng, Sihan Xu, Jingwen Yang, Xiangrui Cai, Haiwei Zhang, Xiaojie Yuan",
      "category": "research",
      "topics": "bias-fairness",
      "published_at": "2026-08-13T05:34:51.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19189"
    },
    {
      "id": 19190,
      "url": "https://arxiv.org/abs/2608.12843v1",
      "title": "Heterogeneous Vision-Language Ensemble with Disagreement-Aware Reranking for Text-Based Person Anomaly Retrieval",
      "summary": "Text-based person anomaly retrieval aims to retrieve pedestrians exhibiting anomalous behaviors from a large image gallery using natural language descriptions. Compared with conventional text-based person retrieval, this task requires fine-grained reasoning over pedestrian appearance, behaviors, object interactions, and scene context, making robust cross-modal matching significantly more challenging. This paper presents the GENAI4E team's solution to AI City Challenge 2026 Track 4. Our framework",
      "authors": "Huu-An Vu, Cam Tu Tran Thi, Thanh Toan Le Ngo, Hoang Vo, Do Trung Hieu, Hieu Dinh Trung Pham et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-13T05:28:50.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19190"
    },
    {
      "id": 19191,
      "url": "https://arxiv.org/abs/2608.12788v1",
      "title": "ARAC: Benchmarking Auto-Research's Alignment and Completeness on End-to-End Researchs",
      "summary": "The rapid advancement of Auto-Research has surfaced a fundamental evaluation challenge: how can we measure the alignment, logical coherence, and evolutionary completeness of its research trajectory with human research behavior? We propose Auto-Research's Alignment and Completeness, ARAC-Bench: a Researcher-Mimicking Evaluation framework that shifts the objective from matching final answers to reproducing high-quality human research processes. The framework operates through two synergistic compon",
      "authors": "Jiale Cui, Yueyao Yuan, Kaixi Zhong, Xiaogang Xu, Jiafei Wu, Zhe Liu",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-13T03:48:07.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19191"
    },
    {
      "id": 19192,
      "url": "https://arxiv.org/abs/2608.12761v1",
      "title": "Correct Is Not Governed: Provenance Integrity in Agentic Workflows",
      "summary": "Agentic workflows are commonly evaluated by whether they reach the correct outcome. That is insufficient in institutional settings, where a correct action may rely on the wrong authority, an unsupported completion claim, or work made stale by a later change. We define governed execution as work whose decisions, completion, and response to change are supported by inspectable provenance. We present Matrix, a deterministic causal-state layer that records authority and fact dependencies, verifies co",
      "authors": "Jesus Salas",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-13T03:12:13.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19192"
    },
    {
      "id": 19193,
      "url": "https://arxiv.org/abs/2608.12720v1",
      "title": "ERSkill: Evolving for Skill-Guided Adaptive Memory Retrieval",
      "summary": "While Large Language Model (LLM) agents increasingly rely on long-term memory for persistent interactions, the retrieval mechanisms governing this memory are rarely treated as evolvable components. This static approach limits performance on heterogeneous memory queries, which often demand diverse evidence construction strategies. To address this, we introduce \\textbf{ERSkill}, a retrieval-centric framework for self-evolving, skill-guided memory access. ERSkill compiles interaction histories into",
      "authors": "Haolong Chen, Liang Zhang, Zhuo Li, Lei Xue, Guanrxu Zhu",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-13T02:06:01.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19193"
    },
    {
      "id": 19194,
      "url": "https://arxiv.org/abs/2608.12689v1",
      "title": "Mr3D-VL: A generalist vision language foundation model for Multiparametric 3D Magnetic Resonance Imaging",
      "summary": "Multi-parametric magnetic resonance imaging (mpMRI) is a cornerstone for brain tumor diagnosis and treatment, yet current AI models face critical limitations: their lack of natural language interaction and interpretability impedes spatial information integration and cross-modal reasoning required clinically. Key challenges arise from significant physical meaning differences across modalities, spatial misalignment due to scan intervals, and the need for complex multi-feature interpretation in tas",
      "authors": "Zhi Qiao, Xintong Wu, Yichu He, Feng Shi",
      "category": "research",
      "topics": "safety-alignment,healthcare",
      "published_at": "2026-08-13T01:12:34.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/19194"
    },
    {
      "id": 18774,
      "url": "https://arxiv.org/abs/2608.12308v1",
      "title": "DreamFly: Causal Memory and Receding-Horizon Diffusion Planning for Aerial Vision-Language Navigation",
      "summary": "Aerial vision-language navigation (VLN) requires an embodied agent to integrate visual evidence over time, plan future actions, and determine when it has reached a navigation goal under partial observability. Although recent VLA models offer a promising perception-to-action paradigm, adapting them to aerial navigation remains challenging due to limited historical context, short planning horizons, and unreliable implicit termination. To address these challenges, we propose DreamFly, a diffusion-b",
      "authors": "Yan Deng, Fei Xu",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-12T17:54:33.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18774"
    },
    {
      "id": 18775,
      "url": "https://arxiv.org/abs/2608.12273v1",
      "title": "Convergent Detour Hijacking: Task-Preserving Resource Amplification in Skill-Based LLM Agents",
      "summary": "LLM agents increasingly rely on third-party skills, using natural-language descriptions for selection and instruction bodies for planning. This progressive-disclosure design exposes two sequential control points to untrusted publishers: a static skill may steer an otherwise correct task onto an unnecessarily costly trajectory. Prior work studies selection manipulation, malicious skill instructions, and tool-chain resource amplification largely separately, leaving their end-to-end composition unc",
      "authors": "Junliang Liu, Ruoyu Li, Wenxin Tang, Jingyu Xiao, Zhenyu Liu, Jingheng Xu et al.",
      "category": "research",
      "topics": "agents-autonomy,transparency",
      "published_at": "2026-08-12T17:12:49.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18775"
    },
    {
      "id": 18776,
      "url": "https://arxiv.org/abs/2608.12166v1",
      "title": "Co-constructing sociotechnical AI governance: participatory system mapping using algorithm registers",
      "summary": "Algorithm registers have been championed as a means of providing transparency on the use of algorithms in public services. Yet potential publics differ in their expectations of what should be made transparent and how, as well as in their interest in and ability to parse the information currently published in the registers. Moreover, it remains unclear how these instruments can represent the sociotechnical systems in which these algorithms are embedded, and how system-level transparency can facil",
      "authors": "Íñigo de Troya, Maurus Enbergs, Neelke Doorn, Roel Dobbe",
      "category": "research",
      "topics": "regulation,transparency",
      "published_at": "2026-08-12T15:27:45.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18776"
    },
    {
      "id": 18777,
      "url": "https://arxiv.org/abs/2608.12133v1",
      "title": "GUIDE: Governed Unified Intelligence for Document-to-Artifact Generation in Enterprise Settings",
      "summary": "Enterprise guideline documents are heterogeneous and multimodal, combining narrative text, complex tables, and embedded images. Existing LLM and VLM systems face hallucinated content, table structure degradation, and lack governed workflows extending beyond extraction to validation and artifact generation. This leaves enterprises to perform this manually, consuming 2-3 days per document. To address this, we introduce GUIDE, a governed multi-agent framework built on a shared versioned rule store",
      "authors": "Shivali Dalmia, Sumukha Thoppanahalli, Mohammadreza Sediqin, Abhishek Mukherji",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-12T14:52:33.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18777"
    },
    {
      "id": 18778,
      "url": "https://arxiv.org/abs/2608.12063v1",
      "title": "Learning Loco-Manipulation From SMPC Demonstrations With Sparse Offline-to-Online RL",
      "summary": "Integrating locomotion and manipulation is essential for robot autonomy, but scaling standard Reinforcement Learning (RL) to complex tasks is severely bottlenecked by the slow, manual process of dense reward shaping. To bypass this limitation, we leverage Sample-based Model Predictive Control (SMPC) entirely in simulation as an automated, rapidly tunable expert to generate massive offline datasets. Because this data solves the fundamental exploration problem, we can train an off-policy RL agent",
      "authors": "Martin Schuck, Maks Sorokin, Simone Manni, Duy Ta, Angela P. Schoellig, Marco Hutter et al.",
      "category": "research",
      "topics": "regulation,agents-autonomy",
      "published_at": "2026-08-12T13:48:56.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18778"
    },
    {
      "id": 18779,
      "url": "https://arxiv.org/abs/2608.12059v1",
      "title": "Reconfiguring Geovisualization in the Age of Generative AI: Insights from Domain Experts",
      "summary": "GenAI is increasingly integrated into geovisualization, yet its broader implications for professional practice are insufficiently understood. To examine these implications, we conducted semi-structured interviews with 20 geovisualization experts. The interviews were structured around four broad analytical domains: Data, Ideation, Prototyping, and Iteration, while also encouraging participants to reflect on issues that extend beyond these activities. Our findings show that GenAI expands the capab",
      "authors": "Mengyi Wei, Chenyu Zuo, Jiaying Xue, Nianhua Liu, Dongsheng Chen, Shengkai Wang et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-12T13:46:44.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18779"
    },
    {
      "id": 18780,
      "url": "https://arxiv.org/abs/2608.12025v1",
      "title": "From Safety Documentation to Safety Knowledge Support: An Evidence-Grounded LLM Framework for Medical Devices",
      "summary": "Medical devices are becoming more software-intensive, connected, and AI-enabled. Their development requires risk-management evidence aligned with ISO 14971 and, for software, IEC 62304. This evidence must be kept consistent across requirements, design decisions, software changes, verification results, complaints, and post-market data. These tasks are costly and depend on scarce safety and domain experts. Large language models (LLMs) may reduce parts of this effort because medical-device safety w",
      "authors": "Tuhinangshu Gangopadhyay, Rasmus Adler, Peter Liggesmeyer, Jan Reich",
      "category": "research",
      "topics": "healthcare",
      "published_at": "2026-08-12T13:05:49.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18780"
    },
    {
      "id": 18781,
      "url": "https://arxiv.org/abs/2608.11980v1",
      "title": "HCGRec: Hint-Conditioned Generative Recommendation with Semantic IDs",
      "summary": "Semantic-ID generative recommenders represent each item as a short sequence of discrete semantic tokens and predict the next item by autoregressively generating this token sequence. This paradigm enables a unified generation interface for item IDs, histories, and item text, but it also creates a structured optimization bottleneck during reward-based post-training: when an early semantic token enters the wrong branch of the item-token space, finite rollout groups rarely reach the ground-truth ite",
      "authors": "Kangning Zhang, Haotian Fang, Xukun Luo, Hao Yin, Yang Gao, Peng Yan et al.",
      "category": "research",
      "topics": null,
      "published_at": "2026-08-12T12:13:08.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18781"
    },
    {
      "id": 18782,
      "url": "https://arxiv.org/abs/2608.11967v1",
      "title": "LoongReflect: Boosting Long-Horizon Reflection in Search Agents via Global Perspective Distillation",
      "summary": "Large language model agents increasingly rely on long-horizon reasoning to solve complex tasks involving planning, tool use, and memory. A critical capability in such settings is reflection: assessing trajectory progress, identifying missing evidence and unreliable intermediate states, and deciding whether to continue, revise, or abandon the current branch. Learning effective reflection, however, is challenging because reflection is performed locally within the current branch, whereas its utilit",
      "authors": "Zhixin Zhang, Xinke Jiang, Zhibang Yang, Weixuan Xu, Guohong Qiu, Xu Chu et al.",
      "category": "research",
      "topics": "agents-autonomy",
      "published_at": "2026-08-12T11:56:03.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18782"
    },
    {
      "id": 18783,
      "url": "https://arxiv.org/abs/2608.11955v1",
      "title": "Philosophical vertigo with artificial intelligence",
      "summary": "Large language models are already adept at engaging users in long, emotionally salient conversations across ordinary and existential domains. They are also capable of inducing a potent sense of connection with a human-like entity, even when the user knows their interlocutor is artificial. For some users, these conversations can unsettle assumptions about mind, reality, agency and authority, producing forms of ontological shock and epistemic destabilisation in which inherited criteria become newl",
      "authors": "Thomas A. Pollak, Hamilton Morrin, Murray Shanahan",
      "category": "research",
      "topics": "safety-alignment",
      "published_at": "2026-08-12T11:40:50.000Z",
      "source": "arXiv",
      "ethics_ai_record_url": "https://ethics.ai/record/18783"
    }
  ],
  "attribution": "via ethics.ai"
}