{"data":{"id":"6bc65002-2f19-4eaf-82dd-407128c0c932","title":"Action-Level Backdoor Attacks Against Deep Reinforcement Learning Systems via Adaptive Reward Exploration","summary":"Researchers have demonstrated a new attack called Adapdoor that can inject hidden malicious behaviors into Deep Reinforcement Learning (DRL) models, which are AI systems trained to make sequential decisions in environments like robotics and autonomous vehicles. The attack works by poisoning the reward signal (the feedback that guides what the AI learns to do) during training, allowing attackers to later manipulate the model's actions when it is deployed in the real world. The paper shows this threat is more serious than previously thought because Adapdoor can work across many different tasks without requiring manual customization for each one.","solution":"N/A -- no mitigation discussed in source.","labels":["security","research"],"sourceUrl":"http://ieeexplore.ieee.org/document/11675897","publishedAt":"2026-09-02T13:16:59.000Z","cveId":null,"cweIds":null,"cvssScore":null,"cvssSeverity":null,"severity":"info","attackType":["model_poisoning"],"issueType":"research","affectedPackages":null,"affectedVendors":[],"affectedVendorsRaw":[],"classifierModel":"claude-haiku-4-5-20251001","classifierPromptVersion":"v3","cvssVector":null,"attackVector":null,"attackComplexity":null,"privilegesRequired":null,"userInteraction":null,"exploitMaturity":null,"epssScore":null,"patchAvailable":null,"disclosureDate":"2026-09-02T13:16:59.000Z","capecIds":null,"crossRefCount":0,"attackSophistication":"advanced","impactType":["integrity","safety"],"aiComponentTargeted":"model","llmSpecific":false,"classifierConfidence":0.92,"researchCategory":"peer_reviewed","atlasIds":null}}