{"data":{"id":"231a3783-0aa2-4d29-bf8f-2d31a1fe666c","title":"Open Sourcing Monitorability Evaluations","summary":"OpenAI released datasets and reference code from its Monitoring Monitorability paper, covering most of its chain-of-thought monitorability evaluation suite, code for computing the g-mean 2 metric, and a cross-fit filtering strategy for noise-dominated intervention instances. Several evals were omitted because they rely on private or restricted data. The scaffold included in the release is an illustration only, and OpenAI will not support it.","solution":"N/A -- no mitigation discussed in source.","labels":["safety","research"],"sourceUrl":"https://alignment.openai.com/monitorability-evals/","publishedAt":"2026-04-23T22:15:00.000Z","cveId":null,"cweIds":null,"cvssScore":null,"cvssSeverity":null,"severity":"info","attackType":[],"issueType":"research","affectedPackages":null,"affectedPackageNames":null,"affectedPackageRefs":null,"affectedVendors":["OpenAI"],"affectedVendorsRaw":["GPT-5.4 thinking","GPT-5.2 thinking","GPT-5 thinking","o3","o1"],"classifierModel":"claude-haiku-5-5","classifierPromptVersion":"v4","summaryPromptVersion":"v2","headline":null,"headlinePromptVersion":null,"cvssVector":null,"attackVector":null,"attackComplexity":null,"privilegesRequired":null,"userInteraction":null,"exploitMaturity":null,"epssScore":null,"epssCheckedAt":null,"kevDateAdded":null,"advisoryAliases":null,"affectedPackagesSource":null,"affectedPackagesCheckedAt":null,"patchAvailable":null,"disclosureDate":"2026-04-23T22:15:00.000Z","capecIds":null,"crossRefCount":0,"attackSophistication":"moderate","impactType":null,"aiComponentTargeted":"model","llmSpecific":true,"classifierConfidence":0.9,"researchCategory":"industry","atlasIds":null}}