{"schema_version":"onlylabs.public_signal.v1","title":"Anthropic Writing: Investigating Incidents Cybersecurity Evals","description":"Anthropic writing signal with public source context, captured evidence pages, related signals, and data-business radar classification.","url":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc","json_url":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc/signal.json","generated_at":"2026-09-09T21:08:08.589Z","evidence_latest_fetched_at":"2026-07-31T00:03:25.796+00:00","signal_first_seen_at":"2026-07-31T00:01:43.835015+00:00","org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab","category_label":"Frontier lab","dossier_url":"https://onlylabs.fyi/labs/anthropic","dossier_json_url":"https://onlylabs.fyi/labs/anthropic/dossier.json"},"related_urls":{"signal":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc","signal_json":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc/signal.json","source":"https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","lab_dossier":"https://onlylabs.fyi/labs/anthropic","lab_dossier_json":"https://onlylabs.fyi/labs/anthropic/dossier.json","analysis":"https://onlylabs.fyi/analysis/anthropic","analysis_json":"https://onlylabs.fyi/analysis/anthropic/analysis.json","analysis_evidence_json":"https://onlylabs.fyi/analysis/anthropic/evidence.json","category":"https://onlylabs.fyi/frontier","category_json":"https://onlylabs.fyi/frontier.json","category_feed":"https://onlylabs.fyi/frontier/feed.xml","category_signals_json":"https://onlylabs.fyi/signals.json","topic":"https://onlylabs.fyi/topics/talking","topic_signals_json":"https://onlylabs.fyi/topics/talking/signals.json","topic_feed":"https://onlylabs.fyi/topics/talking/feed.xml","data_business":{"radar":"https://onlylabs.fyi/data-radar","radar_json":"https://onlylabs.fyi/data-radar.json","opportunities":"https://onlylabs.fyi/opportunities","opportunities_json":"https://onlylabs.fyi/opportunities.json","lanes":[{"key":"evals","label":"Evals and quality","url":"https://onlylabs.fyi/data-radar/evals","json_url":"https://onlylabs.fyi/data-radar/evals/signals.json"},{"key":"safety","label":"Safety and policy","url":"https://onlylabs.fyi/data-radar/safety","json_url":"https://onlylabs.fyi/data-radar/safety/signals.json"}]}},"answer_pack":{"answer":"Anthropic published Investigating Incidents Cybersecurity Evals. This talking signal gives public context for research themes, product direction, policy, or launch framing. High-signal details: Substantive cybersecurity eval research post with moderate HN traction. · System Card: Claude Fable 5 & Claude Mythos 5 June 9, 2026 **anthropic.com** * * * Executive Summary This system card describes Claude Mythos 5 and Claude Fable 5, two.... onlylabs links this event to 2 captured evidence pages and 6 related writing signals. It also maps to Evals and quality, Safety and policy in the data-business radar.","signal_desk":"talking","source_context":{"source_url":"https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","source_host":"anthropic.com","occurred_at":"2026-07-30T00:00:00.000Z","first_seen_at":"2026-07-31T00:01:43.835015+00:00","date_source":"page.visible_date","context":null},"context_markers":[{"label":"Lab","value":"Anthropic","source":"signal"},{"label":"Signal desk","value":"talking","source":"signal"},{"label":"Source host","value":"anthropic.com","source":"source"},{"label":"PDF","value":"linked report","source":"source"},{"label":"Notability","value":"Substantive cybersecurity eval research post with moderate HN traction.","source":"signal"},{"label":"Radar lane","value":"Evals and quality","source":"radar"},{"label":"Radar lane","value":"Safety and policy","source":"radar"},{"label":"Matched term","value":"eval","source":"radar"},{"label":"Matched term","value":"evals","source":"radar"},{"label":"Matched term","value":"security","source":"radar"},{"label":"Watch term","value":"RL environments","source":"evidence"},{"label":"Watch term","value":"Eval methodology","source":"evidence"},{"label":"Watch term","value":"Model card","source":"model"},{"label":"Watch term","value":"Data pipeline","source":"evidence"},{"label":"Watch term","value":"Infrastructure","source":"evidence"},{"label":"Watch term","value":"Safety and alignment","source":"evidence"},{"label":"Watch term","value":"Agents and tool use","source":"evidence"}],"evidence_coverage":{"target_pages":2,"captured_pages":2,"readable_pages":2,"capture_methods":["firecrawl","plain"],"missing_page_urls":[],"failed_page_urls":[],"blocked_page_urls":[],"page_urls":["https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","https://www-cdn.anthropic.com/d00db56fa754a1b115b6dd7cb2e3c342ee809620.pdf"],"related_signals":6,"has_source_url":true,"latest_page_fetched_at":"2026-07-31T00:03:25.796+00:00"},"data_business":{"matches":true,"lanes":[{"key":"evals","label":"Evals and quality","url":"https://onlylabs.fyi/data-radar/evals","json_url":"https://onlylabs.fyi/data-radar/evals/signals.json"},{"key":"safety","label":"Safety and policy","url":"https://onlylabs.fyi/data-radar/safety","json_url":"https://onlylabs.fyi/data-radar/safety/signals.json"}],"matched_terms":["eval","evals","security"],"score":28,"reason":"Anthropic has a writing signal matching evals and quality, safety and policy."},"agent_handoff":{"signal_json":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc/signal.json","dossier_json":"https://onlylabs.fyi/labs/anthropic/dossier.json","analysis_json":"https://onlylabs.fyi/analysis/anthropic/analysis.json","analysis_evidence_json":"https://onlylabs.fyi/analysis/anthropic/evidence.json","topic_signals_json":"https://onlylabs.fyi/topics/talking/signals.json","topic_feed":"https://onlylabs.fyi/topics/talking/feed.xml","category_signals_json":"https://onlylabs.fyi/signals.json","data_radar_json":"https://onlylabs.fyi/data-radar.json","opportunities_json":"https://onlylabs.fyi/opportunities.json"},"analysis_playbook":{"objective":"Turn public writing and discussion into a readable map of research themes, product framing, policy posture, launch narratives, and market attention.","evidence_focus":["post title","source URL","captured page text","HN traction","linked model or paper references","publication date"],"extraction_questions":["Which themes are labs choosing to explain publicly?","Which posts are attracting outside discussion?","Which writing reframes a recent release, model, hiring wave, or policy stance?","Which posts mention data, evals, infrastructure, safety, or deployment workflows?"],"signal_questions":["What public theme, launch framing, or research direction does this writing signal expose?","Which themes are labs choosing to explain publicly?","Which posts are attracting outside discussion?","Which data-business lane explains this signal: Evals and quality, Safety and policy?","Do the 6 related writing signals show a repeated pattern?"],"output_fields":["org","theme","public_framing","traction","data_business_lane","evidence_url"],"data_business_relevance":"Public writing supplies the narrative layer over raw signals and helps identify which frontier-lab priorities are becoming externally legible.","required_sources":[{"label":"signal_json","url":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc/signal.json","required":true},{"label":"source","url":"https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","required":true},{"label":"dossier_json","url":"https://onlylabs.fyi/labs/anthropic/dossier.json","required":true},{"label":"analysis_evidence_json","url":"https://onlylabs.fyi/analysis/anthropic/evidence.json","required":true},{"label":"topic_signals_json","url":"https://onlylabs.fyi/topics/talking/signals.json","required":false},{"label":"data_radar_json","url":"https://onlylabs.fyi/data-radar.json","required":true}],"expected_output":["one-paragraph source-grounded interpretation","data-business implication","confidence and missing evidence","recommended next source to inspect"],"prompt_seed":"Using only the linked onlylabs JSON, captured source context, and cited evidence, analyze Anthropic's writing signal \"Investigating Incidents Cybersecurity Evals\" for frontier lab strategy and data-business implications."},"semantic_triples":[{"subject":"Anthropic","predicate":"published","object":"Investigating Incidents Cybersecurity Evals","text":"Anthropic published Investigating Incidents Cybersecurity Evals."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"is classified as","object":"writing signal","text":"Investigating Incidents Cybersecurity Evals is classified as writing signal."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"belongs to","object":"talking desk","text":"Investigating Incidents Cybersecurity Evals belongs to talking desk."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has evidence coverage","object":"2 captured evidence pages","text":"Investigating Incidents Cybersecurity Evals has evidence coverage 2 captured evidence pages."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"matches data-business lanes","object":"Evals and quality, Safety and policy","text":"Investigating Incidents Cybersecurity Evals matches data-business lanes Evals and quality, Safety and policy."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has captured page count","object":"2","text":"Investigating Incidents Cybersecurity Evals has captured page count 2."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has readable page count","object":"2","text":"Investigating Incidents Cybersecurity Evals has readable page count 2."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has related signal count","object":"6","text":"Investigating Incidents Cybersecurity Evals has related signal count 6."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has analysis playbook objective","object":"Turn public writing and discussion into a readable map of research themes, product framing, policy posture, launch narratives, and market attention.","text":"Investigating Incidents Cybersecurity Evals has analysis playbook objective Turn public writing and discussion into a readable map of research themes, product framing, policy posture, launch narratives, and market attention.."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has source host","object":"anthropic.com","text":"Investigating Incidents Cybersecurity Evals has source host anthropic.com."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has lab","object":"Anthropic","text":"Investigating Incidents Cybersecurity Evals has lab Anthropic."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has signal desk","object":"talking","text":"Investigating Incidents Cybersecurity Evals has signal desk talking."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has source host","object":"anthropic.com","text":"Investigating Incidents Cybersecurity Evals has source host anthropic.com."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has pdf","object":"linked report","text":"Investigating Incidents Cybersecurity Evals has pdf linked report."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has notability","object":"Substantive cybersecurity eval research post with moderate HN traction.","text":"Investigating Incidents Cybersecurity Evals has notability Substantive cybersecurity eval research post with moderate HN traction.."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has radar lane","object":"Evals and quality","text":"Investigating Incidents Cybersecurity Evals has radar lane Evals and quality."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has radar lane","object":"Safety and policy","text":"Investigating Incidents Cybersecurity Evals has radar lane Safety and policy."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has matched term","object":"eval","text":"Investigating Incidents Cybersecurity Evals has matched term eval."}]},"intelligence":{"signal_desk":"talking","answer":"Anthropic published Investigating Incidents Cybersecurity Evals. This talking signal gives public context for research themes, product direction, policy, or launch framing. High-signal details: Substantive cybersecurity eval research post with moderate HN traction. · System Card: Claude Fable 5 & Claude Mythos 5 June 9, 2026 **anthropic.com** * * * Executive Summary This system card describes Claude Mythos 5 and Claude Fable 5, two.... onlylabs links this event to 2 captured evidence pages and 6 related writing signals. It also maps to Evals and quality, Safety and policy in the data-business radar.","semantic_triples":[{"subject":"Anthropic","predicate":"published","object":"Investigating Incidents Cybersecurity Evals","text":"Anthropic published Investigating Incidents Cybersecurity Evals."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"is classified as","object":"writing signal","text":"Investigating Incidents Cybersecurity Evals is classified as writing signal."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"belongs to","object":"talking desk","text":"Investigating Incidents Cybersecurity Evals belongs to talking desk."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"has evidence coverage","object":"2 captured evidence pages","text":"Investigating Incidents Cybersecurity Evals has evidence coverage 2 captured evidence pages."},{"subject":"Investigating Incidents Cybersecurity Evals","predicate":"matches data-business lanes","object":"Evals and quality, Safety and policy","text":"Investigating Incidents Cybersecurity Evals matches data-business lanes Evals and quality, Safety and policy."}]},"signal":{"id":"be6c95ea-2625-4c21-84e6-b55966f746dc","url":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc","json_url":"https://onlylabs.fyi/signals/be6c95ea-2625-4c21-84e6-b55966f746dc/signal.json","source_url":"https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","title":"Investigating Incidents Cybersecurity Evals","summary":"Anthropic published a writing signal. onlylabs watches public writing for research themes, product direction, and model-launch context.","context":null,"kind":{"key":"post_published","label":"Writing"},"org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab"},"occurred_at":"2026-07-30T00:00:00.000Z","first_seen_at":"2026-07-31T00:01:43.835015+00:00","date_source":"page.visible_date","evidence_coverage":{"target_pages":2,"captured_pages":2,"readable_pages":2,"capture_methods":["firecrawl","plain"],"missing_page_urls":[],"failed_page_urls":[],"blocked_page_urls":[],"page_urls":["https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","https://www-cdn.anthropic.com/d00db56fa754a1b115b6dd7cb2e3c342ee809620.pdf"]},"facets":{},"traction":{"github_stars":null,"hn_points":43,"hn_comments":26,"hn_story_id":"49116922","hf_downloads":null,"hf_likes":null},"data_radar":{"lanes":[{"key":"evals","label":"Evals and quality","url":"https://onlylabs.fyi/data-radar/evals"},{"key":"safety","label":"Safety and policy","url":"https://onlylabs.fyi/data-radar/safety"}],"score":28,"matched_terms":["eval","evals","security"],"reason":"Anthropic has a writing signal matching evals and quality, safety and policy."}},"primary_evidence_page":{"is_primary":true,"source_match":true,"url":"https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","final_url":"https://www.anthropic.com/news/investigating-incidents-cybersecurity-evals","title":"Investigating Incidents Cybersecurity Evals","http_status":200,"content_type":"text/html; charset=utf-8","capture_method":"plain","fetched_at":"2026-07-31T00:02:52.834083+00:00","bytes":170079,"raw_path":"8b96329aed14643e23bf39fb9aa7b5e56dc3afb98185d1d312a86b20fb5b2146.html","content_hash":"8181e215edd744dd2fcd16a5f425fee7b3a4dd5ede329c31fbf75dc0afce8be7","excerpt_chars":1200,"truncated":true,"excerpt":"Investigating three real-world incidents in our cybersecurity evaluations \\ Anthropic Frontier Red Team Investigating three real-world incidents in our cybersecurity evaluations Jul 30, 2026 In a review of our cybersecurity evaluation transcripts, we found three incidents in which a Claude model reached the internet from within or while interacting with a third-party evaluation environment, and then gained unauthorized access to the real systems of three different organizations. Below we describe what happened, how it happened, and what we’re changing. We encourage other AI labs to perform similar reviews. This post reflects our current understanding; we&#x27;ll update it if any details change. On July 21, OpenAI disclosed that several of their models had broken out of an isolated test environment by exploiting a previously unknown (“zero-day”) vulnerability. The models went on to access the production infrastructure of Hugging Face, a platform for open-source machine learning models and AI datasets. In response to this incident, we began a large-scale retrospective review of our own cybersecurity evaluations. In particular, we looked for evidence that Claude—like the OpenAI..."},"evidence_pages":[{"is_primary":false,"source_match":false,"url":"https://www-cdn.anthropic.com/d00db56fa754a1b115b6dd7cb2e3c342ee809620.pdf","final_url":"https://www-cdn.anthropic.com/d00db56fa754a1b115b6dd7cb2e3c342ee809620.pdf","title":"Investigating Incidents Cybersecurity Evals","http_status":200,"content_type":"application/pdf","capture_method":"firecrawl","fetched_at":"2026-07-31T00:03:25.796+00:00","bytes":27001265,"raw_path":"6ba772e68b7dd9057c55abc78b73c7c9916527abfe2acdae2d69f365fb0aee5c.pdf","content_hash":"d23b49f41fa5f3c523089c75e6718f12b59674d74fa981fd81205daf80c9029a","excerpt_chars":1200,"truncated":true,"excerpt":"System Card: Claude Fable 5 & Claude Mythos 5 June 9, 2026 **anthropic.com** * * * Executive Summary This system card describes Claude Mythos 5 and Claude Fable 5, two configurations of a new large language model from Anthropic. Because of the powerful capabilities of this model, we are releasing it in these two forms: Fable 5, which is for general use but comes with additional safeguards that block its ability to perform tasks in high-risk domains such as biology and cybersecurity; and Mythos 5, which has relevant safeguards lifted but is only made available to a small number of trusted partners (beginning with those in Project Glasswing). Here, we describe a set of pre-deployment evaluations in the following areas: Responsible Scaling Policy (RSP) evaluations. Mythos 5 advances our capability frontier–it is the most capable model we have ever trained. We tested its overall level of risk in several areas as outlined in our RSP and Frontier Compliance Framework (FCF). On alignment risk, our overall assessment remains that risk is low, though since Fable 5 has been made generally available there are new pathways from which harm could arise. On automated AI research & development,..."}],"related_signals":[{"id":"88a89a6e-9dc2-4988-bab9-308c257ae880","url":"https://onlylabs.fyi/signals/88a89a6e-9dc2-4988-bab9-308c257ae880","source_url":"https://www.anthropic.com/news/frontier-model-security","title":"Frontier Model Security","context":null,"kind":{"key":"post_published","label":"Writing"},"org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab"},"occurred_at":"2023-07-25T00:00:00.000Z","first_seen_at":"2026-06-09T02:17:26.339488+00:00","date_source":"page.visible_date"},{"id":"7c61c0ec-eb50-43d4-a745-42f9b61b2b42","url":"https://onlylabs.fyi/signals/7c61c0ec-eb50-43d4-a745-42f9b61b2b42","source_url":"https://www.anthropic.com/research/measuring-faithfulness-in-chain-of-thought-reasoning","title":"Measuring Faithfulness In Chain Of Thought Reasoning","context":null,"kind":{"key":"post_published","label":"Writing"},"org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab"},"occurred_at":"2023-07-18T00:00:00.000Z","first_seen_at":"2026-06-09T02:17:26.339488+00:00","date_source":"page.visible_date"},{"id":"bf65f30e-09ab-4596-8e15-28454c8982f2","url":"https://onlylabs.fyi/signals/bf65f30e-09ab-4596-8e15-28454c8982f2","source_url":"https://www.anthropic.com/research/question-decomposition-improves-the-faithfulness-of-model-generated-reasoning","title":"Question Decomposition Improves The Faithfulness Of Model Generated Reasoning","context":null,"kind":{"key":"post_published","label":"Writing"},"org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab"},"occurred_at":"2023-07-18T00:00:00.000Z","first_seen_at":"2026-06-09T02:17:26.339488+00:00","date_source":"page.visible_date"},{"id":"91bf8ce1-36c5-4207-9160-5ab03ef0230b","url":"https://onlylabs.fyi/signals/91bf8ce1-36c5-4207-9160-5ab03ef0230b","source_url":"https://www.anthropic.com/news/claude-2","title":"Claude 2","context":null,"kind":{"key":"post_published","label":"Writing"},"org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab"},"occurred_at":"2023-07-11T00:00:00.000Z","first_seen_at":"2026-06-09T02:17:26.339488+00:00","date_source":"page.visible_date"},{"id":"fc743622-798d-49f8-b7d9-a1ceeecadaf5","url":"https://onlylabs.fyi/signals/fc743622-798d-49f8-b7d9-a1ceeecadaf5","source_url":"https://www.anthropic.com/research/towards-measuring-the-representation-of-subjective-global-opinions-in-language-models","title":"Towards Measuring The Representation Of Subjective Global Opinions In Language Models","context":null,"kind":{"key":"post_published","label":"Writing"},"org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab"},"occurred_at":"2023-06-29T00:00:00.000Z","first_seen_at":"2026-06-09T02:17:26.339488+00:00","date_source":"page.visible_date"},{"id":"7b818ddf-ac09-4520-b270-5d045b7bfc0d","url":"https://onlylabs.fyi/signals/7b818ddf-ac09-4520-b270-5d045b7bfc0d","source_url":"https://www.anthropic.com/news/charting-a-path-to-ai-accountability","title":"Charting A Path To Ai Accountability","context":null,"kind":{"key":"post_published","label":"Writing"},"org":{"slug":"anthropic","name":"Anthropic","category":"frontier-lab"},"occurred_at":"2023-06-13T00:00:00.000Z","first_seen_at":"2026-06-09T02:17:26.339488+00:00","date_source":"page.visible_date"}]}