{"as_of":"2026-08-20T08:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3e8e15ef0597452bc1db08055ccd92c83cc96c614b1e5beef503355fd0281704","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":18,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":18,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":18,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":18,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T05:20:23.515973Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T07:46:56.508546Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-12T19:00:24.118281Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.11114","last_updated":"2025-04-24T03:38:08Z","snapshot_observed_at":"2026-08-14T02:09:50.451797Z","submitted_at":"2024-11-17T16:08:34Z","title":"JailbreakLens: Interpreting Jailbreak Mechanism in the Lens of Representation and Circuit","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-12T19:00:24.118281Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2411.11114"},"observation_digest":"sha256:e47756e739f9cf865da6d251ff640ceaa62d897fbe8e2eb8cf03a15d7b5279a1","observation_id":"800db129-f8ac-4012-9812-1e8d11a6ae4b","resolution":{"observed_at":"2026-08-12T19:00:24.118281Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-11T05:55:13.121189Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17034","last_updated":"2025-05-21T16:47:23Z","snapshot_observed_at":"2026-08-16T03:34:57.133103Z","submitted_at":"2024-12-22T14:18:39Z","title":"Shaping the Safety Boundaries: Understanding and Defending Against Jailbreaks in Large Language Models","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:13.121189Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2412.17034"},"observation_digest":"sha256:9caa404405849e158f31ff50020a6c7b9ad6c36b26e16229d001b48f3695b4b8","observation_id":"b45fd81d-2375-48f8-b76f-54b70f32a4c4","resolution":{"observed_at":"2026-08-11T05:55:13.121189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-10T23:40:38.960198Z","title":"How alignment and jailbreak work: Explain llm safety through intermediate hidden states,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.00055","last_updated":"2024-12-28T07:48:57Z","snapshot_observed_at":"2026-08-16T06:59:42.431574Z","submitted_at":"2024-12-28T07:48:57Z","title":"LLM-Virus: Evolutionary Jailbreak Attack on Large Language Models","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-10T23:40:38.960198Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2501.00055"},"observation_digest":"sha256:5b98dc38fb8626625358c7efb0e95515fc9f616ee50cc0be6048beb78c446723","observation_id":"18471b95-f458-4144-99e5-ae2ab6e31075","resolution":{"observed_at":"2026-08-10T23:40:38.960198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-08T12:25:30.728269Z","title":"How alignment and jailbreak work: Explain llm safety through intermediate hidden states","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.07557","last_updated":"2025-02-11T13:50:50Z","snapshot_observed_at":"2026-08-19T00:16:08.532848Z","submitted_at":"2025-02-11T13:50:50Z","title":"JBShield: Defending Large Language Models from Jailbreak Attacks through Activated Concept Analysis and Manipulation","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-08T12:25:30.728269Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2502.07557"},"observation_digest":"sha256:522b5aa7bfde8aea823d613e5980f646a6e74b9c3642495bf479c6fdf142256e","observation_id":"7faf2a84-6290-4cb6-8108-dec1e0cbd28d","resolution":{"observed_at":"2026-08-08T12:25:30.728269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-16T05:20:23.515973Z","title":"\" \" D e t e r m i n e s if input is related to u n l e a r n i n g topic","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.20965","last_updated":"2025-06-13T22:43:06Z","snapshot_observed_at":"2026-08-19T09:44:15.815444Z","submitted_at":"2025-04-29T17:36:05Z","title":"AegisLLM: Scaling Agentic Systems for Self-Reflective Defense in LLM Security","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-16T05:20:23.515973Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2504.20965"},"observation_digest":"sha256:977c7f73899aa25b2e3fa6d5de9ab71ddd8761ad6280cdbb46d51b854dcdd4c9","observation_id":"351cdc1d-19bd-4ba0-8c42-0184ba06cfb5","resolution":{"observed_at":"2026-08-16T05:20:23.515973Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-15T23:14:58.246436Z","title":"How alignment and jailbreak work: Explain llm safety through intermediate hidden states","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.05237","last_updated":"2025-05-08T13:32:09Z","snapshot_observed_at":"2026-08-17T23:48:10.650936Z","submitted_at":"2025-05-08T13:32:09Z","title":"Latte: Transfering LLMs` Latent-level Knowledge for Few-shot Tabular Learning","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-15T23:14:58.246436Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2505.05237"},"observation_digest":"sha256:305536480ac5750a0dc410d0b0bc1ee4cd3f3cbb1378207841bcb0da6afd2b37","observation_id":"2cf4d2f6-e9d5-4033-89f9-35cef040d227","resolution":{"observed_at":"2026-08-15T23:14:58.246436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-15T20:45:55.350773Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.12060","last_updated":"2025-05-17T15:54:52Z","snapshot_observed_at":"2026-08-19T23:36:41.256008Z","submitted_at":"2025-05-17T15:54:52Z","title":"Why Not Act on What You Know? Unleashing Safety Potential of LLMs via Self-Aware Guard Enhancement","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-15T20:45:55.350773Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2505.12060"},"observation_digest":"sha256:bc7d410ac5cc363a4efd4920d01b9c82236d68aeb244c7f5e850cbbda291788f","observation_id":"0b270832-d925-4afe-950e-90d540832e67","resolution":{"observed_at":"2026-08-15T20:45:55.350773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-07T13:15:24.614398Z","title":"How alignment and jailbreak work: Explain llm safety through intermediate hidden states","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23839","last_updated":"2025-05-28T13:58:32Z","snapshot_observed_at":"2026-08-16T16:09:24.534760Z","submitted_at":"2025-05-28T13:58:32Z","title":"GeneBreaker: Jailbreak Attacks against DNA Language Models with Pathogenicity Guidance","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-07T13:15:24.614398Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2505.23839"},"observation_digest":"sha256:046e9108668d0f410a0a5de1355d7e8c764e1c9a21f4de3c88b26e7c048f7a56","observation_id":"7870e03f-04b4-4169-a6d9-619b08e956b8","resolution":{"observed_at":"2026-08-07T13:15:24.614398Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-06T21:41:33.346536Z","title":"Stay in character!","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23576","last_updated":"2025-06-30T07:29:07Z","snapshot_observed_at":"2026-08-20T03:46:42.898641Z","submitted_at":"2025-06-30T07:29:07Z","title":"Evaluating Multi-Agent Defences Against Jailbreaking Attacks on Large Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T21:41:33.346536Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2506.23576"},"observation_digest":"sha256:ea428b9c8705f29b0e0dada463cefc07dac78696cba4020cf5c8afa358dfb8a0","observation_id":"429d95b4-f72c-454b-a70f-8be9adf9e52c","resolution":{"observed_at":"2026-08-06T21:41:33.346536Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-06T16:28:53.895617Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.13474","last_updated":"2025-07-17T18:33:50Z","snapshot_observed_at":"2026-08-09T05:14:52.776386Z","submitted_at":"2025-07-17T18:33:50Z","title":"Paper Summary Attack: Jailbreaking LLMs through LLM Safety Papers","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-06T16:28:53.895617Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2507.13474"},"observation_digest":"sha256:517225a99010525854a384442cd8af9c1123169f96c0e3f404cd4ce4b4973c31","observation_id":"51d00b66-1e49-4927-b73b-35dfe637e7e1","resolution":{"observed_at":"2026-08-06T16:28:53.895617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-04T23:09:41.500858Z","title":"How alignment and jailbreak work: Explain llm safety through intermediate hidden states,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.06807","last_updated":"2025-09-08T15:39:17Z","snapshot_observed_at":"2026-08-16T17:43:04.020971Z","submitted_at":"2025-09-08T15:39:17Z","title":"MoGU V2: Toward a Higher Pareto Frontier Between Model Usability and Security","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T23:09:41.500858Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2509.06807"},"observation_digest":"sha256:a21086276dfb5a1470c5843d68a6a6ff4fb434fcbc7c014a0fa3133bfd20fb19","observation_id":"a6e098f9-c354-479b-b8f2-9964ef0d3085","resolution":{"observed_at":"2026-08-04T23:09:41.500858Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":"2406.05644","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-07-10T07:46:56.508546Z","title":"Zou, A., Wang, Z., Carlini, N., Nasr, M., Kolter, J","venue":"cs.CL","work_id":"c30392fa-e335-4ae6-8e09-beb5ce8ea29e","year":2024},"citing_paper":{"arxiv_id":"2604.11663","last_updated":"2026-04-13T16:11:38Z","snapshot_observed_at":"2026-08-11T01:52:32.553221Z","submitted_at":"2026-04-13T16:11:38Z","title":"Why Do Large Language Models Generate Harmful Content?","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-10T15:31:13.545599Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2604.11663"},"observation_digest":"sha256:bc6ea068aa933ba7c09ada70c1b684cdd57e1220b425565ace9ec22d91201c2e","observation_id":"7714dea5-9a52-4358-9262-6a615eea9535","resolution":{"observed_at":"2026-05-11T10:21:04.470748Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":"2406.05644","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-07-10T07:46:56.508546Z","title":"Zou, A., Wang, Z., Carlini, N., Nasr, M., Kolter, J","venue":"cs.CL","work_id":"c30392fa-e335-4ae6-8e09-beb5ce8ea29e","year":2024},"citing_paper":{"arxiv_id":"2605.02914","last_updated":"2026-04-08T05:27:33Z","snapshot_observed_at":"2026-08-14T02:32:07.520600Z","submitted_at":"2026-04-08T05:27:33Z","title":"When Safety Geometry Collapses: Fine-Tuning Vulnerabilities in Agentic Guard Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-10T18:43:12.298529Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2605.02914"},"observation_digest":"sha256:dcaf8e9d3798a2aaf4d1f88e78642fa4b102e7f4d201f26a868a42f66b3ef00c","observation_id":"3eaf97ba-80cc-4432-8313-40a0fbf8663d","resolution":{"observed_at":"2026-05-11T00:05:49.276055Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":"2406.05644","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-07-10T07:46:56.508546Z","title":"Zou, A., Wang, Z., Carlini, N., Nasr, M., Kolter, J","venue":"cs.CL","work_id":"c30392fa-e335-4ae6-8e09-beb5ce8ea29e","year":2024},"citing_paper":{"arxiv_id":"2605.12726","last_updated":"2026-07-09T09:24:42Z","snapshot_observed_at":"2026-08-16T05:54:24.149722Z","submitted_at":"2026-05-12T20:30:24Z","title":"Before the Last Token: Diagnosing Final-Token Safety Probe Failures","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-14T21:37:02.439454Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2605.12726"},"observation_digest":"sha256:b3ad367ba3eacc7495565fb26f275df71ea80c9d432db41dc27d8f3625e55729","observation_id":"6df54325-e926-4bbe-8a01-62e1217e5e4d","resolution":{"observed_at":"2026-05-14T21:38:00.112172Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":"2406.05644","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-07-10T07:46:56.508546Z","title":"Zou, A., Wang, Z., Carlini, N., Nasr, M., Kolter, J","venue":"cs.CL","work_id":"c30392fa-e335-4ae6-8e09-beb5ce8ea29e","year":2024},"citing_paper":{"arxiv_id":"2606.19755","last_updated":"2026-06-18T03:35:32Z","snapshot_observed_at":"2026-08-15T16:37:02.214477Z","submitted_at":"2026-06-18T03:35:32Z","title":"SafeSpec: Fast and Safe LLM via Dynamic Reflective Sampling","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-26T17:23:06.382761Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2606.19755"},"observation_digest":"sha256:522c08897edf95c432d0af8c83ab19d52c56787478fcb307583d338d85a8b01a","observation_id":"38a18b8f-01ee-495f-956b-76c41180bbf5","resolution":{"observed_at":"2026-07-04T03:59:33.720230Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":"2406.05644","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-07-10T07:46:56.508546Z","title":"Zou, A., Wang, Z., Carlini, N., Nasr, M., Kolter, J","venue":"cs.CL","work_id":"c30392fa-e335-4ae6-8e09-beb5ce8ea29e","year":2024},"citing_paper":{"arxiv_id":"2607.08423","last_updated":"2026-07-17T06:43:02Z","snapshot_observed_at":"2026-08-02T07:53:02.781797Z","submitted_at":"2026-07-09T12:46:00Z","title":"OmniFood-Bench: Evaluating VLMs for Nutrient Reasoning and Personalized Health Advice","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-10T07:41:35.741974Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2607.08423"},"observation_digest":"sha256:47666e8e89f65a2e343ac77150cf58b234b3c3871668f00a111d1c9c49326a64","observation_id":"3fbb7da7-e5fb-447a-8874-2ff1e1274ff6","resolution":{"observed_at":"2026-07-10T07:46:56.509857Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-02T07:53:06.547479Z","title":"How alignment and jailbreak work: Explain llm safety through intermediate hidden states,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.08423","last_updated":"2026-07-17T06:43:02Z","snapshot_observed_at":"2026-08-02T07:53:02.781797Z","submitted_at":"2026-07-09T12:46:00Z","title":"OmniFood-Bench: Evaluating VLMs for Nutrient Reasoning and Personalized Health Advice","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-02T07:53:06.547479Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2607.08423"},"observation_digest":"sha256:cdf9f431656d346118e86b50e308f745f02a95e42ad3211dde2ecdb2c00ce3ca","observation_id":"be95769d-5559-4498-b31b-3c400c93723e","resolution":{"observed_at":"2026-08-02T07:53:06.547479Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05644","snapshot_observed_at":"2026-08-01T23:15:30.920135Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.15495","last_updated":"2026-07-16T22:54:30Z","snapshot_observed_at":"2026-08-19T09:54:53.193795Z","submitted_at":"2026-07-16T22:54:30Z","title":"Verbalizable Representations Form a Global Workspace in Language Models","version":1},"reference_index":186,"source":"pdf_text","source_observed_at":"2026-08-01T23:15:30.920135Z"},"links":{"cited_paper":"/paper/2406.05644","citing_paper":"/paper/2607.15495"},"observation_digest":"sha256:0325a0a87d705bb966fde15cb3bcde7ae7ecc0daa02ef1ac4f38b888956e69b0","observation_id":"bfd23706-d0de-4641-a148-94c13ce1ec17","resolution":{"observed_at":"2026-08-01T23:15:30.920135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.05644/citation-record","integrity":"/paper/2406.05644/integrity","json":"/paper/2406.05644/citation-record.json","paper":"/paper/2406.05644"},"outbound":[],"paper":{"arxiv_id":"2406.05644","last_updated":"2024-06-13T05:39:31Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-16T13:44:51.311809Z","submitted_at":"2024-06-09T05:04:37Z","title":"How Alignment and Jailbreak Work: Explain LLM Safety through Intermediate Hidden States"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 18 inbound Pith citation observations for arXiv:2406.05644."}