{"as_of":"2026-08-09T18:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f96c03b25c0be27e7b633d105bf634e65bf3a188deb7859fdbf1d6bd775d32e7","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":28,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":28,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":28,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":28,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T11:42:07.741050Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-08T11:42:07.741050Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.07771","last_updated":"2025-02-11T18:55:57Z","snapshot_observed_at":"2026-08-08T11:35:32.227506Z","submitted_at":"2025-02-11T18:55:57Z","title":"Breaking Down Bias: On The Limits of Generalizable Pruning Strategies","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-08T11:42:07.741050Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2502.07771"},"observation_digest":"sha256:ec2a41fc9f14d68eb7a9742dd7d5b10bd8aacc19bb1e8949c4c316bb376660e0","observation_id":"0869e5b3-4271-45b0-9692-6b34579f48e1","resolution":{"observed_at":"2026-08-08T11:42:07.741050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-07T15:41:05.918299Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies.arXiv preprint arXiv:2407.17436, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.14215","last_updated":"2026-07-08T13:57:53Z","snapshot_observed_at":"2026-08-08T23:15:53.820935Z","submitted_at":"2025-05-20T11:21:40Z","title":"Safety Degradation in AI Agents","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T15:41:05.918299Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2505.14215"},"observation_digest":"sha256:cae26025cf056cc5df16e354459bd1b459185df9487206624ee66506855b9321","observation_id":"d7045213-da63-4230-becc-db1ecb4d8590","resolution":{"observed_at":"2026-08-07T15:41:05.918299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-07T15:36:27.682170Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.14633","last_updated":"2025-05-20T17:24:09Z","snapshot_observed_at":"2026-08-07T22:06:44.428098Z","submitted_at":"2025-05-20T17:24:09Z","title":"Will AI Tell Lies to Save Sick Children? Litmus-Testing AI Values Prioritization with AIRiskDilemmas","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T15:36:27.682170Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2505.14633"},"observation_digest":"sha256:1714b9bfa36759cbfa6a62041f4cde1ce4c87ef50bafe4b5ec43907252abf7bb","observation_id":"7adf61df-b8bb-4260-86ec-bb937963cec4","resolution":{"observed_at":"2026-08-07T15:36:27.682170Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-07T13:52:03.610676Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies.arXiv preprint arXiv:2407.17436, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20824","last_updated":"2025-05-27T07:34:40Z","snapshot_observed_at":"2026-08-08T13:22:09.551933Z","submitted_at":"2025-05-27T07:34:40Z","title":"MedSentry: Understanding and Mitigating Safety Risks in Medical LLM Multi-Agent Systems","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T13:52:03.610676Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2505.20824"},"observation_digest":"sha256:fec74f4f01e19fd96a19ffd3ac70d436a1c4e0c5b1f4c67fc4ef444e4feb58a1","observation_id":"acbf7256-9e73-4d4b-8725-487d51da02a4","resolution":{"observed_at":"2026-08-07T13:52:03.610676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-07T11:56:56.078889Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.00964","last_updated":"2025-08-09T18:50:29Z","snapshot_observed_at":"2026-08-08T11:49:49.886158Z","submitted_at":"2025-06-01T11:24:23Z","title":"ACCESS DENIED INC: The First Benchmark Environment for Sensitivity Awareness","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T11:56:56.078889Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2506.00964"},"observation_digest":"sha256:f6c383e4f88dee899ef39e1eaf2c57e452d26cca335cc10a10d7145a500ff76a","observation_id":"1b7e9ba2-be7e-4c7e-bfca-b8d16364fb28","resolution":{"observed_at":"2026-08-07T11:56:56.078889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-07T05:40:20.722684Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07402","last_updated":"2025-06-09T03:52:43Z","snapshot_observed_at":"2026-08-07T10:07:41.500171Z","submitted_at":"2025-06-09T03:52:43Z","title":"Beyond Jailbreaks: Revealing Stealthier and Broader LLM Security Risks Stemming from Alignment Failures","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T05:40:20.722684Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2506.07402"},"observation_digest":"sha256:1d47d9e44655da23efef2e5234ee5db2b114e525555b4d15922b3ff507890355","observation_id":"26835b23-1641-4939-ab42-09bdb134a4a5","resolution":{"observed_at":"2026-08-07T05:40:20.722684Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-06T23:00:19.627162Z","title":"Z., Tu, Y ., Mai, Y ., Kly- man, K., Pan, M., Jia, R., Song, D., et al","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.20251","last_updated":"2025-06-25T08:52:22Z","snapshot_observed_at":"2026-08-07T22:24:27.303298Z","submitted_at":"2025-06-25T08:52:22Z","title":"Q-resafe: Assessing Safety Risks and Quantization-aware Safety Patching for Quantized Large Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T23:00:19.627162Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2506.20251"},"observation_digest":"sha256:de26864b8ade7c40a06cc097bcb8bd91bfc186cbfd2014d3d21be6945411c13e","observation_id":"0f2d9e24-efd8-43c9-bf37-44c75b1e2e61","resolution":{"observed_at":"2026-08-06T23:00:19.627162Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-06T21:38:51.524460Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23725","last_updated":"2025-06-30T10:58:36Z","snapshot_observed_at":"2026-08-08T10:23:35.585137Z","submitted_at":"2025-06-30T10:58:36Z","title":"PAC Bench: Do Foundation Models Understand Prerequisites for Executing Manipulation Policies?","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T21:38:51.524460Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2506.23725"},"observation_digest":"sha256:201b25551840416b658b78f4d011185a3debd01a165bcfc875a566dab7777ad6","observation_id":"2638d8e7-5e47-43ad-9163-e84c4210e807","resolution":{"observed_at":"2026-08-06T21:38:51.524460Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T10:39:08.125748Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.03871","last_updated":"2025-09-04T04:12:31Z","snapshot_observed_at":"2026-08-08T03:54:26.940042Z","submitted_at":"2025-09-04T04:12:31Z","title":"A Comprehensive Survey on Trustworthiness in Reasoning with Large Language Models","version":1},"reference_index":242,"source":"pdf_text","source_observed_at":"2026-08-05T10:39:08.125748Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2509.03871"},"observation_digest":"sha256:dc637dbced107c64b01d2ca4eeb804cc23a8d96ccc793c346c9b1b6a571e3d50","observation_id":"32bb5059-d424-47b3-a16e-80c5ac345f4a","resolution":{"observed_at":"2026-08-05T10:39:08.125748Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-04T09:25:58.221798Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.15476","last_updated":"2026-07-04T04:20:15Z","snapshot_observed_at":"2026-08-06T02:45:04.316908Z","submitted_at":"2025-10-17T09:38:54Z","title":"SoK: Systematizing LLM Prompt Security: Taxonomies, Datasets, and Unified Evaluation of Attacks and Defenses","version":3},"reference_index":230,"source":"pdf_text","source_observed_at":"2026-08-04T09:25:58.221798Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2510.15476"},"observation_digest":"sha256:c76ecef05817214c3833c8673096728f41f3d45ab3404af7efb9495cad7f56a7","observation_id":"25b02937-0332-4627-a663-341e2264d9ac","resolution":{"observed_at":"2026-08-04T09:25:58.221798Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2603.13933","last_updated":"2026-04-16T08:02:16Z","snapshot_observed_at":"2026-08-08T16:49:56.947261Z","submitted_at":"2026-03-14T13:04:55Z","title":"OmniCompliance-100K: A Multi-Domain, Rule-Grounded, Real-World Safety Compliance Dataset","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-15T11:32:52.986523Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2603.13933"},"observation_digest":"sha256:18e19f441cae82dc55cac7734fa9aa7806a58ed1c36f408609d64ae972322ccc","observation_id":"fd6579a4-a8d8-4ddf-9357-46a11e763ff8","resolution":{"observed_at":"2026-05-15T11:35:31.636611Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2603.20633","last_updated":"2026-04-17T08:57:17Z","snapshot_observed_at":"2026-07-06T22:49:57.103822Z","submitted_at":"2026-03-21T04:03:45Z","title":"Seed1.8 Model Card: Towards Generalized Real-World Agency","version":3},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-05-15T07:44:02.827006Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2603.20633"},"observation_digest":"sha256:e0adbc36a6c512648f7d35a9074c501c70f724cca2650ebe2aeb7139c4b1c293","observation_id":"93a34e3b-59e6-4fa6-92b9-1d52fe91019f","resolution":{"observed_at":"2026-05-15T07:45:14.326403Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2605.03179","last_updated":"2026-05-04T21:42:10Z","snapshot_observed_at":"2026-08-05T14:29:07.431614Z","submitted_at":"2026-05-04T21:42:10Z","title":"A Validated Prompt Bank for Malicious Code Generation: Separating Executable Weapons from Security Knowledge in 1,554 Consensus-Labeled Prompts","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-08T18:11:29.066362Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2605.03179"},"observation_digest":"sha256:6cec4079cad40bb0329c47d37c23f2261567bf872df1114719c6244e37983308","observation_id":"f3818a60-a2e4-466c-80fa-cf4cc7520b4a","resolution":{"observed_at":"2026-05-09T06:40:43.456003Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2605.06213","last_updated":"2026-05-26T15:14:12Z","snapshot_observed_at":"2026-07-06T23:18:41.400741Z","submitted_at":"2026-05-07T13:15:31Z","title":"Beyond Fixed Benchmarks and Worst-Case Attacks: Dynamic Boundary Evaluation for Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-08T10:13:35.777910Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2605.06213"},"observation_digest":"sha256:99af8055657edbdac976254ecafad7817d0d10aa2cd2df560aa71fc200892213","observation_id":"ad4d6a5d-b8e0-490c-9650-c701483c320c","resolution":{"observed_at":"2026-05-11T20:06:13.738353Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2605.14152","last_updated":"2026-07-07T01:59:55Z","snapshot_observed_at":"2026-08-09T17:15:27.929117Z","submitted_at":"2026-05-13T22:07:22Z","title":"ROK-FORTRESS: Measuring the Effect of Geopolitical Transcreation for National Security and Public Safety","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-15T04:55:36.069767Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2605.14152"},"observation_digest":"sha256:d7ef092ebca8dc846dbec8b1b31a631bdcdd5e7425d5cf8add7a3a2275f216b1","observation_id":"7ebd75bf-b23f-4b84-97de-3bd1b6993d3f","resolution":{"observed_at":"2026-05-15T04:59:45.799831Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-07-12T16:49:32.395383Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.14152","last_updated":"2026-07-07T01:59:55Z","snapshot_observed_at":"2026-08-09T17:15:27.929117Z","submitted_at":"2026-05-13T22:07:22Z","title":"ROK-FORTRESS: Measuring the Effect of Geopolitical Transcreation for National Security and Public Safety","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-12T16:49:32.395383Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2605.14152"},"observation_digest":"sha256:e7829cabb6c745bd50589a3019e27bfc12f1ed19a51105a94845175d29b41f8a","observation_id":"a768578b-b939-42c5-a201-531f761575b0","resolution":{"observed_at":"2026-07-12T16:49:32.395383Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2605.22643","last_updated":"2026-05-22T14:53:30Z","snapshot_observed_at":"2026-07-06T23:32:59.663926Z","submitted_at":"2026-05-21T15:50:18Z","title":"Boiling the Frog: A Multi-Turn Benchmark for Agentic Safety","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-05-22T05:50:28.114140Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2605.22643"},"observation_digest":"sha256:988f4335e938930e656add7d88fde9955138c9e2cdee16f328ebf8f9d250a864","observation_id":"361f769d-7269-4c3e-8ecf-29bfd8626531","resolution":{"observed_at":"2026-05-22T05:51:07.882817Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2605.22643","last_updated":"2026-05-22T14:53:30Z","snapshot_observed_at":"2026-07-06T23:32:59.663926Z","submitted_at":"2026-05-21T15:50:18Z","title":"Boiling the Frog: A Multi-Turn Benchmark for Agentic Safety","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-05-25T06:05:27.736494Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2605.22643"},"observation_digest":"sha256:9f2186598798eeb3a2dd533afe7776bb1bab70a1053fd3312888ea6805be7cf4","observation_id":"5671be64-514b-49dc-a6b5-e5b8aac7edb1","resolution":{"observed_at":"2026-05-25T06:06:42.863331Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2606.01481","last_updated":"2026-05-31T22:46:35Z","snapshot_observed_at":"2026-08-06T16:16:18.407347Z","submitted_at":"2026-05-31T22:46:35Z","title":"SafeGen-Bench: Benchmarking Safety in Image-Conditioned Text-to-Video Generation","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-06-28T17:05:57.685728Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2606.01481"},"observation_digest":"sha256:0d8b5a08ca9f68a75dde13b8def2ae33e4ddec9e4a48f0b113015a3183c5ca47","observation_id":"4b7a78d5-fc87-4552-9dfc-399aa26317b1","resolution":{"observed_at":"2026-06-28T17:12:25.159443Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2606.03648","last_updated":"2026-06-02T13:39:17Z","snapshot_observed_at":"2026-08-09T14:15:59.693564Z","submitted_at":"2026-06-02T13:39:17Z","title":"Safety Measurements for Fine-tuned LLMs Should be Grounded in Capability","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-06-28T10:00:30.904247Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2606.03648"},"observation_digest":"sha256:efd3a7e5c8696cef7eb0ba835eb1f076c45abacc87532c2cf4b2fbb4c8c9ee2c","observation_id":"d8ecb8ee-7117-4d70-ba5a-7ac8bfe58c4c","resolution":{"observed_at":"2026-07-02T03:26:29.809439Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2606.04394","last_updated":"2026-06-03T03:17:26Z","snapshot_observed_at":"2026-08-07T11:58:41.205053Z","submitted_at":"2026-06-03T03:17:26Z","title":"Beyond Single-Policy: Evaluating Composed Organization-Specific Policy Alignment in LLM Chatbots","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-06-28T05:41:16.033862Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2606.04394"},"observation_digest":"sha256:492ffd6eb715cca102e8439ac040128570fd7eec4dbe2657e95864aade8d15b7","observation_id":"a49d69b0-35de-48a1-93c7-5d5d442eaa76","resolution":{"observed_at":"2026-06-28T05:41:40.158587Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2606.08376","last_updated":"2026-06-07T00:06:27Z","snapshot_observed_at":"2026-08-03T08:52:14.234243Z","submitted_at":"2026-06-07T00:06:27Z","title":"RiskNet: A large-scale dataset of AI risk incidents from news with alignment and multi-dimensional annotations","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-27T18:37:24.395892Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2606.08376"},"observation_digest":"sha256:900e22ff05b15cc91aa5cd9544a15f298ae8db72490c41b0eb019cd3102e9b8b","observation_id":"821c38c6-b39a-4790-b855-e6e2e115e6e1","resolution":{"observed_at":"2026-07-02T22:47:26.405643Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2606.09178","last_updated":"2026-06-22T07:46:06Z","snapshot_observed_at":"2026-07-06T23:48:35.190418Z","submitted_at":"2026-06-08T08:17:18Z","title":"Culturally-Adapted Red-Teaming Across East and Southeast Asian Contexts: A Methodological and Comparative Analysis","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-27T16:48:54.802860Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2606.09178"},"observation_digest":"sha256:ddc009aa739b810b03a94205c1bb8943021a31bd6922cc88ce7c6fec4be85aed","observation_id":"cd4ad914-2d01-4bd9-8315-73e47c8324af","resolution":{"observed_at":"2026-07-03T01:07:30.328565Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2606.19887","last_updated":"2026-06-24T07:42:32Z","snapshot_observed_at":"2026-07-06T23:55:05.932014Z","submitted_at":"2026-06-18T07:46:18Z","title":"FinRED: An Expert-Guided Benchmark Generation and Evaluation Framework for Financial LLM Red-Teaming","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-26T17:11:40.088809Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2606.19887"},"observation_digest":"sha256:3fe56946339d8c66b20f0157786d5c80e2f1314b0453f6c3a2df93acdf8954cd","observation_id":"b466eda9-ecf2-4ae4-9d68-1757f587d565","resolution":{"observed_at":"2026-07-04T04:09:34.777211Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":"2407.17436","doi":"10.48550/arxiv.2407.17436","metadata_source":"arxiv_reference","pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A Data Source Details In this section, we will show the detailed infor- mation for regulations and policies collected in OmniCompliance-100K","venue":"arXiv (Cornell University)","work_id":"61c735f0-1aff-45fa-90bc-579e23106404","year":2024},"citing_paper":{"arxiv_id":"2606.20626","last_updated":"2026-05-26T17:35:31Z","snapshot_observed_at":"2026-08-06T08:20:13.528819Z","submitted_at":"2026-05-26T17:35:31Z","title":"Efficient Safety Benchmarking via Item Response Theory","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-01T15:51:00.484343Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2606.20626"},"observation_digest":"sha256:aa1d849a0876f4d2bf7152e2c39089cf5b11f8641d451003b4691540f6381bf1","observation_id":"3eccc86e-7846-4ba5-a478-6d7a2ec44d1e","resolution":{"observed_at":"2026-07-01T15:55:48.974164Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-01T18:54:52.928929Z","title":"arXiv preprint arXiv:2407.17436 , year=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.17152","last_updated":"2026-07-19T09:18:35Z","snapshot_observed_at":"2026-08-09T14:18:14.468847Z","submitted_at":"2026-07-19T09:18:35Z","title":"How Jailbreak Attacks Inform Safety Alignment: A Defender-Centric, Shapley-Based Evaluation of Jailbreak Contributions","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-01T18:54:52.928929Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2607.17152"},"observation_digest":"sha256:7588ce4aaf6710aabdd4184caef2076055d4ac83fa01e4d14bee04c6267620e2","observation_id":"7c3bdb28-b686-421f-ac4d-5bb66a412136","resolution":{"observed_at":"2026-08-01T18:54:52.928929Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-02T08:30:01.720511Z","title":"Air-bench 2024: A safety benchmark based on risk categories from regulations and policies.https://arxiv.org/abs/2407.17436, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.22671","last_updated":"2026-07-06T17:26:46Z","snapshot_observed_at":"2026-08-06T17:18:06.876668Z","submitted_at":"2026-07-06T17:26:46Z","title":"AIR-BENCH Live: An Evolving Safety Benchmark for Foundation Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-02T08:30:01.720511Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2607.22671"},"observation_digest":"sha256:9566d91e749488b2854c0667beb833bc716f828f56c1e6a17cbe693009b1f942","observation_id":"76cd7933-fe44-421e-b840-c1cef0f98e17","resolution":{"observed_at":"2026-08-02T08:30:01.720511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17436","snapshot_observed_at":"2026-08-03T00:55:22.163420Z","title":"arXiv preprint arXiv:2407.17436 , year=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28636","last_updated":"2026-05-19T13:56:13Z","snapshot_observed_at":"2026-08-06T00:38:04.006327Z","submitted_at":"2026-05-19T13:56:13Z","title":"Chain-of-Models: Cross-Model Auditing for Bias-Robust LLM Judges","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-03T00:55:22.163420Z"},"links":{"cited_paper":"/paper/2407.17436","citing_paper":"/paper/2607.28636"},"observation_digest":"sha256:adac8345b326f78d2bfdcb64772f350886c63be1da2d91e1c1199e8790a24ab7","observation_id":"2a2bfb26-0fbe-4c53-955a-2c7da7224d28","resolution":{"observed_at":"2026-08-03T00:55:22.163420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2407.17436/citation-record","integrity":"/paper/2407.17436/integrity","json":"/paper/2407.17436/citation-record.json","paper":"/paper/2407.17436"},"outbound":[],"paper":{"arxiv_id":"2407.17436","last_updated":"2024-08-05T18:12:27Z","latest_version":2,"primary_category":"cs.CY","snapshot_observed_at":"2026-08-08T01:17:55.124326Z","submitted_at":"2024-07-11T21:16:48Z","title":"AIR-Bench 2024: A Safety Benchmark Based on Risk Categories from Regulations and Policies"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 28 inbound Pith citation observations for arXiv:2407.17436."}