{"as_of":"2026-08-22T00:42:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0e163188220125d0bd29779391a375d92f429d94228bd61b1a2b18cba1a47afb","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":6,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":6,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T18:53:15.648678Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T09:29:44.263692Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.11889","last_updated":"2024-06-17T02:48:21Z","snapshot_observed_at":"2026-08-18T10:59:10.728357Z","submitted_at":"2024-02-19T06:58:42Z","title":"ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11889","snapshot_observed_at":"2026-08-11T18:53:15.648678Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.10423","last_updated":"2025-04-14T12:52:24Z","snapshot_observed_at":"2026-08-18T14:27:57.499671Z","submitted_at":"2024-12-10T12:42:33Z","title":"Look Before You Leap: Enhancing Attention and Vigilance Regarding Harmful Content with GuidelineLLM","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-11T18:53:15.648678Z"},"links":{"cited_paper":"/paper/2402.11889","citing_paper":"/paper/2412.10423"},"observation_digest":"sha256:c25323ce8798844ef2f0d23d2af987ace6e9a8460ad2d5621d9cd94ab0a7e03b","observation_id":"3e8d1165-3c6e-4798-9d66-753133cc7da7","resolution":{"observed_at":"2026-08-11T18:53:15.648678Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11889","last_updated":"2024-06-17T02:48:21Z","snapshot_observed_at":"2026-08-18T10:59:10.728357Z","submitted_at":"2024-02-19T06:58:42Z","title":"ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11889","snapshot_observed_at":"2026-08-09T16:18:40.699323Z","title":"Rose doesn’t do that: Boosting the safety of instruction-tuned large language models with reverse prompt contrastive decoding.arXiv preprint arXiv:2402.11889, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01208","last_updated":"2025-06-20T10:54:05Z","snapshot_observed_at":"2026-08-16T10:10:26.213081Z","submitted_at":"2025-02-03T09:59:32Z","title":"On Almost Surely Safe Alignment of Large Language Models at Inference-Time","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-09T16:18:40.699323Z"},"links":{"cited_paper":"/paper/2402.11889","citing_paper":"/paper/2502.01208"},"observation_digest":"sha256:e5eceabc84a8c278a6727fe8cbe318fa32b7656d797b6f5671a23c315d874838","observation_id":"123dee81-6968-497b-8966-f3b1c026d33a","resolution":{"observed_at":"2026-08-09T16:18:40.699323Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11889","last_updated":"2024-06-17T02:48:21Z","snapshot_observed_at":"2026-08-18T10:59:10.728357Z","submitted_at":"2024-02-19T06:58:42Z","title":"ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","version":2},"cited_work":{"arxiv_id":"2402.11889","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.11889","snapshot_observed_at":"2026-07-04T09:29:44.263692Z","title":"ROSE Doesn’t Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","venue":null,"work_id":"5124c447-663b-47bf-99a9-cae99f50124b","year":2024},"citing_paper":{"arxiv_id":"2502.17419","last_updated":"2025-06-25T02:24:46Z","snapshot_observed_at":"2026-08-20T08:13:05.630967Z","submitted_at":"2025-02-24T18:50:52Z","title":"From System 1 to System 2: A Survey of Reasoning Large Language Models","version":6},"reference_index":214,"source":"pdf_text","source_observed_at":"2026-05-13T01:36:23.845366Z"},"links":{"cited_paper":"/paper/2402.11889","citing_paper":"/paper/2502.17419"},"observation_digest":"sha256:b2561db47b96c6f4b0634fd8447e0d6f9114fe93956b864350f9a4b967b530e1","observation_id":"1253af54-4495-4576-919c-9a0b5f21ff5b","resolution":{"observed_at":"2026-05-13T01:36:24.324707Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11889","last_updated":"2024-06-17T02:48:21Z","snapshot_observed_at":"2026-08-18T10:59:10.728357Z","submitted_at":"2024-02-19T06:58:42Z","title":"ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11889","snapshot_observed_at":"2026-08-05T13:11:15.385658Z","title":"Zhong, L","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.00869","last_updated":"2025-08-31T14:29:54Z","snapshot_observed_at":"2026-08-20T12:51:03.922770Z","submitted_at":"2025-08-31T14:29:54Z","title":"Exploring and Mitigating Fawning Hallucinations in Large Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T13:11:15.385658Z"},"links":{"cited_paper":"/paper/2402.11889","citing_paper":"/paper/2509.00869"},"observation_digest":"sha256:9c4d5099c7fc3cc3941b7a58e8f6bb999e0eae9e1d153f6394ba026abda91e63","observation_id":"2ed13e14-0892-4b18-a2a6-46c16f8511df","resolution":{"observed_at":"2026-08-05T13:11:15.385658Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11889","last_updated":"2024-06-17T02:48:21Z","snapshot_observed_at":"2026-08-18T10:59:10.728357Z","submitted_at":"2024-02-19T06:58:42Z","title":"ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","version":2},"cited_work":{"arxiv_id":"2402.11889","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.11889","snapshot_observed_at":"2026-07-04T09:29:44.263692Z","title":"ROSE Doesn’t Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","venue":null,"work_id":"5124c447-663b-47bf-99a9-cae99f50124b","year":2024},"citing_paper":{"arxiv_id":"2606.22686","last_updated":"2026-06-30T05:48:55Z","snapshot_observed_at":"2026-08-16T10:51:03.316399Z","submitted_at":"2026-06-21T22:04:48Z","title":"The Geometry of Refusal: Linear Instability in Safety-Aligned LLMs","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-06-26T09:52:52.603708Z"},"links":{"cited_paper":"/paper/2402.11889","citing_paper":"/paper/2606.22686"},"observation_digest":"sha256:823cf61ad9ebf81d40b397305441d452a14087568f38262e727c21179455b76b","observation_id":"03628607-c82e-43f1-b156-b6511d09dde8","resolution":{"observed_at":"2026-07-04T09:29:44.265560Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11889","last_updated":"2024-06-17T02:48:21Z","snapshot_observed_at":"2026-08-18T10:59:10.728357Z","submitted_at":"2024-02-19T06:58:42Z","title":"ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","version":2},"cited_work":{"arxiv_id":"2402.11889","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.11889","snapshot_observed_at":"2026-07-04T09:29:44.263692Z","title":"ROSE Doesn’t Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding","venue":null,"work_id":"5124c447-663b-47bf-99a9-cae99f50124b","year":2024},"citing_paper":{"arxiv_id":"2606.22686","last_updated":"2026-06-30T05:48:55Z","snapshot_observed_at":"2026-08-16T10:51:03.316399Z","submitted_at":"2026-06-21T22:04:48Z","title":"The Geometry of Refusal: Linear Instability in Safety-Aligned LLMs","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-07-01T07:07:45.177242Z"},"links":{"cited_paper":"/paper/2402.11889","citing_paper":"/paper/2606.22686"},"observation_digest":"sha256:35b6815187758bfdfc2e82d05d4755c1309fe576ba26b5524f90e11b7ca7f64c","observation_id":"7d9809cf-d8ef-40c9-a3a9-8b10e5f846e8","resolution":{"observed_at":"2026-07-01T08:55:35.054465Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2402.11889/citation-record","integrity":"/paper/2402.11889/integrity","json":"/paper/2402.11889/citation-record.json","paper":"/paper/2402.11889"},"outbound":[],"paper":{"arxiv_id":"2402.11889","last_updated":"2024-06-17T02:48:21Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-18T10:59:10.728357Z","submitted_at":"2024-02-19T06:58:42Z","title":"ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 6 inbound Pith citation observations for arXiv:2402.11889."}