{"as_of":"2026-08-10T16:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d869c3a16998afd8b31f7c9799e397b34e7fedf18924199406e137cba66d4594","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":12,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":12,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":12,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":12,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:25:22.449090Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-08T19:55:33.960847Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-08-07T14:25:22.449090Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19073","last_updated":"2025-07-20T15:35:43Z","snapshot_observed_at":"2026-08-10T00:41:11.948345Z","submitted_at":"2025-05-25T10:17:57Z","title":"Towards Harmonized Uncertainty Estimation for Large Language Models","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T14:25:22.449090Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2505.19073"},"observation_digest":"sha256:f3bbd830faa4dffe94c70a3ac546a64f7f3282f07fb88a3b40c9f75fe8c50879","observation_id":"1555fe95-6146-4ff3-bf68-eb5919e48133","resolution":{"observed_at":"2026-08-07T14:25:22.449090Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-08-07T12:08:37.100267Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.00519","last_updated":"2025-06-03T09:58:17Z","snapshot_observed_at":"2026-08-08T14:54:48.673691Z","submitted_at":"2025-05-31T11:35:31Z","title":"CausalAbstain: Enhancing Multilingual LLMs with Causal Reasoning for Trustworthy Abstention","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T12:08:37.100267Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2506.00519"},"observation_digest":"sha256:960b16f001739deae68e672f14b892d580ad879357dc343a724f76a608f90cc0","observation_id":"ddee5a53-226a-47fc-9163-449121d25e26","resolution":{"observed_at":"2026-08-07T12:08:37.100267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-08-06T22:49:23.531540Z","title":"and Hashemi, M., 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.20598","last_updated":"2025-06-25T16:37:46Z","snapshot_observed_at":"2026-08-09T07:52:39.034038Z","submitted_at":"2025-06-25T16:37:46Z","title":"Fine-Tuning and Prompt Engineering of LLMs, for the Creation of Multi-Agent AI for Addressing Sustainable Protein Production Challenges","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T22:49:23.531540Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2506.20598"},"observation_digest":"sha256:d31b4f78829daae904f0ba738a131a693167143ae7505b529085f211aab37361","observation_id":"4f7557c1-9c72-41b1-9710-27353326a2ed","resolution":{"observed_at":"2026-08-06T22:49:23.531540Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-08-06T14:57:08.233389Z","title":"Do llms know when to not answer? investigating abstention abilities of large lan- guage models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.17262","last_updated":"2025-07-23T07:00:19Z","snapshot_observed_at":"2026-08-08T12:09:57.704230Z","submitted_at":"2025-07-23T07:00:19Z","title":"VisionTrap: Unanswerable Questions On Visual Data","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T14:57:08.233389Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2507.17262"},"observation_digest":"sha256:d693cac3d5c90bdb282488605aae2c5f643eadfb4062d0044c3a8566ba25a2f7","observation_id":"5aa7c029-4be7-4f30-81db-8b7951dcb419","resolution":{"observed_at":"2026-08-06T14:57:08.233389Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-08-05T23:23:47.223084Z","title":"Do llms know when to not answer? investigating abstention abilities of large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.10022","last_updated":"2025-08-07T16:46:47Z","snapshot_observed_at":"2026-08-07T17:33:40.545319Z","submitted_at":"2025-08-07T16:46:47Z","title":"Conformal P-Value in Multiple-Choice Question Answering Tasks with Provable Risk Control","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-05T23:23:47.223084Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2508.10022"},"observation_digest":"sha256:86ee2faffd1eb9f9dc6075aab2a12c6a147321271c3dbf9483ed45e77781190e","observation_id":"a58d7397-4577-4cc5-9ca4-b5aab31dee69","resolution":{"observed_at":"2026-08-05T23:23:47.223084Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-08-04T20:09:16.900567Z","title":"https://arxiv.org/abs/2407.16221","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.08803","last_updated":"2025-09-10T17:36:25Z","snapshot_observed_at":"2026-08-06T13:54:36.806572Z","submitted_at":"2025-09-10T17:36:25Z","title":"Scaling Truth: The Confidence Paradox in AI Fact-Checking","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-04T20:09:16.900567Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2509.08803"},"observation_digest":"sha256:02a17c748e629d644a62846efbeab7283a870b80d3620070e50339cd1e573922","observation_id":"2418c629-0792-4380-a047-235bada9e429","resolution":{"observed_at":"2026-08-04T20:09:16.900567Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":"2407.16221","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-07-08T19:55:33.960847Z","title":"Do llms know when to not answer? investigating abstention abilities of large language models.arXiv preprint arXiv:2407.16221","venue":"cs.CL","work_id":"fbe2f53d-5ccf-4007-8582-5d2688add694","year":2024},"citing_paper":{"arxiv_id":"2603.22161","last_updated":"2026-05-19T14:46:51Z","snapshot_observed_at":"2026-08-06T02:25:28.537505Z","submitted_at":"2026-03-23T16:23:31Z","title":"Causal Evidence that Language Models use Confidence to Drive Behavior","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-21T09:43:05.524088Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2603.22161"},"observation_digest":"sha256:d2c7c74fabc525896886daf7126660057077a9b144815756eb7aa0a521c284e9","observation_id":"f091fd4c-f072-403c-a475-1c593a7ef720","resolution":{"observed_at":"2026-05-21T09:44:05.697748Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":"2407.16221","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-07-08T19:55:33.960847Z","title":"Do llms know when to not answer? investigating abstention abilities of large language models.arXiv preprint arXiv:2407.16221","venue":"cs.CL","work_id":"fbe2f53d-5ccf-4007-8582-5d2688add694","year":2024},"citing_paper":{"arxiv_id":"2604.03216","last_updated":"2026-04-03T17:44:32Z","snapshot_observed_at":"2026-07-06T22:52:25.307426Z","submitted_at":"2026-04-03T17:44:32Z","title":"BAS: A Decision-Theoretic Approach to Evaluating Large Language Model Confidence","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-13T19:48:13.133733Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2604.03216"},"observation_digest":"sha256:1b4de97435402eb2f39bcde5db2f15357e9fc7132a63c940ca3547091c13e180","observation_id":"a5f1f369-6d9a-4308-a79e-158928104152","resolution":{"observed_at":"2026-05-13T19:53:11.971607Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":"2407.16221","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-07-08T19:55:33.960847Z","title":"Do llms know when to not answer? investigating abstention abilities of large language models.arXiv preprint arXiv:2407.16221","venue":"cs.CL","work_id":"fbe2f53d-5ccf-4007-8582-5d2688add694","year":2024},"citing_paper":{"arxiv_id":"2604.16752","last_updated":"2026-04-17T23:54:34Z","snapshot_observed_at":"2026-07-06T23:04:01.558812Z","submitted_at":"2026-04-17T23:54:34Z","title":"Don't Start What You Can't Finish: A Counterfactual Audit of Support-State Triage in LLM Agents","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T07:55:07.930578Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2604.16752"},"observation_digest":"sha256:ea6b2df8169ef7d9e99914695015fb2bf50eef87c55c393dd7764706f1799f42","observation_id":"c22f51ba-b88e-46f3-8a36-4a272bcd8648","resolution":{"observed_at":"2026-05-10T07:57:15.529664Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":"2407.16221","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-07-08T19:55:33.960847Z","title":"Do llms know when to not answer? investigating abstention abilities of large language models.arXiv preprint arXiv:2407.16221","venue":"cs.CL","work_id":"fbe2f53d-5ccf-4007-8582-5d2688add694","year":2024},"citing_paper":{"arxiv_id":"2605.05379","last_updated":"2026-05-06T19:01:29Z","snapshot_observed_at":"2026-07-31T06:35:39.314747Z","submitted_at":"2026-05-06T19:01:29Z","title":"Partial Evidence Bench: Benchmarking Authorization-Limited Evidence in Agentic Systems","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-08T17:28:51.943383Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2605.05379"},"observation_digest":"sha256:f55b6b730b9e93e3369976c39b1b036decfda87f1c233145f842f788ab381f28","observation_id":"2ef30ac9-5799-416e-bae7-7a688aaa7a63","resolution":{"observed_at":"2026-05-11T17:31:07.879312Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":"2407.16221","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-07-08T19:55:33.960847Z","title":"Do llms know when to not answer? investigating abstention abilities of large language models.arXiv preprint arXiv:2407.16221","venue":"cs.CL","work_id":"fbe2f53d-5ccf-4007-8582-5d2688add694","year":2024},"citing_paper":{"arxiv_id":"2606.23937","last_updated":"2026-06-22T20:57:11Z","snapshot_observed_at":"2026-08-03T07:55:40.358007Z","submitted_at":"2026-06-22T20:57:11Z","title":"When Retrieval Metrics Mislead: Measuring Policy Signal in Long-Horizon Tool-Use Agents","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-06-26T08:04:57.334657Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2606.23937"},"observation_digest":"sha256:65609464d067efa4c5d433476d42c35ecd13d3c982fa5278a958931497c1aa53","observation_id":"f21201c5-3228-4cca-be59-44147ff29bac","resolution":{"observed_at":"2026-07-04T11:19:49.568196Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models","version":2},"cited_work":{"arxiv_id":"2407.16221","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.16221","snapshot_observed_at":"2026-07-08T19:55:33.960847Z","title":"Do llms know when to not answer? investigating abstention abilities of large language models.arXiv preprint arXiv:2407.16221","venue":"cs.CL","work_id":"fbe2f53d-5ccf-4007-8582-5d2688add694","year":2024},"citing_paper":{"arxiv_id":"2607.05985","last_updated":"2026-07-07T08:17:49Z","snapshot_observed_at":"2026-08-09T01:50:13.589992Z","submitted_at":"2026-07-07T08:17:49Z","title":"Auto-DSM Under the Lens: A Black-Box Evaluation Framework for LLM-Based DSM Generation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-08T19:45:46.002113Z"},"links":{"cited_paper":"/paper/2407.16221","citing_paper":"/paper/2607.05985"},"observation_digest":"sha256:f18f892242cf52322a80fc0b9ce9cf4b6355c13a140283e310981fdf741cf35c","observation_id":"7cbb30f9-6e33-4599-8211-fc5229f886a2","resolution":{"observed_at":"2026-07-08T19:55:33.962308Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2407.16221/citation-record","integrity":"/paper/2407.16221/integrity","json":"/paper/2407.16221/citation-record.json","paper":"/paper/2407.16221"},"outbound":[],"paper":{"arxiv_id":"2407.16221","last_updated":"2024-09-24T14:25:58Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-08T06:05:05.745431Z","submitted_at":"2024-07-23T06:56:54Z","title":"Do LLMs Know When to NOT Answer? Investigating Abstention Abilities of Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 12 inbound Pith citation observations for arXiv:2407.16221."}