{"as_of":"2026-08-10T15:34:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9350be6f2018f743933c6dbccc52c6d58cc3de4331090c43b240f18426bdc073","coverage":[{"denominator":99,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":99,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T12:42:39.941106Z","state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T04:48:25.173042Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.02023","snapshot_observed_at":"2026-08-01T04:48:25.173042Z","title":"arXiv:2601.02023","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.22448","last_updated":"2026-07-27T23:08:28Z","snapshot_observed_at":"2026-08-08T03:02:00.797726Z","submitted_at":"2026-07-24T16:10:21Z","title":"Where Facts Go Missing: A Layerwise Taxonomy and Per-Layer Attribution of Information Omission in Air-Gapped LLMAgent Pipelines","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-01T04:48:25.173042Z"},"links":{"cited_paper":"/paper/2601.02023","citing_paper":"/paper/2607.22448"},"observation_digest":"sha256:03fd9b568314c2ab64dfbbb38d412e15fba306f1dc768a3f22125dda9cf9b312","observation_id":"15fffd33-73d3-4656-8fc1-4708bf4a229f","resolution":{"observed_at":"2026-08-01T04:48:25.173042Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2601.02023/citation-record","integrity":"/paper/2601.02023/integrity","json":"/paper/2601.02023/citation-record.json","paper":"/paper/2601.02023"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.10149","last_updated":"2024-11-06T14:50:40Z","snapshot_observed_at":"2026-08-09T01:19:53.754048Z","submitted_at":"2024-06-14T16:00:29Z","title":"BABILong: Testing the Limits of LLMs with Long Context Reasoning-in-a-Haystack","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10149","snapshot_observed_at":"2026-08-03T12:42:27.870418Z","title":"BABILong: Testing the limits of LLMs with long context reasoning-in-a-haystack,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:27.870418Z"},"links":{"cited_paper":"/paper/2406.10149","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:4b2516c0459a72b5d9fe025d388494676506428bc6ebde0071804a0a00c92376","observation_id":"e56728d6-6922-42a7-8896-6aef5a8a41dd","resolution":{"observed_at":"2026-08-03T12:42:27.870418Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:27.933008Z","title":"Needlebench: Can llms do retrieval and reasoning in 1 million context window?","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:27.933008Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:8318afeefee2ce4a3ff21387f350242eaf2402b195b3c051f7bf2ecb8e5a7020","observation_id":"9392e76b-b182-4655-baa6-d111f6678e85","resolution":{"observed_at":"2026-08-03T12:42:27.933008Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.13718","last_updated":"2024-02-24T15:07:55Z","snapshot_observed_at":"2026-08-02T23:59:10.592798Z","submitted_at":"2024-02-21T11:30:29Z","title":"$\\infty$Bench: Extending Long Context Evaluation Beyond 100K Tokens","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.13718","snapshot_observed_at":"2026-08-03T12:42:28.024235Z","title":"∞bench: Extending long context evaluation beyond 100k tokens,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.024235Z"},"links":{"cited_paper":"/paper/2402.13718","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:7d22d08e7e2b252b7f2c87bc5b48c3e89ad00a5cf1a7e7f4bdfdea7a6f6e8558","observation_id":"1b681f5a-3dfe-4d30-964b-97071f816c23","resolution":{"observed_at":"2026-08-03T12:42:28.024235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14488","last_updated":"2024-02-22T12:26:07Z","snapshot_observed_at":"2026-08-05T01:58:48.384538Z","submitted_at":"2024-02-22T12:26:07Z","title":"Does the Generator Mind its Contexts? An Analysis of Generative Model Faithfulness under Context Transfer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14488","snapshot_observed_at":"2026-08-03T12:42:28.201572Z","title":"Rethinking context length in large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.201572Z"},"links":{"cited_paper":"/paper/2402.14488","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:0f1376292648583cc5f167971ac9b0a20647b8e313e6dde1cf7ee520e2686630","observation_id":"5c32c017-281b-41cc-b11e-c9fc52f77fb2","resolution":{"observed_at":"2026-08-03T12:42:28.201572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:28.258057Z","title":"Lost in the middle: How language models use long contexts,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.258057Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:e601c14be98bd2020b8438f5f8126e19968a694983dfb6fdf88e1e0835c29d55","observation_id":"b93e4054-dbeb-44ae-a758-e39eaad75f5b","resolution":{"observed_at":"2026-08-03T12:42:28.258057Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.06654","last_updated":"2024-08-06T21:48:58Z","snapshot_observed_at":"2026-08-10T09:20:57.804343Z","submitted_at":"2024-04-09T23:41:27Z","title":"RULER: What's the Real Context Size of Your Long-Context Language Models?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.06654","snapshot_observed_at":"2026-08-03T12:42:28.451858Z","title":"RULER: What’s the real context size of your long-context language models?","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.451858Z"},"links":{"cited_paper":"/paper/2404.06654","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:e530a13bf0d65450ff69217ec4e16d1dab8a5a8004dc769a80d1fecca1c18e2b","observation_id":"db891dfb-12f3-469b-aa78-9454e26c06c4","resolution":{"observed_at":"2026-08-03T12:42:28.451858Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:28.499347Z","title":"Lv-eval: A balanced long-context benchmark with 5 length levels up to 256k,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.499347Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:c5d4df3253ac7be122f18f857da62b37e6ca07b17adf9032ea2d81388863f047","observation_id":"89016f02-1806-41d2-a19a-6ee07c167311","resolution":{"observed_at":"2026-08-03T12:42:28.499347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05424","last_updated":"2025-07-07T19:13:20Z","snapshot_observed_at":"2026-08-09T07:23:08.553839Z","submitted_at":"2025-07-07T19:13:20Z","title":"\"Lost-in-the-Later\": Framework for Quantifying Contextual Grounding in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05424","snapshot_observed_at":"2026-08-03T12:42:28.632427Z","title":"Lost-in-the-later: Framework for quantifying contextual grounding in large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.632427Z"},"links":{"cited_paper":"/paper/2507.05424","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:2945c932ce8aad24a88f7cf67b1a3056f4206180fc5801825e62327972d96cf9","observation_id":"d7e61b3e-d44f-482a-9b4a-8b9a998b34a6","resolution":{"observed_at":"2026-08-03T12:42:28.632427Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12641","last_updated":"2024-11-11T05:48:35Z","snapshot_observed_at":"2026-08-10T14:57:13.162234Z","submitted_at":"2024-06-18T14:08:01Z","title":"DetectBench: Can Large Language Model Detect and Piece Together Implicit Evidence?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12641","snapshot_observed_at":"2026-08-03T12:42:28.777774Z","title":"Detectbench: Can large language model detect and piece together implicit evidence?","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.777774Z"},"links":{"cited_paper":"/paper/2406.12641","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:78163d0b20bca77a2d19e4d4e675201d53ec4b01554295518bab24320db2e439","observation_id":"18c824d7-dab0-431b-aa3a-643d2e5ede78","resolution":{"observed_at":"2026-08-03T12:42:28.777774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18006","last_updated":"2024-10-12T20:11:52Z","snapshot_observed_at":"2026-08-09T07:55:17.057124Z","submitted_at":"2024-09-26T16:15:14Z","title":"Evaluating Multilingual Long-Context Models for Retrieval and Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18006","snapshot_observed_at":"2026-08-03T12:42:28.914371Z","title":"Evaluating multilingual long- context models for retrieval and reasoning,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:28.914371Z"},"links":{"cited_paper":"/paper/2409.18006","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:f1f2ecc84a1700654acedaa01858622be35021c8d9b835a884d63f56b4701fb3","observation_id":"7e8587c5-622e-48d5-bef2-2323dcd70883","resolution":{"observed_at":"2026-08-03T12:42:28.914371Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:29.105564Z","title":"The two-hop curse: LLMs trained on 𝐴→𝐵,𝐵→𝐶 fail to learn𝐴→𝐶,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:29.105564Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:15f6f2f5ca4709183d13baf7fa58d206eeb7e79fb5ee54041173d24b8112b905","observation_id":"9677f9e4-6276-43cb-9c9b-a068d93fcb63","resolution":{"observed_at":"2026-08-03T12:42:29.105564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.16679","last_updated":"2025-05-31T11:16:44Z","snapshot_observed_at":"2026-08-09T07:56:12.121579Z","submitted_at":"2024-11-25T18:59:30Z","title":"Do Large Language Models Perform Latent Multi-Hop Reasoning without Exploiting Shortcuts?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.16679","snapshot_observed_at":"2026-08-03T12:42:29.274642Z","title":"Do large language models perform latent multi-hop reasoning without exploiting shortcuts?","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:29.274642Z"},"links":{"cited_paper":"/paper/2411.16679","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:bac675ad285a4d04a7a17ae9283252b1e3b1c2a7d1602e67bf9b5c810111aa6d","observation_id":"317a9de4-fe3b-4838-80bb-0fb85b5bd291","resolution":{"observed_at":"2026-08-03T12:42:29.274642Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:29.526186Z","title":"Generating wikipedia by summarizing long sequences,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:29.526186Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:63b241386aadf35439cbdba0112e3f107784a21ad5cdcdf158b20e0fc45a359e","observation_id":"efb1b6a2-9bda-4029-9076-d882199f5887","resolution":{"observed_at":"2026-08-03T12:42:29.526186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.22257","last_updated":"2025-01-08T02:33:39Z","snapshot_observed_at":"2026-08-05T18:34:03.849161Z","submitted_at":"2024-10-29T17:19:56Z","title":"FactBench: A Dynamic Benchmark for In-the-Wild Language Model Factuality Evaluation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.22257","snapshot_observed_at":"2026-08-03T12:42:29.787906Z","title":"FactBench: A dynamic benchmark for in-the-wild language model factuality evaluation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:29.787906Z"},"links":{"cited_paper":"/paper/2410.22257","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:a511ecf307da174cb05e06c18dc48b9c310f32ee05e14e197036b4e110679bbd","observation_id":"3c48e851-65d2-4ad9-97cf-8c65f758a34f","resolution":{"observed_at":"2026-08-03T12:42:29.787906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.05133","last_updated":"2024-12-09T04:21:08Z","snapshot_observed_at":"2026-08-06T13:42:00.057517Z","submitted_at":"2024-02-06T04:18:58Z","title":"Personalized Language Modeling from Personalized Human Feedback","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.05133","snapshot_observed_at":"2026-08-03T12:42:29.904718Z","title":"Lv-eval: A balanced long-context benchmark with 5 length levels up to 256k,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:29.904718Z"},"links":{"cited_paper":"/paper/2402.05133","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:41c7ea57b54b47ea3d2ec2e8e71e889d2b68f4d9d1db0802fa0fc846c69542ea","observation_id":"1fee4a06-e7f3-4d75-a4b1-4865469298c8","resolution":{"observed_at":"2026-08-03T12:42:29.904718Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03727","last_updated":"2025-04-24T22:33:02Z","snapshot_observed_at":"2026-08-09T07:56:36.279428Z","submitted_at":"2024-09-30T06:27:53Z","title":"FaithEval: Can Your Language Model Stay Faithful to Context, Even If \"The Moon is Made of Marshmallows\"","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.03727","snapshot_observed_at":"2026-08-03T12:42:29.966722Z","title":"Faitheval: Can your language model stay faithful to context,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:29.966722Z"},"links":{"cited_paper":"/paper/2410.03727","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:8b82b50c017d6af49084b6bbb11649bb00ccff3aa4cd998fef8d4e4038692926","observation_id":"f7cad873-edfb-48ef-a837-31c57af71204","resolution":{"observed_at":"2026-08-03T12:42:29.966722Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.14508","last_updated":"2024-06-19T04:00:32Z","snapshot_observed_at":"2026-08-08T03:49:18.086396Z","submitted_at":"2023-08-28T11:53:40Z","title":"LongBench: A Bilingual, Multitask Benchmark for Long Context Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.14508","snapshot_observed_at":"2026-08-03T12:42:30.024622Z","title":"Longbench: A bilingual, multitask benchmark for long context understanding,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.024622Z"},"links":{"cited_paper":"/paper/2308.14508","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:b8eca13c6360ec58732f1ca7786e7fe99cd0d218f7ed6af1eff75ba765e83f71","observation_id":"210c117c-f6b6-4e41-81f6-e8eb94a9f489","resolution":{"observed_at":"2026-08-03T12:42:30.024622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.11088","last_updated":"2023-10-04T10:04:25Z","snapshot_observed_at":"2026-08-09T07:56:50.464493Z","submitted_at":"2023-07-20T17:59:41Z","title":"L-Eval: Instituting Standardized Evaluation for Long Context Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.11088","snapshot_observed_at":"2026-08-03T12:42:30.058516Z","title":"L-eval: Instituting standardized evaluation for long context language models,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.058516Z"},"links":{"cited_paper":"/paper/2307.11088","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:d909561ee864790864382c1cc08285ba4ea5a1d81c967a92e081e292ffe681b6","observation_id":"8e835c7a-4d13-4570-86c9-3b8f29da5fc3","resolution":{"observed_at":"2026-08-03T12:42:30.058516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-03T12:42:30.149171Z","title":"GPT-4 technical report,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.149171Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:d29d924188994770d503ffaf5cab52f8d38a6a4471f4fb9cbd41fe4376f5dbaa","observation_id":"e24b6808-5b13-4d13-8a8c-6bb1466e176f","resolution":{"observed_at":"2026-08-03T12:42:30.149171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.05452","last_updated":"2026-04-15T14:48:07Z","snapshot_observed_at":"2026-07-06T22:09:28.204393Z","submitted_at":"2025-08-07T14:46:30Z","title":"LLMEval-Fair: A Large-Scale Longitudinal Study on Robust and Fair Evaluation of Large Language Models","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.05452","snapshot_observed_at":"2026-08-03T12:42:30.244092Z","title":"Llmeval-3: A large-scale longitudinal study on robust and fair evaluation of large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.244092Z"},"links":{"cited_paper":"/paper/2508.05452","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:bf6d603575badba1a0d79302ff88b7b184c26ec320c14e6999ca6482ee90a899","observation_id":"913f1fd9-15cc-44c9-90b9-087529a19240","resolution":{"observed_at":"2026-08-03T12:42:30.244092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.04422","last_updated":"2025-08-26T08:55:10Z","snapshot_observed_at":"2026-08-09T07:55:39.009426Z","submitted_at":"2024-10-06T09:29:19Z","title":"Long-context Language Models Fail in Basic Retrieval Tasks Without Sufficient Reasoning Steps","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.04422","snapshot_observed_at":"2026-08-03T12:42:30.325461Z","title":"Long-context language models fail in basic retrieval tasks without sufficient reasoning steps,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.325461Z"},"links":{"cited_paper":"/paper/2410.04422","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:fc6efc86f4fcc737e4b2ad034d69ca9d9ad8f349294e04d59c98f3af29f5d504","observation_id":"455df315-ac27-4f1d-998f-f9a14271a753","resolution":{"observed_at":"2026-08-03T12:42:30.325461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.00109","last_updated":"2025-07-31T19:00:11Z","snapshot_observed_at":"2026-08-10T03:44:51.516487Z","submitted_at":"2025-07-31T19:00:11Z","title":"FACTORY: A Challenging Human-Verified Prompt Set for Long-Form Factuality","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.00109","snapshot_observed_at":"2026-08-03T12:42:30.445868Z","title":"FACTORY: A challenging human-verified prompt set for long-form factuality,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.445868Z"},"links":{"cited_paper":"/paper/2508.00109","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:74bac81adcbe94eb2b965e9e8d9ced11f98c94853165488751e223730eb3bb97","observation_id":"c5a0722d-0f4b-4756-be1b-ce80bc465473","resolution":{"observed_at":"2026-08-03T12:42:30.445868Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:30.548515Z","title":"Investigating factuality in long-form text generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.548515Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:33f188bdf9a3f7b59946ebea4c304e32b9fa0b9d67a4d16c0111a407f48b3e28","observation_id":"57a6ff05-a917-4826-b2cd-0a1107c09ab7","resolution":{"observed_at":"2026-08-03T12:42:30.548515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03651","last_updated":"2024-07-14T22:47:13Z","snapshot_observed_at":"2026-08-09T07:55:20.839463Z","submitted_at":"2024-07-04T05:46:20Z","title":"Evaluating Language Model Context Windows: A \"Working Memory\" Test and Inference-time Correction","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.03651","snapshot_observed_at":"2026-08-03T12:42:30.601801Z","title":"Evaluating language model context windows: A “working memory","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.601801Z"},"links":{"cited_paper":"/paper/2407.03651","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:a98c27f081c320dc3910a2c99ba3445b4f40a95fe658d71018c59db3daafb6a0","observation_id":"710a47ae-1039-4e15-b53e-07dbf435cd87","resolution":{"observed_at":"2026-08-03T12:42:30.601801Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.06120","last_updated":"2025-05-09T15:21:44Z","snapshot_observed_at":"2026-08-08T19:35:51.229758Z","submitted_at":"2025-05-09T15:21:44Z","title":"LLMs Get Lost In Multi-Turn Conversation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.06120","snapshot_observed_at":"2026-08-03T12:42:30.762720Z","title":"LLMs get lost in multi-turn conversation,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.762720Z"},"links":{"cited_paper":"/paper/2505.06120","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:5667aaeaec1c6fc4d6ee7f7a660530fe8e1d892bb72ddb98181a50e86122d10b","observation_id":"ce664d0d-eed2-4955-aae1-cd1e914745b3","resolution":{"observed_at":"2026-08-03T12:42:30.762720Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.17588","last_updated":"2025-08-13T04:43:27Z","snapshot_observed_at":"2026-08-09T07:56:49.445341Z","submitted_at":"2024-06-25T14:31:26Z","title":"LongIns: A Challenging Long-context Instruction-based Exam for LLMs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.17588","snapshot_observed_at":"2026-08-03T12:42:30.921429Z","title":"Longins: A challenging long-context instruction-based exam for llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:30.921429Z"},"links":{"cited_paper":"/paper/2406.17588","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:55ef0df86ff883518f864c63bc6a56e2bef6241865024e94ea210ec9803cbc67","observation_id":"53e7d78c-1f1a-4838-9deb-91e99737a786","resolution":{"observed_at":"2026-08-03T12:42:30.921429Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:31.061360Z","title":"Needle in a haystack - pressure testing LLMs,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.061360Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:7bee30f35c186e877b364265f0812af2594e7973e0f8127bc34f7b0fcfa56e1f","observation_id":"31957f20-d04e-46f0-857e-25d9df01d77e","resolution":{"observed_at":"2026-08-03T12:42:31.061360Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:31.201438Z","title":"The needle in a haystack test: Evaluating the performance of LLM RAG systems,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.201438Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:1eabe288ff790c6c8ca4f395a3f6a8c16af842735da7a9a1ca4f8e72bc23ad91","observation_id":"460d326a-1aca-42c0-9bdb-7e9f4c86d1c1","resolution":{"observed_at":"2026-08-03T12:42:31.201438Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:31.324939Z","title":"Sequential-NIAH: A needle-in-a-haystack benchmark for extracting sequential needles from long contexts,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.324939Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:03a67e1a6381c38b19a4ebcb8c3d191b4d4ed4ff8e54f8d5ddb288f6b9108062","observation_id":"6d065bd3-eec9-466c-8ff0-24436c8c8603","resolution":{"observed_at":"2026-08-03T12:42:31.324939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.05167","last_updated":"2025-07-09T14:35:23Z","snapshot_observed_at":"2026-08-08T19:59:08.845558Z","submitted_at":"2025-02-07T18:49:46Z","title":"NoLiMa: Long-Context Evaluation Beyond Literal Matching","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.05167","snapshot_observed_at":"2026-08-03T12:42:31.510188Z","title":"NoLiMa: Long-context evaluation beyond literal matching,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.510188Z"},"links":{"cited_paper":"/paper/2502.05167","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:10273763566558e7963ec42b3b29d1f2d3d1c22d20e885c731384269562cfcf5","observation_id":"7fa9bf73-431d-44e8-9143-c25062bc943f","resolution":{"observed_at":"2026-08-03T12:42:31.510188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.02076","last_updated":"2025-01-23T00:52:08Z","snapshot_observed_at":"2026-08-09T15:51:57.712199Z","submitted_at":"2024-09-03T17:25:54Z","title":"LongGenBench: Benchmarking Long-Form Generation in Long Context LLMs","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.02076","snapshot_observed_at":"2026-08-03T12:42:31.601779Z","title":"LongGenBench: Benchmarking long-form generation in long context LLMs,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.601779Z"},"links":{"cited_paper":"/paper/2409.02076","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:e5336c32dfe30fd617ae8c77a3bd1e377df2f4625e58e1cd8f2dc239d8916b7b","observation_id":"95045da3-68d3-42c6-8b7e-66fa5829bce2","resolution":{"observed_at":"2026-08-03T12:42:31.601779Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.08435","last_updated":"2024-11-21T02:31:15Z","snapshot_observed_at":"2026-08-08T14:54:27.189128Z","submitted_at":"2024-09-13T00:03:19Z","title":"When Context Leads but Parametric Memory Follows in Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.08435","snapshot_observed_at":"2026-08-03T12:42:31.678880Z","title":"When context leads but parametric memory follows in large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.678880Z"},"links":{"cited_paper":"/paper/2409.08435","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:768dd38604d503303ad2d950e6ac98cdcace5c9f935745ea72b0b786b5e12245","observation_id":"c8cd46c7-c063-4845-b551-2ba0ac7a3b28","resolution":{"observed_at":"2026-08-03T12:42:31.678880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.10151","last_updated":"2024-08-19T17:02:06Z","snapshot_observed_at":"2026-08-09T07:54:12.583479Z","submitted_at":"2024-08-19T17:02:06Z","title":"Multilingual Needle in a Haystack: Investigating Long-Context Behavior of Multilingual Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.10151","snapshot_observed_at":"2026-08-03T12:42:31.786033Z","title":"Multilingual needle in a haystack: Investigating long-context behavior of multilingual large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.786033Z"},"links":{"cited_paper":"/paper/2408.10151","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:f2fefc09b7272ffe97e888b6e52a376aee08e9782be9f0c1cbf2225bfba4178f","observation_id":"f78528d1-f79b-434f-99f4-57b32011e6f2","resolution":{"observed_at":"2026-08-03T12:42:31.786033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:31.929789Z","title":"Premise order matters in reasoning with large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:31.929789Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:18df4b93a874a2e0046114d778c1235dcf4bebe54915625dd7d6804224cb9bdb","observation_id":"13b2bafb-bb75-4c4e-ac28-63a4bc9516ea","resolution":{"observed_at":"2026-08-03T12:42:31.929789Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:32.131684Z","title":"A survey on hallucination in large language models: Principles, taxonomy, challenges, and open questions,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:32.131684Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:c7267f2371f891333fb7fe247c960b321f74bc6842ff777369e31fa47f800f33","observation_id":"b47b553c-418e-4539-8c10-d01dc8888f7b","resolution":{"observed_at":"2026-08-03T12:42:32.131684Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.03538","last_updated":"2024-11-05T22:37:43Z","snapshot_observed_at":"2026-08-09T07:56:12.857879Z","submitted_at":"2024-11-05T22:37:43Z","title":"Long Context RAG Performance of Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.03538","snapshot_observed_at":"2026-08-03T12:42:32.334383Z","title":"Long context RAG performance of large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:32.334383Z"},"links":{"cited_paper":"/paper/2411.03538","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:0bdb87bdd6055a54a348a8313b9f47f3e04b8b4515602d9440915bffce89e0fd","observation_id":"2ceedd2a-2e2e-4eaf-99a3-7e0fb54901c7","resolution":{"observed_at":"2026-08-03T12:42:32.334383Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:32.483182Z","title":"Understanding and addressing ai hallucinations in healthcare and life sciences,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:32.483182Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:94c0c038400d901cb27b787f27c8d879935164dd4243842e97ed8c5a6d643b36","observation_id":"3c5603f2-039e-4a8a-8bf8-6e05298667fb","resolution":{"observed_at":"2026-08-03T12:42:32.483182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:32.623439Z","title":"A survey on hallucination in large language and foundation models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:32.623439Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:171c28f2526722719c560bd780a01e10e7a848038dcc34080a16df2c36c846ad","observation_id":"d795bb48-63df-4703-8818-08e98bb84a8a","resolution":{"observed_at":"2026-08-03T12:42:32.623439Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.01463","last_updated":"2023-09-26T20:52:46Z","snapshot_observed_at":"2026-08-07T23:09:45.569340Z","submitted_at":"2023-09-26T20:52:46Z","title":"Creating Trustworthy LLMs: Dealing with Hallucinations in Healthcare AI","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.01463","snapshot_observed_at":"2026-08-03T12:42:32.861808Z","title":"Creating trustworthy llms: Dealing with hallucinations in healthcare ai,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:32.861808Z"},"links":{"cited_paper":"/paper/2311.01463","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:ef96e18ff8febdcec48e317b752043eab31dfdd1cc0b01251124d177838ba1a3","observation_id":"b6183e3a-4e35-451d-a55c-690402e9bd76","resolution":{"observed_at":"2026-08-03T12:42:32.861808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:32.949477Z","title":"Unravelling the mysteries of hallucination in large language models: Strategies for precision in artificial intelligence language generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:32.949477Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:81b361e071a6cd52fb2dacab00eafbbf5b3c113b04d278691b1ce3651d257893","observation_id":"9e7b3a2b-1aaa-45ef-af49-e63e2c680c8d","resolution":{"observed_at":"2026-08-03T12:42:32.949477Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.10011","last_updated":"2024-09-18T20:03:43Z","snapshot_observed_at":"2026-08-08T18:36:18.134813Z","submitted_at":"2024-09-16T05:50:39Z","title":"HALO: Hallucination Analysis and Learning Optimization to Empower LLMs with Retrieval-Augmented Context for Guided Clinical Decision Making","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.10011","snapshot_observed_at":"2026-08-03T12:42:33.116506Z","title":"Halo: Hallucination analysis and learning optimization to empower llms with retrieval-augmented context for guided clinical decision making,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:33.116506Z"},"links":{"cited_paper":"/paper/2409.10011","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:05c06e007f248e7c4096aa6fd6b395372679d1eb520b60bebee5f460dcf300ec","observation_id":"7683181c-f19f-45d2-8798-995906641b62","resolution":{"observed_at":"2026-08-03T12:42:33.116506Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:33.301104Z","title":"Dual process theory for large language models: An overview of using psychology to address hallucination and reliability issues,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:33.301104Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:58892f3227462d58b791cf4fde0adc70af72ae24f17ca83036409d83b03e8652","observation_id":"b2520d3f-e312-4d09-9091-ba8f230c8360","resolution":{"observed_at":"2026-08-03T12:42:33.301104Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:33.498985Z","title":"Factchd: Benchmarking fact-conflicting hallucination detection,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:33.498985Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:c81cf2c6dc33ead2367af8ac34e497dd3350a90e8c5d5b7eb4aeacc5310d53c8","observation_id":"7accfd72-9f62-41eb-86f6-afb3c53627e6","resolution":{"observed_at":"2026-08-03T12:42:33.498985Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:33.585138Z","title":"Explainable hallucination mitigation in large language models: A survey,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:33.585138Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:b939f4f4499684b4e72f687f3a03aff83bba41122e5e6db03f7345643fa706ed","observation_id":"0ce3a900-5ec7-406d-888f-09bb0f7e5ae4","resolution":{"observed_at":"2026-08-03T12:42:33.585138Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:33.631475Z","title":"Zero-resource hallucination detection for text generation via graph- based contextual knowledge triples modeling,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:33.631475Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:622db9f8fb78bd8658fa2a5d76d5bb434120c0efa1e04f7298551e565fe64ed0","observation_id":"74617589-6273-4fe9-ac58-33ef90206025","resolution":{"observed_at":"2026-08-03T12:42:33.631475Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.18344","last_updated":"2023-10-22T14:45:14Z","snapshot_observed_at":"2026-08-10T15:01:44.358201Z","submitted_at":"2023-10-22T14:45:14Z","title":"Chainpoll: A high efficacy method for LLM hallucination detection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.18344","snapshot_observed_at":"2026-08-03T12:42:33.768951Z","title":"Chainpoll: A high efficacy method for llm hallucination detection,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:33.768951Z"},"links":{"cited_paper":"/paper/2310.18344","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:222ad151805880230c334c6d71142545c8bb6cc4a5b07ec6f39979d9a0469187","observation_id":"0ae3a80a-4ac6-4f8d-8820-f1d83c47f1e0","resolution":{"observed_at":"2026-08-03T12:42:33.768951Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:33.945972Z","title":"Zero-knowledge llm hallucination detection and mitigation through fine-grained cross-model consistency,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:33.945972Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:1b1a4a7322ce6b103861c3315862acadd8fd3a7d08effab7059bbcf3e306af79","observation_id":"9617f5ce-a0f7-4fc5-909e-ca9d96e5c604","resolution":{"observed_at":"2026-08-03T12:42:33.945972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:34.076782Z","title":"Detecting and preventing hallucinations in large vision language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:34.076782Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:282792f8cfd0a30354c2ce67bfe2913f6f0e6cad78a975dce4b179917cb54fd2","observation_id":"81e22626-cffb-4f92-b4ed-388a4f111ae6","resolution":{"observed_at":"2026-08-03T12:42:34.076782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:34.203036Z","title":"Beyond probabilities: Unveiling the delicate dance of large language models (llms) and ai-hallucination,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:34.203036Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:ce6684aecc12e102f8271e166f6c4ebeb5b1981453716a1c834e3af4a72c9927","observation_id":"e108ea99-376a-4986-a8fb-fb9b66f873e2","resolution":{"observed_at":"2026-08-03T12:42:34.203036Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.03847","last_updated":"2025-08-21T01:34:15Z","snapshot_observed_at":"2026-08-09T18:33:34.790031Z","submitted_at":"2025-07-05T00:55:15Z","title":"KEA Explain: Explanations of Hallucinations using Graph Kernel Analysis","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.03847","snapshot_observed_at":"2026-08-03T12:42:34.323746Z","title":"Kea explain: Explanations of hallucinations using graph kernel analysis,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:34.323746Z"},"links":{"cited_paper":"/paper/2507.03847","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:546c6d7e27ec4b4c0143573e834bb8864fc4c6b28f65b7d67730e34c666cbdc3","observation_id":"9376b2d0-9f36-402f-bc3c-1787b8668d95","resolution":{"observed_at":"2026-08-03T12:42:34.323746Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:34.491923Z","title":"Mitigating hallucinations in large language models for educational application,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:34.491923Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:9ef1a813637185cb171889b6fc03de5b9360564673b7a37bda5523e6ffc55805","observation_id":"c9297b5a-1651-43c4-99a1-6141320b9419","resolution":{"observed_at":"2026-08-03T12:42:34.491923Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.08285","last_updated":"2025-08-13T22:09:11Z","snapshot_observed_at":"2026-08-09T07:54:52.693645Z","submitted_at":"2025-08-01T20:34:01Z","title":"The Illusion of Progress: Re-evaluating Hallucination Detection in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.08285","snapshot_observed_at":"2026-08-03T12:42:34.626786Z","title":"The illusion of progress: Re-evaluating hallucination detection in llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:34.626786Z"},"links":{"cited_paper":"/paper/2508.08285","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:a584a91ae47f641b7069c50d57d5e9fe261e7566a6c8c42658e21fbcc4cb77b4","observation_id":"b25b8640-ec49-4c07-b20c-3c002161f5c8","resolution":{"observed_at":"2026-08-03T12:42:34.626786Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:34.734659Z","title":"Hallucinations in large language models (llm’s): challenges in mitigation, trust, and future directions,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:34.734659Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:7ae81be6599c7c58f31fb3d696d7cde5f341607ddf9d3353024b34393826368a","observation_id":"cf3c88d3-4efa-48d9-9c0e-e882d09a07a0","resolution":{"observed_at":"2026-08-03T12:42:34.734659Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:34.966239Z","title":"Detecting llm hallucinations using monte carlo simulations on token probabilities,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:34.966239Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:ca0027db593c5efe5c3473c0cb8618ea3d100d2a7c558dfadfbdc15970fd5d0b","observation_id":"50405243-a0fd-4d5e-83e4-3e7539aa92c7","resolution":{"observed_at":"2026-08-03T12:42:34.966239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11747","last_updated":"2023-10-23T01:49:32Z","snapshot_observed_at":"2026-08-10T13:04:51.592649Z","submitted_at":"2023-05-19T15:36:27Z","title":"HaluEval: A Large-Scale Hallucination Evaluation Benchmark for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11747","snapshot_observed_at":"2026-08-03T12:42:35.085906Z","title":"Halueval: A large-scale hallucination evaluation benchmark for large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.085906Z"},"links":{"cited_paper":"/paper/2305.11747","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:de4ababe1fcfd7f58b68ff62e166a5d4f8342e18da6ad505503423a747dc424d","observation_id":"c14198a1-4884-413e-8e8d-e25c7366b5f9","resolution":{"observed_at":"2026-08-03T12:42:35.085906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.02870","last_updated":"2025-06-06T10:50:08Z","snapshot_observed_at":"2026-08-08T01:09:03.836574Z","submitted_at":"2025-06-06T10:50:08Z","title":"Loki's Dance of Illusions: A Comprehensive Survey of Hallucination in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.02870","snapshot_observed_at":"2026-08-03T12:42:35.151076Z","title":"Loki’s dance of illusions: A comprehensive survey of hallucination in large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.151076Z"},"links":{"cited_paper":"/paper/2507.02870","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:553ec3dfd37f7f2bd5fb3896b818436536e58268ac2eac604432eac7d4bdc53d","observation_id":"4191dd39-8dc3-407a-8495-1529dbfdc863","resolution":{"observed_at":"2026-08-03T12:42:35.151076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.15449","last_updated":"2024-01-27T16:19:30Z","snapshot_observed_at":"2026-08-05T02:06:00.890233Z","submitted_at":"2024-01-27T16:19:30Z","title":"Learning to Trust Your Feelings: Leveraging Self-awareness in LLMs for Hallucination Mitigation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.15449","snapshot_observed_at":"2026-08-03T12:42:35.249637Z","title":"Learning to trust your feelings: Leveraging self-awareness in llms for hallucination mitigation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.249637Z"},"links":{"cited_paper":"/paper/2401.15449","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:f68bc518c262a990578d696b2c5c6394138cc9914b74387023aea281112af552","observation_id":"d9a1085d-4a43-466b-8935-08c6268c7b47","resolution":{"observed_at":"2026-08-03T12:42:35.249637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:35.399305Z","title":"Attention-guided self-reflection for zero-shot hallucination detection in large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.399305Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:7669c72e6990f5203a29a89c51631f232a7f5268fc3b70604665d93ea8be0bd7","observation_id":"8f36fb3c-99e2-42b4-b491-6bdc6cb7abac","resolution":{"observed_at":"2026-08-03T12:42:35.399305Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:35.553414Z","title":"Roberta with low-rank adaptation and hierarchical attention for hallucination detection in llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.553414Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:727e423cd58b4892a39de03ade9d3ad660f784dbbf7f744e3ff29b2f212b2504","observation_id":"b1a65246-1597-4ee3-84f3-c95157cf0a39","resolution":{"observed_at":"2026-08-03T12:42:35.553414Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08896","last_updated":"2023-10-11T17:43:28Z","snapshot_observed_at":"2026-07-06T15:03:55.982628Z","submitted_at":"2023-03-15T19:31:21Z","title":"SelfCheckGPT: Zero-Resource Black-Box Hallucination Detection for Generative Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08896","snapshot_observed_at":"2026-08-03T12:42:35.719806Z","title":"Selfcheckgpt: Zero-resource black-box hallucination detection for generative large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.719806Z"},"links":{"cited_paper":"/paper/2303.08896","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:1f171535df1d5bcac67ee5a3c11bff3313d820cf438c25eb42d502353a2ee1d4","observation_id":"0ec7bc13-6cb7-47da-807f-717a4d7628b6","resolution":{"observed_at":"2026-08-03T12:42:35.719806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:35.825354Z","title":"Hallucination detox: Sensitivity dropout (send) for large language model training,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.825354Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:2d54ce7ca3f36e198ebf30cf8f5124d49c65148dede7bd2f5be6bb9e4f25ac84","observation_id":"1901cc84-ce60-4eed-a353-0407e564061b","resolution":{"observed_at":"2026-08-03T12:42:35.825354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.03745","last_updated":"2024-08-12T14:13:15Z","snapshot_observed_at":"2026-07-06T17:55:53.304768Z","submitted_at":"2024-04-04T18:34:32Z","title":"Fakes of Varying Shades: How Warning Affects Human Perception and Engagement Regarding LLM Hallucinations","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.03745","snapshot_observed_at":"2026-08-03T12:42:35.972895Z","title":"Fakes of varying shades: How warning affects human perception and engagement regarding llm hallucinations,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:35.972895Z"},"links":{"cited_paper":"/paper/2404.03745","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:5931ec50e934fd5ff5166ef0f4bca615d133b68f13a06bc6d10b347b1e4e0c7a","observation_id":"0cacecfb-a64f-4919-b4c7-672f89c4fe15","resolution":{"observed_at":"2026-08-03T12:42:35.972895Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.04485","last_updated":"2024-07-05T13:08:58Z","snapshot_observed_at":"2026-08-09T07:54:47.435616Z","submitted_at":"2024-07-05T13:08:58Z","title":"Leveraging Graph Structures to Detect Hallucinations in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.04485","snapshot_observed_at":"2026-08-03T12:42:36.166326Z","title":"Leveraging graph structures to detect hallucinations in large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:36.166326Z"},"links":{"cited_paper":"/paper/2407.04485","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:2ca3315893e7ecd0c28e69e356aa960cc99ceef523d817f98ba8978dcbe80630","observation_id":"1e342b18-ccee-4a2e-98c7-95758ef00614","resolution":{"observed_at":"2026-08-03T12:42:36.166326Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05266","last_updated":"2024-11-03T18:38:50Z","snapshot_observed_at":"2026-08-09T07:56:29.391895Z","submitted_at":"2024-03-08T12:42:36Z","title":"ERBench: An Entity-Relationship based Automatically Verifiable Hallucination Benchmark for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05266","snapshot_observed_at":"2026-08-03T12:42:36.293010Z","title":"Erbench: An entity-relationship based automatically verifiable hallucination benchmark for large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:36.293010Z"},"links":{"cited_paper":"/paper/2403.05266","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:91e92ab7d104b272b561192add07aeb57eda88f03d4d3b03832f1fffb2b3d711","observation_id":"14a96be2-2df9-4db1-95d5-0b386e248d02","resolution":{"observed_at":"2026-08-03T12:42:36.293010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02707","last_updated":"2025-05-18T11:18:56Z","snapshot_observed_at":"2026-08-04T21:25:41.967552Z","submitted_at":"2024-10-03T17:31:31Z","title":"LLMs Know More Than They Show: On the Intrinsic Representation of LLM Hallucinations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02707","snapshot_observed_at":"2026-08-03T12:42:36.406342Z","title":"Llms know more than they show: On the intrinsic representation of llm hallucinations,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:36.406342Z"},"links":{"cited_paper":"/paper/2410.02707","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:0b4cb2cb52a42be28bc63b29c360f1fc0d9310a06de7d855109ecea40803823b","observation_id":"39be5c72-f95e-40b1-abd9-c7337be3e6cc","resolution":{"observed_at":"2026-08-03T12:42:36.406342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:36.595582Z","title":"Mitigating hallucinations in large language models via semantic enrichment of prompts: Insights from biobert and ontological integration,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:36.595582Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:9c818ea4531adbc047279c6660c1b5c6bcb1f25b726d20ceeca78a8cffa9771f","observation_id":"8e835405-9176-466d-aa02-eb6008f81574","resolution":{"observed_at":"2026-08-03T12:42:36.595582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:36.773630Z","title":"Hallusafe at semeval-2024 task 6: An nli-based approach to make llms safer by better detecting hallucinations and overgeneration mistakes,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:36.773630Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:4651b6908cb818dbd53c2fe9507f6e41777572e5d6a0a7f0bd66b0d950c901b1","observation_id":"1d251080-1a8c-478c-a026-4caf55bb4658","resolution":{"observed_at":"2026-08-03T12:42:36.773630Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.05922","last_updated":"2023-09-12T02:34:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-12T02:34:06Z","title":"A Survey of Hallucination in Large Foundation Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.05922","snapshot_observed_at":"2026-08-03T12:42:36.885996Z","title":"A survey of hallucination in large foundation models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:36.885996Z"},"links":{"cited_paper":"/paper/2309.05922","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:83b38410371ac5c0a8c98c4387b8bb8e4fe1b93917894e25accf1a506076bd78","observation_id":"8c5be357-b6cc-48e1-ac63-46c62cd8489c","resolution":{"observed_at":"2026-08-03T12:42:36.885996Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:36.925471Z","title":"Delucionqa: Detecting hallucinations in domain-specific question answering,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:36.925471Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:f7c9782cd13da96b16cd9b407074bb9086b42c1261f22518edcea9fb20680a36","observation_id":"c4349128-d1a8-4b29-9ad9-c0f7c9c7acf5","resolution":{"observed_at":"2026-08-03T12:42:36.925471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:37.018426Z","title":"Mitigation of hallucinations in language models in education: A new approach of comparative and cross-verification,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.018426Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:ae6da1b7d04a46c6dfa5721827e07e826f4510ab6d9b0138873e50ead574665d","observation_id":"982f03dd-a1ad-4405-bec3-2e2c39c5f620","resolution":{"observed_at":"2026-08-03T12:42:37.018426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.06448","last_updated":"2024-06-10T05:48:30Z","snapshot_observed_at":"2026-07-06T17:42:24.853273Z","submitted_at":"2024-03-11T05:51:03Z","title":"Unsupervised Real-Time Hallucination Detection based on the Internal States of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.06448","snapshot_observed_at":"2026-08-03T12:42:37.137620Z","title":"Unsupervised real-time hallucination detection based on the internal states of large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.137620Z"},"links":{"cited_paper":"/paper/2403.06448","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:c2ff61672080d32851b8c2c4ee42644a165075de5abf48ee789548a83ac1056e","observation_id":"144a7740-1d41-47d6-9000-e267cbaf9457","resolution":{"observed_at":"2026-08-03T12:42:37.137620Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.04175","last_updated":"2024-06-25T18:37:19Z","snapshot_observed_at":"2026-08-09T07:54:50.551762Z","submitted_at":"2024-06-06T15:32:29Z","title":"Confabulation: The Surprising Value of Large Language Model Hallucinations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.04175","snapshot_observed_at":"2026-08-03T12:42:37.288954Z","title":"Confabulation: The surprising value of large language model hallucinations,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.288954Z"},"links":{"cited_paper":"/paper/2406.04175","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:673326834add7429fe379c8fc76d47789b1eae18c0726840b033ec4aaf93c10a","observation_id":"e9f584bd-ba29-419a-8e9d-782d0d6915c2","resolution":{"observed_at":"2026-08-03T12:42:37.288954Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.01313","last_updated":"2024-01-08T16:19:17Z","snapshot_observed_at":"2026-07-06T17:10:56.398607Z","submitted_at":"2024-01-02T17:56:30Z","title":"A Comprehensive Survey of Hallucination Mitigation Techniques in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.01313","snapshot_observed_at":"2026-08-03T12:42:37.377620Z","title":"A comprehensive survey of hallucination mitigation techniques in large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.377620Z"},"links":{"cited_paper":"/paper/2401.01313","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:7ab42a4a9c3f876959c5dba7d356a3b9d3a3476d7ed20d338f0ba340794437d4","observation_id":"432dcd8d-2352-4c01-9de2-16185e1c96cb","resolution":{"observed_at":"2026-08-03T12:42:37.377620Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:37.478912Z","title":"Investigating hallucination tendencies of large language models in japanese and english,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.478912Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:05a096018a313db28af44aa113b6d483b46ed29209fb319512723413c49f48e5","observation_id":"49ca31cc-19e8-49fc-89c9-85f742d66937","resolution":{"observed_at":"2026-08-03T12:42:37.478912Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.03987","last_updated":"2023-08-12T14:57:37Z","snapshot_observed_at":"2026-08-09T07:54:49.290541Z","submitted_at":"2023-07-08T14:25:57Z","title":"A Stitch in Time Saves Nine: Detecting and Mitigating Hallucinations of LLMs by Validating Low-Confidence Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.03987","snapshot_observed_at":"2026-08-03T12:42:37.546861Z","title":"A stitch in time saves nine: Detecting and mitigating hallucinations of llms by validating low-confidence generation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.546861Z"},"links":{"cited_paper":"/paper/2307.03987","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:2a85e78a8074f42fb051d47ed8d3ad8b67e23129edf45c79cf4e9272561a7040","observation_id":"dbadf06a-0d02-4769-9d21-68a0e55def89","resolution":{"observed_at":"2026-08-03T12:42:37.546861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.18715","last_updated":"2024-06-05T13:53:42Z","snapshot_observed_at":"2026-08-09T07:55:47.071847Z","submitted_at":"2024-03-27T16:04:47Z","title":"Mitigating Hallucinations in Large Vision-Language Models with Instruction Contrastive Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.18715","snapshot_observed_at":"2026-08-03T12:42:37.680719Z","title":"Mitigating hallucinations in large vision-language models with instruction contrastive decoding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.680719Z"},"links":{"cited_paper":"/paper/2403.18715","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:b7b450aabd2d01eaba4cc15e2b74b88314a22053ad87ab396a6d61351f599e67","observation_id":"382aaa0c-f068-4db3-9fcc-8eead70e1569","resolution":{"observed_at":"2026-08-03T12:42:37.680719Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.06601","last_updated":"2025-04-26T02:00:33Z","snapshot_observed_at":"2026-08-09T07:54:50.022836Z","submitted_at":"2024-09-10T15:51:15Z","title":"LaMsS: When Large Language Models Meet Self-Skepticism","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.06601","snapshot_observed_at":"2026-08-03T12:42:37.789483Z","title":"Alleviating hallucinations in large language models with scepticism modeling,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.789483Z"},"links":{"cited_paper":"/paper/2409.06601","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:2b4e6e7b9d234def3b9e79907c25e854eeed4dec4d6b103e05732fdcde9ceec1","observation_id":"7112456d-5ecc-4f9e-b9bc-87a377499d5e","resolution":{"observed_at":"2026-08-03T12:42:37.789483Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:37.900100Z","title":"Detecting and reducing the factual hallucinations of large language models with metamorphic testing,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:37.900100Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:dfbea85c7c8bee988107560fdd431eda97997de69c88b9bf1253f48a1a81d513","observation_id":"689a9a5d-4c70-422a-818a-57ea9a3c432c","resolution":{"observed_at":"2026-08-03T12:42:37.900100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.09801","last_updated":"2024-09-23T02:05:02Z","snapshot_observed_at":"2026-08-10T10:54:54.776313Z","submitted_at":"2024-02-15T08:58:03Z","title":"EFUF: Efficient Fine-grained Unlearning Framework for Mitigating Hallucinations in Multimodal Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.09801","snapshot_observed_at":"2026-08-03T12:42:38.006298Z","title":"Efuf: Efficient fine-grained unlearning framework for mitigating hallucinations in multimodal large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.006298Z"},"links":{"cited_paper":"/paper/2402.09801","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:7bd2ed04ce55bd94d20d0aee37cac4f1805b6bcceaa25a9eb86ac2be501026bf","observation_id":"0a52ba67-cec8-4bea-a320-f134e0de08ec","resolution":{"observed_at":"2026-08-03T12:42:38.006298Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.02889","last_updated":"2024-08-19T07:53:17Z","snapshot_observed_at":"2026-08-10T10:11:57.529324Z","submitted_at":"2024-03-05T11:50:01Z","title":"InterrogateLLM: Zero-Resource Hallucination Detection in LLM-Generated Answers","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.02889","snapshot_observed_at":"2026-08-03T12:42:38.096256Z","title":"Interrogatellm: Zero-resource hallucination detection in llm-generated answers,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.096256Z"},"links":{"cited_paper":"/paper/2403.02889","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:1aea7985580e803058ca21233e15fabad905776d75f02d4bc359db82d16503a2","observation_id":"fb68b6f6-33e3-4286-976a-35847502be34","resolution":{"observed_at":"2026-08-03T12:42:38.096256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:38.175846Z","title":"Siren’s song in the ai ocean: A survey on hallucination in large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.175846Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:b53c09beec5ce6589a2ce4d0f21368c2e30a6809abbfe392476b56862a00d46f","observation_id":"5dba7273-2578-4879-b2c8-f6029720d81b","resolution":{"observed_at":"2026-08-03T12:42:38.175846Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.04699","last_updated":"2025-08-06T17:58:36Z","snapshot_observed_at":"2026-08-09T07:54:59.815573Z","submitted_at":"2025-08-06T17:58:36Z","title":"Hop, Skip, and Overthink: Diagnosing Why Reasoning Models Fumble during Multi-Hop Analysis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.04699","snapshot_observed_at":"2026-08-03T12:42:38.241011Z","title":"Hop, skip, and overthink: Diagnosing why reasoning models fumble during multi-hop analysis,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.241011Z"},"links":{"cited_paper":"/paper/2508.04699","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:61c4f042a64e098603d10e7a258d176c8bd7af8402ed86f8be76e27135bb378e","observation_id":"835c536e-5752-4126-b27a-d1e8a7a512c0","resolution":{"observed_at":"2026-08-03T12:42:38.241011Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:38.324828Z","title":"Prompting for faithfulness: When “don’t make it up","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.324828Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:6cf2a098f84afdee2b9352503d648c6ff89ebea020338a5ebf32f40aedaf141e","observation_id":"954c3029-27b7-42b8-a9fd-14374147bc1f","resolution":{"observed_at":"2026-08-03T12:42:38.324828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:38.448802Z","title":"Aspects of human memory and large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.448802Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:316da8e36b4a9150f9330ecd250477e1593cb8eaad05429eaf6987e1468cb207","observation_id":"63bf5cec-390a-4172-bda1-ace5d7036eb9","resolution":{"observed_at":"2026-08-03T12:42:38.448802Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:38.513703Z","title":"More is less: Increased processing of unwanted memories facilitates forgetting,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.513703Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:39674cb546ddd05398cc8145ebf73acce4ad7fd444d250f54733866a067c2acf","observation_id":"b927ae67-16c3-4d6f-b788-0e3c66fcb232","resolution":{"observed_at":"2026-08-03T12:42:38.513703Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.19578","last_updated":"2025-08-27T05:23:22Z","snapshot_observed_at":"2026-08-08T08:00:50.125294Z","submitted_at":"2025-08-27T05:23:22Z","title":"Towards a Holistic and Automated Evaluation Framework for Multi-Level Comprehension of LLMs in Book-Length Contexts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.19578","snapshot_observed_at":"2026-08-03T12:42:38.614626Z","title":"Towards a holistic and automated evaluation framework for multi-level comprehension of LLMs in book-length contexts,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.614626Z"},"links":{"cited_paper":"/paper/2508.19578","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:941a90d300951d713c4057070a182d4a68aa9657f567e1142c123b254dbbcd0c","observation_id":"013e6555-b428-4ea6-90f3-0a00f04d1e91","resolution":{"observed_at":"2026-08-03T12:42:38.614626Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:38.715156Z","title":"Abductive commonsense reasoning,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.715156Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:e0e6d15a6b51402a6487d10ac5af1e115ab69dc999495bd38573c111ecc9aa01","observation_id":"e0036ca5-c46d-44df-82b1-a931b6c71c16","resolution":{"observed_at":"2026-08-03T12:42:38.715156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20947","last_updated":"2025-06-15T21:44:25Z","snapshot_observed_at":"2026-07-06T18:23:23.565502Z","submitted_at":"2024-05-31T15:44:33Z","title":"OR-Bench: An Over-Refusal Benchmark for Large Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20947","snapshot_observed_at":"2026-08-03T12:42:38.838889Z","title":"OR-Bench: An over-refusal benchmark for large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.838889Z"},"links":{"cited_paper":"/paper/2405.20947","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:35a0555ac175652deb884bf77d7368f366f66427f77363b71587e8b0185cf00e","observation_id":"93cf701d-1912-4476-9664-f2189ee9d5b2","resolution":{"observed_at":"2026-08-03T12:42:38.838889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:38.944759Z","title":"Evaluating long-context language models on distributed evidence reasoning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:38.944759Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:002e8c5b70b01394d1271113233486c4ebba7b1065f4100ce11109170076fb22","observation_id":"05c81e98-d1aa-4fca-8bfb-4d122a477b71","resolution":{"observed_at":"2026-08-03T12:42:38.944759Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.045955Z","title":"Context rot: How increasing input tokens impacts llm performance,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.045955Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:61c7b503171915d3b4fdbc59b1baf81eaa4cda75aa319fd3974345242dee3e9c","observation_id":"9463721d-2880-4c77-be62-dc231d4e55c5","resolution":{"observed_at":"2026-08-03T12:42:39.045955Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.137754Z","title":"Scrolls: Standardized comparison over long language sequences,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.137754Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:d2144fcf14e77001fd8ba2b986a0696db484ee8242f9299162c568378db76a7e","observation_id":"0fd7e218-193c-41e9-be39-319faf7e1ddc","resolution":{"observed_at":"2026-08-03T12:42:39.137754Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09296","last_updated":"2024-07-01T03:38:57Z","snapshot_observed_at":"2026-07-06T15:43:03.878575Z","submitted_at":"2023-06-15T17:20:46Z","title":"KoLA: Carefully Benchmarking World Knowledge of Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.09296","snapshot_observed_at":"2026-08-03T12:42:39.212896Z","title":"Kola: Carefully benchmarking world knowledge of large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.212896Z"},"links":{"cited_paper":"/paper/2306.09296","citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:edbf2c0f5ad111e206c6b00e5691df3cf3ce802cbc4375f253389550f64c0788","observation_id":"709b7117-57a4-4aac-9e3a-d5769017d943","resolution":{"observed_at":"2026-08-03T12:42:39.212896Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.485621Z","title":"Provide your answers in the following format: Question 1: [YOUR ANSWER] Question 2: [YOUR ANSWER]","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":96,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.485621Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:9deea69d8d602fe386eea86277455bcac474f995d8a288035a6f665e39c59694","observation_id":"06f3ae5c-4724-401b-b88b-67541c20c0eb","resolution":{"observed_at":"2026-08-03T12:42:39.485621Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.554368Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.554368Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:cd93604f72f4809e721bb9e93f411bb43e76bc0bd072d276619e465780d938bf","observation_id":"d1967022-c456-4ee6-9cb6-d90fe852f863","resolution":{"observed_at":"2026-08-03T12:42:39.554368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.617906Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.617906Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:c92b0c004a850773e8acf520519d872d6c54a01ff3bac31e251ab8c9346c08b0","observation_id":"b5f6ecbd-b53c-43a6-b952-7081c963156c","resolution":{"observed_at":"2026-08-03T12:42:39.617906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.721849Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":99,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.721849Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:7c32f9ee2274617fdf358736240447249843febfe91ed925fe81df362795c034","observation_id":"550d1c51-808d-4baa-8004-5ea19886332d","resolution":{"observed_at":"2026-08-03T12:42:39.721849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.811946Z","title":"Not mentioned in the text or story","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.811946Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:963c2ad4911a06be1158b0a18abd3f167d9a2a5b864521f0d623c9e65d850a89","observation_id":"cf550189-204f-4515-af23-9ba96030bbce","resolution":{"observed_at":"2026-08-03T12:42:39.811946Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.848743Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":101,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.848743Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:4f3af8cca19ae3e5a37227e079ff0b558a56c3e4da08c03ea220c9ffe2b8278f","observation_id":"b7054414-b5d4-437a-8abb-e74c3f18aea8","resolution":{"observed_at":"2026-08-03T12:42:39.848743Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:42:39.941106Z","title":"don’t make it up","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs","version":2},"reference_index":102,"source":"pdf_text","source_observed_at":"2026-08-03T12:42:39.941106Z"},"links":{"citing_paper":"/paper/2601.02023"},"observation_digest":"sha256:ccc11d923cabfbf63b075634091cee5969ac178c0ad26c0848efdbd34838f387","observation_id":"2de45785-57b0-42c6-8c68-dd0a2f86fd2f","resolution":{"observed_at":"2026-08-03T12:42:39.941106Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2601.02023","last_updated":"2026-07-14T21:47:31Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-08T04:32:46.940287Z","submitted_at":"2026-01-05T11:30:56Z","title":"Not All Needles Are Found: How Fact Distribution and Don't Make It Up Prompts Shape Retrieval, Reasoning, and Hallucination in Long-Context LLMs"},"reference_resolution":{"displayed":99,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":99,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":99},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 99 of 99 outbound references and 1 inbound Pith citation observation for arXiv:2601.02023."}