{"as_of":"2026-08-14T02:43:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c76e5598c52732f670888338cad369e0d2c5763c4370c49f553c0545eecb5bfe","coverage":[{"denominator":57,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":57,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T05:55:14.902112Z","state":"measured"},{"denominator":61,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":61,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:19:14.992316Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.17032","snapshot_observed_at":"2026-08-07T15:19:14.992316Z","title":"Mintqa: A multi-hop ques- tion answering benchmark for evaluating llms on new and tail knowledge","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15872","last_updated":"2025-05-23T10:16:01Z","snapshot_observed_at":"2026-08-13T21:14:11.003292Z","submitted_at":"2025-05-21T14:44:40Z","title":"InfoDeepSeek: Benchmarking Agentic Information Seeking for Retrieval-Augmented Generation","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T15:19:14.992316Z"},"links":{"cited_paper":"/paper/2412.17032","citing_paper":"/paper/2505.15872"},"observation_digest":"sha256:abd9b5ff0178901a94ec8ea9789b4cc314db8726a8086b657a1baad293679eda","observation_id":"0eee85ad-94c3-4cb4-99d2-2ce6b8a33120","resolution":{"observed_at":"2026-08-07T15:19:14.992316Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.17032","snapshot_observed_at":"2026-08-06T18:00:38.663920Z","title":"In Advances in Neural Information Processing Systems, volume 36, pages 45870–45894","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.09477","last_updated":"2025-07-16T15:44:18Z","snapshot_observed_at":"2026-08-13T04:23:58.061031Z","submitted_at":"2025-07-13T03:29:41Z","title":"Towards Agentic RAG with Deep Reasoning: A Survey of RAG-Reasoning Systems in LLMs","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T18:00:38.663920Z"},"links":{"cited_paper":"/paper/2412.17032","citing_paper":"/paper/2507.09477"},"observation_digest":"sha256:e65045df90cfe5b9320994303b0ff0d06b58b75100d4f8bb5c903d3aa7067522","observation_id":"24d508a4-30e9-4d4d-8781-b77d9e3ff792","resolution":{"observed_at":"2026-08-06T18:00:38.663920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"cited_work":{"arxiv_id":"2412.17032","doi":"10.48550/arxiv.2412.17032","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.17032","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2412.17032 , year=","venue":"Research Explorer (The University of Manchester)","work_id":"23fa36cc-876f-46af-9218-0bb0a11e152a","year":2024},"citing_paper":{"arxiv_id":"2606.11470","last_updated":"2026-08-10T19:13:02Z","snapshot_observed_at":"2026-08-14T02:09:32.579515Z","submitted_at":"2026-06-09T21:59:37Z","title":"The Periodic Table of LLM Reasoning: A Structured Survey of Reasoning Paradigms, Methods, and Failure Modes","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-06-27T12:59:51.091008Z"},"links":{"cited_paper":"/paper/2412.17032","citing_paper":"/paper/2606.11470"},"observation_digest":"sha256:21b857d5a24da1c90cdc4490771f55b3fdddbbc91960926614f820fe515c0eee","observation_id":"9c9cb728-9e24-481d-9046-4501a99e1a5d","resolution":{"observed_at":"2026-06-27T13:00:56.135207Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"cited_work":{"arxiv_id":"2412.17032","doi":"10.48550/arxiv.2412.17032","metadata_source":"arxiv_reference","pith_arxiv_id":"2412.17032","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2412.17032 , year=","venue":"Research Explorer (The University of Manchester)","work_id":"23fa36cc-876f-46af-9218-0bb0a11e152a","year":2024},"citing_paper":{"arxiv_id":"2606.12087","last_updated":"2026-06-10T13:49:11Z","snapshot_observed_at":"2026-08-02T15:03:50.335992Z","submitted_at":"2026-06-10T13:49:11Z","title":"FORT-Searcher: Synthesizing Shortcut-Resistant Search Tasks for Training Deep Search Agents","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-06-27T10:01:45.332920Z"},"links":{"cited_paper":"/paper/2412.17032","citing_paper":"/paper/2606.12087"},"observation_digest":"sha256:da96d8dea2001b7326795a7da2db5384c55eac9cc8f0f60eb3c41e58cf9893ab","observation_id":"e47dcf84-bf68-4a0b-ac11-8d4ce0bb4789","resolution":{"observed_at":"2026-07-03T10:27:56.385466Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.17032/citation-record","integrity":"/paper/2412.17032/integrity","json":"/paper/2412.17032/citation-record.json","paper":"/paper/2412.17032"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-08-10T14:07:02.234322Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-11T05:55:14.648021Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.648021Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:400105fb1fb286f2a8d3db6414fa302c2d44148e5743d1b08ce8a5b2c8146860","observation_id":"c3fa33c9-9fff-4ece-937b-ff73f8c67a48","resolution":{"observed_at":"2026-08-11T05:55:14.648021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.02393","last_updated":"2020-04-06T03:54:38Z","snapshot_observed_at":"2026-08-12T02:54:04.248312Z","submitted_at":"2020-04-06T03:54:38Z","title":"Learning to Recover Reasoning Chains for Multi-Hop Question Answering via Cooperative Games","version":1},"cited_work":{"arxiv_id":"2004.02393","doi":null,"metadata_source":"pith","pith_arxiv_id":"2004.02393","snapshot_observed_at":"2026-08-11T05:55:15.504437Z","title":"Learning to Recover Reasoning Chains for Multi-Hop Question Answering via Cooperative Games","venue":"cs.CL","work_id":"e068da3e-47b2-4a24-b9f4-56d345ff6511","year":2020},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.654283Z"},"links":{"cited_paper":"/paper/2004.02393","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:17bc15666be81e0010e9deccea446512979c4c5df8a0242c4be014d4fbdeeaed","observation_id":"1ccf4b38-807d-498b-ae17-62afec84d6b8","resolution":{"observed_at":"2026-08-11T05:55:15.509178Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-11T05:55:14.659377Z","title":"Christian Keller etc","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.659377Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:1c08c2611b88f9d6eb925d5ac32eaf75fec1d1676f763ca4515558c4aa8b8240","observation_id":"15774753-2e12-4a10-9856-cf60599fbd4a","resolution":{"observed_at":"2026-08-11T05:55:14.659377Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2011.01060","last_updated":"2020-11-12T07:47:48Z","snapshot_observed_at":"2026-07-06T10:10:56.466018Z","submitted_at":"2020-11-02T15:42:40Z","title":"Constructing A Multi-hop QA Dataset for Comprehensive Evaluation of Reasoning Steps","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2011.01060","snapshot_observed_at":"2026-08-11T05:55:14.664341Z","title":"Nguyen, Saku Sugawara, and Akiko Aizawa","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.664341Z"},"links":{"cited_paper":"/paper/2011.01060","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:6973b57ee7977286dd494585d0caea4f682b0fc0452866d11dbd872f8c8a9652","observation_id":"04c5b463-0714-4264-a42a-5631638fdc3a","resolution":{"observed_at":"2026-08-11T05:55:14.664341Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.669252Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.669252Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:864819f441326c996c530395ef85434199b2f4cf341c19a0225b93e175ea823a","observation_id":"500c559b-11b4-4c74-8f0e-2d9a29db9ef0","resolution":{"observed_at":"2026-08-11T05:55:14.669252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12186","last_updated":"2024-11-12T13:24:25Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-09-18T17:57:57Z","title":"Qwen2.5-Coder Technical Report","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12186","snapshot_observed_at":"2026-08-11T05:55:14.674445Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.674445Z"},"links":{"cited_paper":"/paper/2409.12186","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:9ec01ba21b267344e4b41dd1a2cdceca0d2aebfc59a15bea1a0bce3799ef7ab3","observation_id":"06437c8e-b0eb-469b-93d5-ccc896513f73","resolution":{"observed_at":"2026-08-11T05:55:14.674445Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.01782","last_updated":"2024-10-02T17:37:18Z","snapshot_observed_at":"2026-08-13T02:11:42.935780Z","submitted_at":"2024-10-02T17:37:18Z","title":"Open-RAG: Enhanced Retrieval-Augmented Reasoning with Open-Source Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.01782","snapshot_observed_at":"2026-08-11T05:55:14.679583Z","title":"Joty, and Md","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.679583Z"},"links":{"cited_paper":"/paper/2410.01782","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:a4d3bdf690d59bed1e2d87e9babb21516285b1263ddc44de44c5517d19dd74c3","observation_id":"3b02736f-ad64-4f45-9448-9412ecd0ec97","resolution":{"observed_at":"2026-08-11T05:55:14.679583Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.684533Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.684533Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:3700173cc0a22688a3c5360ea81d31f91c5d307d3ae5bd3393c27c46f707242f","observation_id":"cd17f135-a0c9-4e03-b441-1586e96e3c1f","resolution":{"observed_at":"2026-08-11T05:55:14.684533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.811137Z","title":null,"venue":null,"work_id":"7d5e3d9e-3bf3-4b4b-8d71-7ec39c6ef85b","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.688361Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:7be9fb1592ce84edde7fb0ed2d2a45bdcaceb88ecbcd2fdf9130c19d999320f1","observation_id":"5c031ec2-67df-42d5-852b-7bb713bb8a1a","resolution":{"observed_at":"2026-08-11T05:55:15.815764Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06825","last_updated":"2023-10-10T17:54:58Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-10T17:54:58Z","title":"Mistral 7B","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06825","snapshot_observed_at":"2026-08-11T05:55:14.692321Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.692321Z"},"links":{"cited_paper":"/paper/2310.06825","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:129eba34d87897a8df8fa3777f7efa6f9c455fcf56f5533750883ac8e46144bb","observation_id":"fbad84af-98c5-481e-bdbb-9f151e4b04d1","resolution":{"observed_at":"2026-08-11T05:55:14.692321Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.06984","last_updated":"2023-07-06T18:52:08Z","snapshot_observed_at":"2026-08-13T11:44:16.668662Z","submitted_at":"2023-05-11T17:14:33Z","title":"Evaluating Open-Domain Question Answering in the Era of Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.06984","snapshot_observed_at":"2026-08-11T05:55:14.696605Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.696605Z"},"links":{"cited_paper":"/paper/2305.06984","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:a96ff9b654fe5bfb12d87919ee97740030be4cce84db240b32053cb2a483c78a","observation_id":"3a00bcdd-b499-4755-a9fd-2bb4ff41f467","resolution":{"observed_at":"2026-08-11T05:55:14.696605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.797889Z","title":null,"venue":null,"work_id":"7b5efc12-898a-41e9-a327-1f8fe081600f","year":2019},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.700990Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:b30a1e7c656eb03c80177ec282d2245dcfff2df416b712060cb6e70e150f7525","observation_id":"8851b154-f772-4be2-92ea-668b71633222","resolution":{"observed_at":"2026-08-11T05:55:15.802024Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.705536Z","title":"Parikh, Chris Alberti, Danielle Epstein, Illia Polosukhin, Jacob Devlin, Kenton Lee, Kristina Toutanova, Llion Jones, Matthew Kelcey, Ming-Wei Chang, Andrew M","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.705536Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:f7b1e8f0530af2e2a61e9b274cf3d8e9379fdecaf288569144de3d106af9b5ae","observation_id":"0955d653-7425-44d3-94e7-ab48c022ca2f","resolution":{"observed_at":"2026-08-11T05:55:14.705536Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.709849Z","title":"Gonzalez, Haotong Zhang, and Ion Stoica","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.709849Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:429bbb8b7ae0ea45078fcaeaa6bd978ab62ae68e11a409745e630d6052463cb7","observation_id":"d43c95b4-fa1f-4398-8c3a-b04ee565d55b","resolution":{"observed_at":"2026-08-11T05:55:14.709849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.11401","last_updated":"2021-04-12T15:42:18Z","snapshot_observed_at":"2026-08-07T05:44:30.677502Z","submitted_at":"2020-05-22T21:34:34Z","title":"Retrieval-Augmented Generation for Knowledge-Intensive NLP Tasks","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.11401","snapshot_observed_at":"2026-08-11T05:55:14.714204Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.714204Z"},"links":{"cited_paper":"/paper/2005.11401","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:6ec2581ab1155c721038ed71f1e9ab55186758d8d701fe3dd158ed2fbae3d3bf","observation_id":"bb878223-5a60-4e33-b24d-03755f0d0a05","resolution":{"observed_at":"2026-08-11T05:55:14.714204Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05110","last_updated":"2022-11-09T18:58:29Z","snapshot_observed_at":"2026-08-13T13:45:48.754123Z","submitted_at":"2022-11-09T18:58:29Z","title":"Large Language Models with Controllable Working Memory","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05110","snapshot_observed_at":"2026-08-11T05:55:14.719947Z","title":"Yu, and Surinder Kumar","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.719947Z"},"links":{"cited_paper":"/paper/2211.05110","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:f8c084c951c3312f6a704eb93efd021993240db601a959c921c1a56ff850c0ab","observation_id":"02b048db-2b15-4816-8818-97996dcbb148","resolution":{"observed_at":"2026-08-11T05:55:14.719947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.766474Z","title":null,"venue":null,"work_id":"5783b642-caa0-47b3-a988-5062e18b1b7f","year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.724903Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:444e312f581881a8591ca432c31fd626ea636210141ed725b2ea080c7f656bbb","observation_id":"55862471-7f16-4987-a417-96e3dfdfbf31","resolution":{"observed_at":"2026-08-11T05:55:15.770632Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.753371Z","title":null,"venue":null,"work_id":"3ee72fd1-483b-4b2f-89c3-5922ffbc7fe9","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.732120Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:09fa1972decc9b6713c02b01eea87c1784e37ebd4cd909c9bbcdcb9626c13da6","observation_id":"c8441089-dcbf-443e-ace2-e2863a413b03","resolution":{"observed_at":"2026-08-11T05:55:15.757538Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.13492","last_updated":"2024-03-27T18:48:34Z","snapshot_observed_at":"2026-08-13T04:14:03.675012Z","submitted_at":"2024-02-21T03:05:50Z","title":"Retrieval Helps or Hurts? A Deeper Dive into the Efficacy of Retrieval Augmentation to Language Models","version":3},"cited_work":{"arxiv_id":"2402.13492","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.13492","snapshot_observed_at":"2026-08-11T05:55:15.375352Z","title":"Retrieval Helps or Hurts? A Deeper Dive into the Efficacy of Retrieval Augmentation to Language Models","venue":"cs.CL","work_id":"02c7df66-84ef-4285-aa9f-1256ef735796","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.735869Z"},"links":{"cited_paper":"/paper/2402.13492","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:e5e2f0368df2658c65637a9d7774cbd89117dc6315167e4e906cc8d01b72dc84","observation_id":"6e0d918c-de0a-447f-9e7d-c2d6a31eef0f","resolution":{"observed_at":"2026-08-11T05:55:15.381405Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.739545Z","title":null,"venue":null,"work_id":"75fe4b5f-06c6-46fc-b98c-d103d0dfbb1a","year":2022},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.739708Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:4c2f7f0555054ca4a124d8bc01774967b2649c5de4f455393a1412127fc0f514","observation_id":"ef7d9c5f-36a5-43c4-8394-17cb944a296f","resolution":{"observed_at":"2026-08-11T05:55:15.743955Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.725914Z","title":null,"venue":null,"work_id":"04cafb79-0ecf-45be-b932-ff4f73555e2e","year":2019},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.743987Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:18d3b920072b024e32dd49a7ad3b90206b0e2f876dfa8500dd12f1fbf1335256","observation_id":"0fa58127-4eac-4b8f-ac01-1566fc47bb09","resolution":{"observed_at":"2026-08-11T05:55:15.730181Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.07899","last_updated":"2021-12-15T05:33:27Z","snapshot_observed_at":"2026-08-13T17:15:03.424893Z","submitted_at":"2021-12-15T05:33:27Z","title":"Large Dual Encoders Are Generalizable Retrievers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.07899","snapshot_observed_at":"2026-08-11T05:55:14.748838Z","title":"Hall, Ming-Wei Chang, and Yinfei Yang","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.748838Z"},"links":{"cited_paper":"/paper/2112.07899","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:42510d1ce31f4d364a481ac2796a3797214094c0be375c1bf6fd0f9c5e21b09c","observation_id":"2c7b9b02-5790-4afd-8d9b-fe30b38ab018","resolution":{"observed_at":"2026-08-11T05:55:14.748838Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.711888Z","title":"Guo, and Xueqi Cheng","venue":null,"work_id":"8c6837a4-3e95-47ba-8580-c029a2669219","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.754170Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:77525f219236560205555ca3821718bcea49c6b92f3f11e0f3994267b377855b","observation_id":"3ec105f1-b80f-4935-bc89-fca35185e6d7","resolution":{"observed_at":"2026-08-11T05:55:15.716433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.11019","last_updated":"2024-11-19T05:35:02Z","snapshot_observed_at":"2026-08-13T10:51:02.271954Z","submitted_at":"2023-07-20T16:46:10Z","title":"Investigating the Factual Knowledge Boundary of Large Language Models with Retrieval Augmentation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.11019","snapshot_observed_at":"2026-08-11T05:55:14.758418Z","title":"Liu, Hao Tian, Huaqin Wu, Ji rong Wen, and Haifeng Wang","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.758418Z"},"links":{"cited_paper":"/paper/2307.11019","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:4d80f578cbd253384de6a10741b9ca7c0802e91d8cc808efb2ebe43fd98d9dcb","observation_id":"a05b6d6b-a9dc-4058-8d9b-a2fb603eff54","resolution":{"observed_at":"2026-08-11T05:55:14.758418Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.762778Z","title":"Robertson and Hugo Zaragoza","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.762778Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:934470bf2356e6e84e845acc1cd8cc9809c5b9c0db3da6ece0956c3fc3144a85","observation_id":"7e030baa-9c5a-4ac3-a45e-0d9f5340b4e0","resolution":{"observed_at":"2026-08-11T05:55:14.762778Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.01613","last_updated":"2022-10-04T13:54:29Z","snapshot_observed_at":"2026-08-13T14:13:22.331195Z","submitted_at":"2022-10-04T13:54:29Z","title":"Mintaka: A Complex, Natural, and Multilingual Dataset for End-to-End Question Answering","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.01613","snapshot_observed_at":"2026-08-11T05:55:14.767197Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.767197Z"},"links":{"cited_paper":"/paper/2210.01613","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:92bc84486eae5a1e55f9299aa64dd4347ac195ead1333ac79f68f2855080649f","observation_id":"1afe389f-61dd-4943-85fd-8433934924f2","resolution":{"observed_at":"2026-08-11T05:55:14.767197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.690449Z","title":null,"venue":null,"work_id":"a0b3f3a6-7a8b-4acc-a53b-b1c52854f8ce","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.771553Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:4af1aeb7505ce2e564a98986c53af92f31017c3f1a73bf6e1dc126c4763e9585","observation_id":"219b7b2a-d79e-4d49-a115-73da98a3c1d0","resolution":{"observed_at":"2026-08-11T05:55:15.694290Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.677138Z","title":null,"venue":null,"work_id":"83f45372-768c-4a31-a1fb-361ee885f38e","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.776241Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:640c7f073c7ec403220c2107f4f13065813d02be0c7efcf81bdb50c958d4f14d","observation_id":"968fa1e3-3d56-43a0-af76-4db2c8f9dc94","resolution":{"observed_at":"2026-08-11T05:55:15.681754Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.780421Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.780421Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:aa2b393299019fd23ad37caa618c016fc7992b65d2baa55bfd993043773597dd","observation_id":"5ad18645-bd1e-461f-8481-c266fa38f5a0","resolution":{"observed_at":"2026-08-11T05:55:14.780421Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.09741","last_updated":"2023-05-30T15:22:50Z","snapshot_observed_at":"2026-08-13T13:18:45.782206Z","submitted_at":"2022-12-19T18:57:05Z","title":"One Embedder, Any Task: Instruction-Finetuned Text Embeddings","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.09741","snapshot_observed_at":"2026-08-11T05:55:14.784568Z","title":"Smith, Luke Zettlemoyer, and Tao Yu","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.784568Z"},"links":{"cited_paper":"/paper/2212.09741","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:5c45b0ed15778931acd17bcd6535bcc02fea0eca55df2b9ed769ba3dc6883c8c","observation_id":"2632d840-75c1-46a7-80a1-555e5e21e3dd","resolution":{"observed_at":"2026-08-11T05:55:14.784568Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10168","last_updated":"2024-04-03T00:25:39Z","snapshot_observed_at":"2026-08-13T10:31:03.133336Z","submitted_at":"2023-08-20T05:31:03Z","title":"Head-to-Tail: How Knowledgeable are Large Language Models (LLMs)? A.K.A. Will LLMs Replace Knowledge Graphs?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10168","snapshot_observed_at":"2026-08-11T05:55:14.789259Z","title":"Xu, Hanwen Zha, Yue Liu, and Xinhsuai Dong","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.789259Z"},"links":{"cited_paper":"/paper/2308.10168","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:c23061ab6f40935add093e1ecba9deb391dc6bbd24a839c9817b4687205cc92f","observation_id":"550c42f2-7140-49fe-bea9-de6241d3e40c","resolution":{"observed_at":"2026-08-11T05:55:14.789259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.15391","last_updated":"2024-01-27T11:41:48Z","snapshot_observed_at":"2026-08-13T01:06:03.395715Z","submitted_at":"2024-01-27T11:41:48Z","title":"MultiHop-RAG: Benchmarking Retrieval-Augmented Generation for Multi-Hop Queries","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.15391","snapshot_observed_at":"2026-08-11T05:55:14.793714Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.793714Z"},"links":{"cited_paper":"/paper/2401.15391","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:7507aa4045a0520c2e559e70a2c05be77b29b1f78fb0aa06fff4f6f759a9b447","observation_id":"2c5cc6f1-7512-4a28-93d6-cc3f8b5542ca","resolution":{"observed_at":"2026-08-11T05:55:14.793714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00118","last_updated":"2024-10-02T15:22:49Z","snapshot_observed_at":"2026-08-02T16:20:09.773989Z","submitted_at":"2024-07-31T19:13:07Z","title":"Gemma 2: Improving Open Language Models at a Practical Size","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00118","snapshot_observed_at":"2026-08-11T05:55:14.798033Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.798033Z"},"links":{"cited_paper":"/paper/2408.00118","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:0faf3327a16b11e9a0a333e3fac9926bef1d8f3df63e8e163a8f6867cf9c1659","observation_id":"bd6bcbcd-51f6-4416-ba5d-e533e5e6260f","resolution":{"observed_at":"2026-08-11T05:55:14.798033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.664366Z","title":"Trivedi, Niranjan Balasubramanian, Tushar Khot, and Ashish Sabharwal","venue":null,"work_id":"66825b40-b762-4fe2-b4c6-cd5e97eb1855","year":2021},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.802177Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:7e5068a03c01c20ac52ef9d484e55408c291b9e5acc158e596c5912a090f9503","observation_id":"f6b3bcba-d562-4e67-904f-6eae0c253f65","resolution":{"observed_at":"2026-08-11T05:55:15.668466Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.805847Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.805847Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:ee117c0aedd8a9e67f603dedb12f94ad83a7719d88f7aa816fab4c86bf1f9fbb","observation_id":"8f51e6fb-cd0f-4d0a-8aa2-06f12269e604","resolution":{"observed_at":"2026-08-11T05:55:14.805847Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.809422Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.809422Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:c6b7eaf1d939c9e8b26e64d0b06ef5d3a970056133b2f5f8c68d017d5cbe70bb","observation_id":"034e0e7b-aedc-43f0-9c26-b3fce9532a38","resolution":{"observed_at":"2026-08-11T05:55:14.809422Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13552","last_updated":"2023-10-23T05:42:42Z","snapshot_observed_at":"2026-08-14T00:54:45.457434Z","submitted_at":"2023-10-20T14:51:10Z","title":"Self-prompted Chain-of-Thought on Large Language Models for Open-domain Multi-hop Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13552","snapshot_observed_at":"2026-08-11T05:55:14.814055Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.814055Z"},"links":{"cited_paper":"/paper/2310.13552","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:976424fd80ef11ea5a92de2eba95fb3d07ed42c4492033c1d421832be973f0bf","observation_id":"dac56400-d0a1-4af8-9c5c-c8eca28261bd","resolution":{"observed_at":"2026-08-11T05:55:14.814055Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.641636Z","title":null,"venue":null,"work_id":"ad03faf8-d21f-455c-8bf3-7d802ce6f838","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.818452Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:92d0c4eb22c51c8f14f1044d3ee3eb437988a6e5b7c0015b2df3547cde608741","observation_id":"06dc3a3d-0616-4137-9848-230da6870e7c","resolution":{"observed_at":"2026-08-11T05:55:15.645885Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.628187Z","title":null,"venue":null,"work_id":"afa89613-9735-4301-93a4-e6fcd33329f8","year":2022},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.823098Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:acbf433e920d32a0b92fbd5fe7e2118d36a8a80fd0e26ec36e74956af603e55a","observation_id":"c120a452-40e0-4e3c-b2a9-211633a44dbf","resolution":{"observed_at":"2026-08-11T05:55:15.632440Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.827330Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.827330Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:7a637509a126069bd890598ee08e452bb104b0c12a54bc4b2945765e5241c852","observation_id":"d376a72f-5f67-4854-83dc-c55de4e41a2f","resolution":{"observed_at":"2026-08-11T05:55:14.827330Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-08-13T07:04:41.220509Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.11903","snapshot_observed_at":"2026-08-11T05:55:14.831416Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.831416Z"},"links":{"cited_paper":"/paper/2201.11903","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:96354226589170ace99898dffd82a9c131817741f9313d7b7cf164ec2f9c88e0","observation_id":"96b9467a-0c67-42a6-bbeb-f1efb39208df","resolution":{"observed_at":"2026-08-11T05:55:14.831416Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.11136","last_updated":"2024-09-17T12:42:55Z","snapshot_observed_at":"2026-08-12T22:43:16.274443Z","submitted_at":"2024-09-17T12:42:55Z","title":"Promptriever: Instruction-Trained Retrievers Can Be Prompted Like Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.11136","snapshot_observed_at":"2026-08-11T05:55:14.836046Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.836046Z"},"links":{"cited_paper":"/paper/2409.11136","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:581869f2d736a716f8015d48085aa150cfd7afee9c3e7260bfa61348f7698c39","observation_id":"b0cfbe42-b200-4ed5-a25b-255c5e4a5cf4","resolution":{"observed_at":"2026-08-11T05:55:14.836046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.13534","last_updated":"2023-12-08T16:45:22Z","snapshot_observed_at":"2026-08-13T05:19:31.115953Z","submitted_at":"2023-11-22T17:14:54Z","title":"LM-Cocktail: Resilient Tuning of Language Models via Model Merging","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.13534","snapshot_observed_at":"2026-08-11T05:55:14.841141Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.841141Z"},"links":{"cited_paper":"/paper/2311.13534","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:97eaeb8bf0fe8a0b5ec4b78ff33ee825a3da9e5bd6b03c7cc24417ad3e2d46a0","observation_id":"b86cd3b5-149c-4dd1-aa52-7d208e5707f3","resolution":{"observed_at":"2026-08-11T05:55:14.841141Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2025.acl-long.1540","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.932681Z","title":null,"venue":null,"work_id":"2219e68e-4048-4aef-9c15-304e5af5bd9f","year":2025},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.845683Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:8adebee73177e98deaba0d43dfb390f22c4b9cdadb756eb1ee283c853e21d96e","observation_id":"2e64a726-be0e-4ea9-8590-bacd53e343a1","resolution":{"observed_at":"2026-08-11T05:55:14.939744Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.12756","last_updated":"2021-02-19T22:15:03Z","snapshot_observed_at":"2026-08-08T13:42:50.719443Z","submitted_at":"2020-09-27T06:12:29Z","title":"Answering Complex Open-Domain Questions with Multi-Hop Dense Retrieval","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.12756","snapshot_observed_at":"2026-08-11T05:55:14.850001Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.850001Z"},"links":{"cited_paper":"/paper/2009.12756","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:38734d242abc459e4b05e152ebf4e471aff95b539706218c23ac2e1b7627b53a","observation_id":"59b3c53c-3d12-4743-acb5-a38b61c4ffb7","resolution":{"observed_at":"2026-08-11T05:55:14.850001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.604735Z","title":null,"venue":null,"work_id":"c47480df-7e69-4f0f-84a5-a01072a05ad2","year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.854280Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:92c598a70cdd15d017951354564716928bb0104068824491dc84dfd7daeefbaa","observation_id":"79a03d29-3262-41a0-b30c-5d87dfe99e0d","resolution":{"observed_at":"2026-08-11T05:55:15.609094Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.591526Z","title":"Cohen, Ruslan Salakhutdinov, and Christopher D","venue":null,"work_id":"8a3f9204-2f48-4a3a-91b1-391f6d0d8243","year":2018},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.858421Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:e83a9a4c181e9be5c622e58b709d2a6a1feb1b4f66c9c53352ef12f40db6310b","observation_id":"69b744c9-ad5f-4f5b-ae1d-60f4bf7c28c2","resolution":{"observed_at":"2026-08-11T05:55:15.595699Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.578287Z","title":null,"venue":null,"work_id":"5bf04f77-7a2d-49bb-b9a8-e40ed5d0f2f3","year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.862741Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:05fd81c40fbf8c84f6b073496077f9d3dcff32fcce81a8a6959396d4a69f2a4d","observation_id":"07a26606-07bf-4f5e-a3d6-63e8dd4602e7","resolution":{"observed_at":"2026-08-11T05:55:15.582715Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.564734Z","title":"Yu, Wenhao Yu, Chenguang Zhu, Zaitang Li, Zhiting Hu, Qingyun Wang, Heng Ji, and Meng Jiang","venue":null,"work_id":"452ca396-cb1b-4a6a-b84f-e0d640dbf5bf","year":2020},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.867077Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:a2b2924d080cf11c82ddd1ace3d8ca2d43fadb1b41b1ef86e16ec9c832ce5e5e","observation_id":"cb29f5ca-c30b-4377-8b3e-2a55bd819688","resolution":{"observed_at":"2026-08-11T05:55:15.569316Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.16457","last_updated":"2024-06-05T05:23:21Z","snapshot_observed_at":"2026-08-13T23:22:35.889667Z","submitted_at":"2024-02-26T09:59:04Z","title":"RetrievalQA: Assessing Adaptive Retrieval-Augmented Generation for Short-form Open-Domain Question Answering","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.16457","snapshot_observed_at":"2026-08-11T05:55:14.871403Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.871403Z"},"links":{"cited_paper":"/paper/2402.16457","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:f683fe401fb8f7d9264e7131ecf37e6967adbf9d834eb32a10d858e7da634abb","observation_id":"0004aca3-7805-4f9e-966e-2ea550cf9e33","resolution":{"observed_at":"2026-08-11T05:55:14.871403Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.03268","last_updated":"2023-05-05T03:49:14Z","snapshot_observed_at":"2026-08-13T11:48:53.553044Z","submitted_at":"2023-05-05T03:49:14Z","title":"Verify-and-Edit: A Knowledge-Enhanced Chain-of-Thought Framework","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.03268","snapshot_observed_at":"2026-08-11T05:55:14.875132Z","title":"Joty, Chengwei Qin, and Lidong Bing","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.875132Z"},"links":{"cited_paper":"/paper/2305.03268","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:534961b627778cb44f5a200dd51f440a5b19cbd13b967ef29229c3668bcd1e30","observation_id":"e7161233-33ae-4eeb-af2f-c2a0691163f1","resolution":{"observed_at":"2026-08-11T05:55:14.875132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:15.550798Z","title":null,"venue":null,"work_id":"24e58d12-c7d7-41b2-83e3-62b34f8759b9","year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.878899Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:581d0f8e52ee3492b9286f6e7c64227df3db2459670bb553703732275702277e","observation_id":"45b8792c-1047-4b93-9b26-99ed99da4b89","resolution":{"observed_at":"2026-08-11T05:55:15.555225Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16178","last_updated":"2024-05-25T11:10:04Z","snapshot_observed_at":"2026-08-13T14:50:10.382772Z","submitted_at":"2024-05-25T11:10:04Z","title":"Accelerating Inference of Retrieval-Augmented Generation via Sparse Context Selection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16178","snapshot_observed_at":"2026-08-11T05:55:14.883477Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.883477Z"},"links":{"cited_paper":"/paper/2405.16178","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:9e9ea2ab83e37e02b6e18d8f4226f14f53c7ca0dbc6359daf58991d00fb51b49","observation_id":"4954c2a4-3254-45a3-9c22-9d1d0b6196c0","resolution":{"observed_at":"2026-08-11T05:55:14.883477Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.888105Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.888105Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:9d2afab0307b74596ca518ce36abb78000fb4ca5391a95d5e835cf4e2f439149","observation_id":"d3217fc6-19cf-4db7-b636-3b5574086b6c","resolution":{"observed_at":"2026-08-11T05:55:14.888105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.04259","last_updated":"2024-09-26T11:42:35Z","snapshot_observed_at":"2026-08-12T23:06:46.583561Z","submitted_at":"2024-08-08T06:57:49Z","title":"EfficientRAG: Efficient Retriever for Multi-Hop Question Answering","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.04259","snapshot_observed_at":"2026-08-11T05:55:14.892561Z","title":"Rajmohan, Dongmei Zhang, and Qi Zhang","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.892561Z"},"links":{"cited_paper":"/paper/2408.04259","citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:e59a7e0c8f9d3252a599ae62e5e077df5dc1bcfece9cc2ed9e4f744af7eb4d81","observation_id":"c182206e-d2c1-45f7-a7a2-daa558bbfe16","resolution":{"observed_at":"2026-08-11T05:55:14.892561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.897152Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.897152Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:41e71a4f9ab9002245df001e429eb77e8f37d3a27c4db097c39d1d1f2fddfaea","observation_id":"0b4d1837-57d9-4514-a222-7d94ae615d56","resolution":{"observed_at":"2026-08-11T05:55:14.897152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T05:55:14.902112Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge","version":3},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-11T05:55:14.902112Z"},"links":{"citing_paper":"/paper/2412.17032"},"observation_digest":"sha256:0f85de09adc25bb45443b6a571dd97b87f56f8a938f292935163933c0a959b13","observation_id":"4af56952-683f-4c5e-a41d-42f68a081186","resolution":{"observed_at":"2026-08-11T05:55:14.902112Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.17032","last_updated":"2025-08-22T02:06:08Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-13T01:06:43.034748Z","submitted_at":"2024-12-22T14:17:12Z","title":"MINTQA: A Multi-Hop Question Answering Benchmark for Evaluating LLMs on New and Tail Knowledge"},"reference_resolution":{"displayed":57,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":50,"verified_exact":3,"verified_fuzzy":4},"total_outbound_references":57},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 57 of 57 outbound references and 4 inbound Pith citation observations for arXiv:2412.17032."}