{"as_of":"2026-08-13T23:22:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1185bb8015dd473cf432e5b820182e21866af0d0151f964bba7748b21525b8bb","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":11,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":11,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":11,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":11,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T04:20:21.524087Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T11:56:55.163345Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":"2402.14865","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-07-02T11:56:55.163345Z","title":"Dyval 2: Dynamic evaluation of large language models by meta probing agents","venue":null,"work_id":"4ad8c494-31e6-4433-800e-ec38ad95927a","year":2024},"citing_paper":{"arxiv_id":"2404.11584","last_updated":"2024-04-17T17:32:41Z","snapshot_observed_at":"2026-08-04T23:37:27.678120Z","submitted_at":"2024-04-17T17:32:41Z","title":"The Landscape of Emerging AI Agent Architectures for Reasoning, Planning, and Tool Calling: A Survey","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-16T23:16:41.679855Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2404.11584"},"observation_digest":"sha256:e39c6166e6d74a02ec8d036711eb511169d9402f9750c71a8bc4364fd7d0d161","observation_id":"774029bf-e886-44f0-bf9c-2c7fb90d0874","resolution":{"observed_at":"2026-05-16T23:16:41.768894Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":"2402.14865","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-07-02T11:56:55.163345Z","title":"Dyval 2: Dynamic evaluation of large language models by meta probing agents","venue":null,"work_id":"4ad8c494-31e6-4433-800e-ec38ad95927a","year":2024},"citing_paper":{"arxiv_id":"2406.04244","last_updated":"2024-06-06T16:41:39Z","snapshot_observed_at":"2026-07-30T15:43:06.151242Z","submitted_at":"2024-06-06T16:41:39Z","title":"Benchmark Data Contamination of Large Language Models: A Survey","version":1},"reference_index":190,"source":"pdf_text","source_observed_at":"2026-05-22T23:10:40.420241Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2406.04244"},"observation_digest":"sha256:25c37a183923ab88b308e99380754025a0f7137f360e451fbe395fd90e9c16e8","observation_id":"e5667eed-8736-453e-a12b-878b49375758","resolution":{"observed_at":"2026-05-22T23:10:41.165139Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":"2402.14865","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-07-02T11:56:55.163345Z","title":"Dyval 2: Dynamic evaluation of large language models by meta probing agents","venue":null,"work_id":"4ad8c494-31e6-4433-800e-ec38ad95927a","year":2024},"citing_paper":{"arxiv_id":"2406.12708","last_updated":"2026-05-10T21:28:38Z","snapshot_observed_at":"2026-08-12T08:44:03.251322Z","submitted_at":"2024-06-18T15:22:12Z","title":"AgentReview: Exploring Peer Review Dynamics with LLM Agents","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-23T23:38:28.005028Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2406.12708"},"observation_digest":"sha256:92ef3e351773806eedaaceccc55877810235f02f0a58c51bc8a4d01f3e13dd85","observation_id":"46ab20e4-e10a-4530-b7ea-acb9a9d7cd62","resolution":{"observed_at":"2026-05-23T23:38:37.076842Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-08-12T04:20:21.524087Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01526","last_updated":"2024-12-02T14:18:32Z","snapshot_observed_at":"2026-08-13T17:12:49.685789Z","submitted_at":"2024-12-02T14:18:32Z","title":"Addressing Data Leakage in HumanEval Using Combinatorial Test Design","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T04:20:21.524087Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2412.01526"},"observation_digest":"sha256:2d8653a15b3ec12384ec686dfb28160a6e3c96b03b14a341503966bb00ea8ff0","observation_id":"41996491-cd86-4db3-8569-cc51fdad9ac3","resolution":{"observed_at":"2026-08-12T04:20:21.524087Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-08-09T18:09:01.717256Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01683","last_updated":"2025-02-02T06:36:01Z","snapshot_observed_at":"2026-08-12T20:49:55.785732Z","submitted_at":"2025-02-02T06:36:01Z","title":"LLM-Powered Benchmark Factory: Reliable, Generic, and Efficient","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-09T18:09:01.717256Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2502.01683"},"observation_digest":"sha256:67c265b177f933412dc548fc865349e309bf5addeae7aa3c7ee8d4ba107eaa5d","observation_id":"1ee7a73b-7456-4d9b-8939-e8246680a985","resolution":{"observed_at":"2026-08-09T18:09:01.717256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-08-07T10:19:12.218289Z","title":"Dynamic evaluation of large language models by meta probing agents,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02870","last_updated":"2025-06-06T10:50:08Z","snapshot_observed_at":"2026-08-08T01:09:03.836574Z","submitted_at":"2025-06-06T10:50:08Z","title":"Loki's Dance of Illusions: A Comprehensive Survey of Hallucination in Large Language Models","version":1},"reference_index":205,"source":"pdf_text","source_observed_at":"2026-08-07T10:19:12.218289Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2507.02870"},"observation_digest":"sha256:56e7aa4aeb470e6bbcced510576b5e2266271a41e82d6296843d9c8ab5fee55c","observation_id":"c5e53e17-3ea4-4387-8b95-c6b375353b8a","resolution":{"observed_at":"2026-08-07T10:19:12.218289Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-08-05T15:43:12.868181Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.19570","last_updated":"2025-08-27T05:04:07Z","snapshot_observed_at":"2026-08-09T05:32:43.331268Z","submitted_at":"2025-08-27T05:04:07Z","title":"Generative Models for Synthetic Data: Transforming Data Mining in the GenAI Era","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T15:43:12.868181Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2508.19570"},"observation_digest":"sha256:bc41e5b476a5c288fda0ab2b904a817136aa3b3d4389a77f7d5a6b6297c9372e","observation_id":"3d656faf-f8bc-4921-a481-5c2ec615bb6c","resolution":{"observed_at":"2026-08-05T15:43:12.868181Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":"2402.14865","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-07-02T11:56:55.163345Z","title":"Dyval 2: Dynamic evaluation of large language models by meta probing agents","venue":null,"work_id":"4ad8c494-31e6-4433-800e-ec38ad95927a","year":2024},"citing_paper":{"arxiv_id":"2604.08571","last_updated":"2026-07-21T03:29:11Z","snapshot_observed_at":"2026-08-11T03:30:40.958370Z","submitted_at":"2026-03-26T22:19:33Z","title":"Robust Reasoning Benchmark","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-15T00:06:08.545334Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2604.08571"},"observation_digest":"sha256:63adf3d8260d5c5531b72b270d6a0f3b9bbd30b0bd6ceb730419e79474e830aa","observation_id":"697935c8-f6d7-416e-80ab-1f179bb4f393","resolution":{"observed_at":"2026-05-15T00:08:21.747145Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":"2402.14865","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-07-02T11:56:55.163345Z","title":"Dyval 2: Dynamic evaluation of large language models by meta probing agents","venue":null,"work_id":"4ad8c494-31e6-4433-800e-ec38ad95927a","year":2024},"citing_paper":{"arxiv_id":"2604.08571","last_updated":"2026-07-21T03:29:11Z","snapshot_observed_at":"2026-08-11T03:30:40.958370Z","submitted_at":"2026-03-26T22:19:33Z","title":"Robust Reasoning Benchmark","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-22T11:18:51.642405Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2604.08571"},"observation_digest":"sha256:b94118780527060fc172f2227893c36d64a524eefe95fc7f99d68436c808e48c","observation_id":"a005f15b-86d9-4f86-bd55-bf92d3178a95","resolution":{"observed_at":"2026-05-22T11:21:28.774630Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-08-02T17:25:30.971119Z","title":"Dynamic evaluation of large language models by meta probing agents","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.08571","last_updated":"2026-07-21T03:29:11Z","snapshot_observed_at":"2026-08-11T03:30:40.958370Z","submitted_at":"2026-03-26T22:19:33Z","title":"Robust Reasoning Benchmark","version":3},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-02T17:25:30.971119Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2604.08571"},"observation_digest":"sha256:c032cfc7b639c72ad7c66594037d904f49383cc188a9a58e620d2357225c0057","observation_id":"e70dc3cf-cd69-4e8f-a7fc-6d8597385c83","resolution":{"observed_at":"2026-08-02T17:25:30.971119Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents","version":2},"cited_work":{"arxiv_id":"2402.14865","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.14865","snapshot_observed_at":"2026-07-02T11:56:55.163345Z","title":"Dyval 2: Dynamic evaluation of large language models by meta probing agents","venue":null,"work_id":"4ad8c494-31e6-4433-800e-ec38ad95927a","year":2024},"citing_paper":{"arxiv_id":"2606.06546","last_updated":"2026-06-04T07:40:12Z","snapshot_observed_at":"2026-07-31T21:40:27.810095Z","submitted_at":"2026-06-04T07:40:12Z","title":"Elmes*: Automated Construction of Fine-Grained Evaluation Rubrics for Large Language Models in Long-Tail Educational Scenarios","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-06-28T02:49:01.527149Z"},"links":{"cited_paper":"/paper/2402.14865","citing_paper":"/paper/2606.06546"},"observation_digest":"sha256:5eb3744b13566bafcecf601dc48bd29a79e1e23186d30a260c5b5d4489f64f7b","observation_id":"62e7752a-db7f-4d7c-9e70-ab0064f0e99f","resolution":{"observed_at":"2026-07-02T11:56:55.164979Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2402.14865/citation-record","integrity":"/paper/2402.14865/integrity","json":"/paper/2402.14865/citation-record.json","paper":"/paper/2402.14865"},"outbound":[],"paper":{"arxiv_id":"2402.14865","last_updated":"2024-06-07T09:19:45Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-13T04:13:58.496742Z","submitted_at":"2024-02-21T06:46:34Z","title":"Dynamic Evaluation of Large Language Models by Meta Probing Agents"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 11 inbound Pith citation observations for arXiv:2402.14865."}