{"as_of":"2026-08-18T01:48:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:4d8f5358c36cdbe07a5d8a4d843b66668dcc56ee9135376c33948056a786db8a","coverage":[{"denominator":47,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":47,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T22:23:09.158810Z","state":"measured"},{"denominator":47,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":47,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.21783/citation-record","integrity":"/paper/2506.21783/integrity","json":"/paper/2506.21783/citation-record.json","paper":"/paper/2506.21783"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2205.12665","last_updated":"2023-05-29T06:16:41Z","snapshot_observed_at":"2026-08-17T21:54:33.110402Z","submitted_at":"2022-05-25T11:21:30Z","title":"QAMPARI: An Open-domain Question Answering Benchmark for Questions with Many Answers from Multiple Paragraphs","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12665","snapshot_observed_at":"2026-08-06T22:23:04.374692Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:04.374692Z"},"links":{"cited_paper":"/paper/2205.12665","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:0fec748b0dfb43685a5c176a67881088005f05594494dd64417f24248d1f164c","observation_id":"3a0efc02-e767-49ed-9fcf-6f280db2b926","resolution":{"observed_at":"2026-08-06T22:23:04.374692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:14.032763Z","title":null,"venue":null,"work_id":"cf3f4a37-2b2d-4697-9588-a6656ff843c7","year":2020},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:05.561369Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:1bc26404fe7a6987e87bfd9d26d4a50a675ae6d5c7c322a5883cdd1530813960","observation_id":"b3ba15a7-7a5f-4542-9ead-7c2b6eae4446","resolution":{"observed_at":"2026-08-06T22:23:14.120962Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1145/2619088","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"Jorge, and Adam Jatowt","venue":"ACM Computing Surveys","work_id":"e3c2640e-f79e-4146-8418-6d8cde6b91ed","year":2014},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:05.611183Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:e0ace3ed28dca49294c37cbca1ec9de41b345614e83e72ad46b31012897cca7c","observation_id":"fa96eb93-49db-46f4-aea5-9d4860d68f15","resolution":{"observed_at":"2026-08-06T22:23:10.295044Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2001.04828","last_updated":"2020-01-10T01:43:54Z","snapshot_observed_at":"2026-08-14T21:34:24.909027Z","submitted_at":"2020-01-10T01:43:54Z","title":"TableQnA: Answering List Intent Queries With Web Tables","version":1},"cited_work":{"arxiv_id":"2001.04828","doi":null,"metadata_source":"pith","pith_arxiv_id":"2001.04828","snapshot_observed_at":"2026-08-06T22:23:12.504466Z","title":"TableQnA: Answering List Intent Queries With Web Tables","venue":"cs.IR","work_id":"a9f149f9-a733-4fe2-869b-cdbef6f214fa","year":2020},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:05.871209Z"},"links":{"cited_paper":"/paper/2001.04828","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:bac0c72b69c1c24328db1715e6ee9b3bc3d54df758e353e83d63bda911c7739d","observation_id":"8345b6ac-4fa9-4c66-8c6a-2c9e7bf574bf","resolution":{"observed_at":"2026-08-06T22:23:12.579978Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"8433.2023","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:12.305747Z","title":null,"venue":null,"work_id":"1e5b1083-a1a0-4ab1-ada1-fe98bad1a4f0","year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.055738Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:7c45774ed3daea8729e74a77ff1cb4ed5317eb9ddd2254f7a0eb8c14924a743b","observation_id":"cf6909dd-43b9-432a-8c63-ef15d1566065","resolution":{"observed_at":"2026-08-06T22:23:12.406175Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2108.06314","last_updated":"2021-10-25T01:10:38Z","snapshot_observed_at":"2026-08-16T18:03:21.632822Z","submitted_at":"2021-08-13T16:42:25Z","title":"A Dataset for Answering Time-Sensitive Questions","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2108.06314","snapshot_observed_at":"2026-08-06T22:23:06.148897Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.148897Z"},"links":{"cited_paper":"/paper/2108.06314","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:bfe6aeee7f2ee5ff7361b6b4a663408a0a0acb87181aa86e1f270f1906ccc230","observation_id":"c27269b5-1a6a-4770-8e18-91cd0f58ec82","resolution":{"observed_at":"2026-08-06T22:23:06.148897Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:13.862323Z","title":null,"venue":null,"work_id":"efd9753d-6d4e-410e-972c-f3203edcee50","year":null},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.297704Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:18bd69a9395c7b044cac18a5190eea46cb80f5e59b127aea4e05451dc422278f","observation_id":"b3d9f569-4c75-4803-a254-140c7cde4077","resolution":{"observed_at":"2026-08-06T22:23:13.929477Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:06.403423Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.403423Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:c9f9409411aeb88456778c1b3c0bce4b2930149f36e30bb7a9f3ad08555d871e","observation_id":"f7a29210-4235-4e06-a178-9d27c11a3a5c","resolution":{"observed_at":"2026-08-06T22:23:06.403423Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:13.686246Z","title":"Cole, Julian Martin Eisenschlos, Daniel Gillick, Jacob Eisenstein, and William W","venue":null,"work_id":"9d70b0f8-f48e-46b4-a16f-f82ff1562112","year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.467023Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:2e31900a1cc2f9fcc91c89e1cf28da15579a28b30e28d2a911a4a21f3505da5a","observation_id":"ab0fadb3-6915-4b0c-8e19-7b6c21d43dc0","resolution":{"observed_at":"2026-08-06T22:23:13.778622Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1804.04526","last_updated":"2018-04-12T14:12:48Z","snapshot_observed_at":"2026-08-14T19:26:36.205395Z","submitted_at":"2018-04-12T14:12:48Z","title":"EventKG: A Multilingual Event-Centric Temporal Knowledge Graph","version":1},"cited_work":{"arxiv_id":"1804.04526","doi":null,"metadata_source":"pith","pith_arxiv_id":"1804.04526","snapshot_observed_at":"2026-08-06T22:23:12.012730Z","title":"EventKG: A Multilingual Event-Centric Temporal Knowledge Graph","venue":"cs.CL","work_id":"bad180a3-71f8-4956-b5e4-af39ed1a9d55","year":2018},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.536218Z"},"links":{"cited_paper":"/paper/1804.04526","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:a8b464215bf5395c5cc1d5e54f092df30124a05394150f6a34cdb817abae0ac7","observation_id":"928b88b2-8585-4b72-b1a1-ccf867f2cbc2","resolution":{"observed_at":"2026-08-06T22:23:12.079214Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"0690.14607","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:11.851240Z","title":"Green, Alice K","venue":null,"work_id":"5853d522-ae9b-46bb-a28d-571ce2a4b77e","year":1961},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.587959Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:840d0209566d8fe66fa202391b9f300fce5db3aea23c367a1a4ba539c02e8d94","observation_id":"d9892656-f368-49cf-90e2-5750fff698c3","resolution":{"observed_at":"2026-08-06T22:23:11.925666Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:06.680084Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.680084Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:53d1e25c58b9214059ba4199604e6a00c22cbe81688aa65135febed8425b000b","observation_id":"c470392f-e806-4520-b8d0-eac31003d9ba","resolution":{"observed_at":"2026-08-06T22:23:06.680084Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:06.820451Z","title":"Suchanek, Klaus Berberich, Edwin Lewis-Kelham, Gerard de Melo, and Gerhard Weikum","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.820451Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:6f06e31a6bcbdcf509d9ca265832ebb12291ab0d8b876593231c06645e53314d","observation_id":"90bad459-b73e-40d5-aa9e-80d4e18b7681","resolution":{"observed_at":"2026-08-06T22:23:06.820451Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:06.906128Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.906128Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:6e437578d49d4ff7dc2569c99059dc6a74d298abcc4ef12a6907a787063a76c7","observation_id":"bb30cd6c-b36a-401b-8c8b-ac7e8ccf231d","resolution":{"observed_at":"2026-08-06T22:23:06.906128Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:06.950255Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.950255Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:c8137ff17a7f728e0aaec0117db1da318a81e2327c232c496e95ca74265dcc9b","observation_id":"9122243d-35f3-4a79-89fb-b2e4d1744337","resolution":{"observed_at":"2026-08-06T22:23:06.950255Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:07.047940Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.047940Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:021f60cff03c179b80dddf97dfbe6a68976b9572e959bf729bdf43ea25d9c269","observation_id":"effe2f33-5dc0-4f02-9115-660913202d97","resolution":{"observed_at":"2026-08-06T22:23:07.047940Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"7948.25790","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:11.280596Z","title":null,"venue":null,"work_id":"fcc057fc-94fd-42c0-9962-38dc1ec805ae","year":2014},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.137674Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:762f51ef73271dc3be957ec533b30dfebc2b9a00bfef880f708aeea67e338523","observation_id":"977c243c-8b94-42d2-a1b9-1bd7095710e0","resolution":{"observed_at":"2026-08-06T22:23:11.380446Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:07.236731Z","title":null,"venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.236731Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:5adb8f51be12794a5472352ce3528b3debcc332386cc4d1b71ce5c5f4ebbe741","observation_id":"647821a5-589f-4cc5-b028-1631a4c025f8","resolution":{"observed_at":"2026-08-06T22:23:07.236731Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:07.294674Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.294674Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:6e7321baf53b5c2ddaf0440d44731c311a451d0d569e68c3a84841af135ed83e","observation_id":"b910ddc5-7b90-4118-b052-35d4e2bb5383","resolution":{"observed_at":"2026-08-06T22:23:07.294674Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:07.384743Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.384743Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:717f97c5df855d5514f5f6d712aaa30c8dc0925f799dd7d4ef5eee6b0e14d882","observation_id":"665e3ea5-d2f6-4daa-9bb3-5770be3f6a06","resolution":{"observed_at":"2026-08-06T22:23:07.384743Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.13332","last_updated":"2024-02-28T07:45:37Z","snapshot_observed_at":"2026-08-16T16:43:03.622616Z","submitted_at":"2022-07-27T07:26:01Z","title":"RealTime QA: What's the Answer Right Now?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.13332","snapshot_observed_at":"2026-08-06T22:23:07.545920Z","title":"Radev, Noah A","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.545920Z"},"links":{"cited_paper":"/paper/2207.13332","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:d43312a51cc81d2903179559d4f0e3f3257659a6d7dafc32a8b624ed9894dc99","observation_id":"d0f254f8-7446-4502-8a82-47efda87ed73","resolution":{"observed_at":"2026-08-06T22:23:07.545920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:13.516614Z","title":null,"venue":null,"work_id":"3c2f388a-924c-422a-a93f-b4e69ddd4005","year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.625786Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:be71f86f52952f3e1221d8a3240e242804d28577051d60ae26329ba665778378","observation_id":"c871ed96-9558-43fe-b452-08b0dee3b1aa","resolution":{"observed_at":"2026-08-06T22:23:13.595082Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:13.294682Z","title":null,"venue":null,"work_id":"1ccbff06-a41e-4c43-9d17-525f2e0cd9bb","year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.679804Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:4d071889eb8cd7203b0641bb3eb4975afd6d41183452c304a81c2df6c3bc41e7","observation_id":"c11ed72f-86f6-4b4e-ba3a-7e92c2b2c00d","resolution":{"observed_at":"2026-08-06T22:23:13.370000Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1609/aaai.v37i11.26529","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":null,"venue":"Proceedings of the AAAI Conference on Artificial Intelligence","work_id":"ab8cf2d7-2489-45fe-a9f3-38d47b3d5440","year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.744746Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:1d0417824bc0429158e9373f2d32d2d9a87ec86b4330f2340d31f45ff441eafb","observation_id":"a40664e9-9593-46ad-a8fb-b1645d79d82e","resolution":{"observed_at":"2026-08-06T22:23:13.211975Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.07396","last_updated":"2024-11-11T22:06:51Z","snapshot_observed_at":"2026-08-16T13:01:00.229516Z","submitted_at":"2024-11-11T22:06:51Z","title":"Toward Optimal Search and Retrieval for RAG","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.07396","snapshot_observed_at":"2026-08-06T22:23:07.807419Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.807419Z"},"links":{"cited_paper":"/paper/2411.07396","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:c5573a97d11f4cafe61cb6e25c234949000d2dced3f0ebc30951ff2aa4471b99","observation_id":"42aa87fb-2ada-4525-9ed4-3cd332ceb0ab","resolution":{"observed_at":"2026-08-06T22:23:07.807419Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:07.873260Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.873260Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:0db656c0d831948fb95b300c6deb2d48b6a19269cadf6e0861283c4947c431eb","observation_id":"c7c1943c-058f-4c38-a98c-75ab910d71fa","resolution":{"observed_at":"2026-08-06T22:23:07.873260Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.12297","last_updated":"2023-02-23T19:24:55Z","snapshot_observed_at":"2026-08-16T15:53:07.936534Z","submitted_at":"2023-02-23T19:24:55Z","title":"Dynamic Benchmarking of Masked Language Models on Temporal Concept Drift with Multiple Views","version":1},"cited_work":{"arxiv_id":"2302.12297","doi":null,"metadata_source":"pith","pith_arxiv_id":"2302.12297","snapshot_observed_at":"2026-08-06T22:23:10.861420Z","title":"Dynamic Benchmarking of Masked Language Models on Temporal Concept Drift with Multiple Views","venue":"cs.CL","work_id":"2df22484-cc47-48dc-bc00-7b1c8dfab6f8","year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:07.992648Z"},"links":{"cited_paper":"/paper/2302.12297","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:c9e454acb1a191adad28a75597d9c640f45b495d0458c6ba00752294e9c8f24e","observation_id":"6911bb7b-f755-41dc-92b2-cce2dbcbbf9d","resolution":{"observed_at":"2026-08-06T22:23:10.914947Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.09140","last_updated":"2024-05-31T14:28:40Z","snapshot_observed_at":"2026-08-17T03:01:33.957468Z","submitted_at":"2022-04-19T21:55:18Z","title":"Multi-hop Question Answering","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.09140","snapshot_observed_at":"2026-08-06T22:23:08.041771Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.041771Z"},"links":{"cited_paper":"/paper/2204.09140","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:f9b3ef6cb88b71259a27d960c0813789fcdcb7a6e304b8f010ab1c28f541920c","observation_id":"4aaa981b-091d-4d05-8058-b7983625b549","resolution":{"observed_at":"2026-08-06T22:23:08.041771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.05785","last_updated":"2021-12-10T23:59:14Z","snapshot_observed_at":"2026-08-16T17:35:28.835016Z","submitted_at":"2021-12-10T23:59:14Z","title":"TempoQR: Temporal Question Reasoning over Knowledge Graphs","version":1},"cited_work":{"arxiv_id":"2112.05785","doi":null,"metadata_source":"pith","pith_arxiv_id":"2112.05785","snapshot_observed_at":"2026-08-06T22:23:10.679743Z","title":"TempoQR: Temporal Question Reasoning over Knowledge Graphs","venue":"cs.CL","work_id":"f0737770-ad9d-41ff-ae47-a8a5d50d52e7","year":2021},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.103377Z"},"links":{"cited_paper":"/paper/2112.05785","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:1e56006c092cd7d3e101b3a746f4738732279eb48b8b54545c46ffad7d273200","observation_id":"9da521f5-07c8-41e1-aff7-ed6657825655","resolution":{"observed_at":"2026-08-06T22:23:10.725257Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-06T22:23:08.161871Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.161871Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:be0b9a89bc93d4acd038ee9180d9de5cd44eb0078b90738b985b7edcc484468b","observation_id":"c2a1d91f-051a-4e92-8c9e-6105e9a3d355","resolution":{"observed_at":"2026-08-06T22:23:08.161871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1606.05250","last_updated":"2016-10-11T02:42:36Z","snapshot_observed_at":"2026-08-17T14:32:22.468812Z","submitted_at":"2016-06-16T16:36:00Z","title":"SQuAD: 100,000+ Questions for Machine Comprehension of Text","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.05250","snapshot_observed_at":"2026-08-06T22:23:08.216974Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.216974Z"},"links":{"cited_paper":"/paper/1606.05250","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:5cc8aaed7d0430dc958e32d481ba3c5c17951319c8bdd7b6fe23141fbeeb6f10","observation_id":"988294b0-b3a9-4de8-88a5-a6ed2ed30b36","resolution":{"observed_at":"2026-08-06T22:23:08.216974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:08.282782Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.282782Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:f3c8becd629b3cb6b6713636f9385760567407f26e1ea8a4dcef9b7b335d6d5c","observation_id":"e6c1511b-5139-4939-8b4a-09cb2bcbd26b","resolution":{"observed_at":"2026-08-06T22:23:08.282782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:08.328916Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.328916Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:24cca1929943a4e69a3a84c247ddd9c8c898b69685e4782d24301a68ded4cbf4","observation_id":"a6701a25-a601-4b7a-92c8-ddd525834bf6","resolution":{"observed_at":"2026-08-06T22:23:08.328916Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:12.912125Z","title":null,"venue":null,"work_id":"7ed6d02a-fca1-4d60-aace-86a808cd1bbd","year":2021},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.393176Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:a36ce6b87c532940af0ea3fc681833035a98b18eda3d3a38f591f8a0bc5a8070","observation_id":"f5676930-7256-4063-a39e-deb212180a2c","resolution":{"observed_at":"2026-08-06T22:23:13.019342Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:08.444683Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.444683Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:a50724207ac0522a4ba84fc84c14a6dbcd6101f823ba6b44e5296bb3afd92f64","observation_id":"5f1201af-67f7-418e-b6f7-4842fc9e92b8","resolution":{"observed_at":"2026-08-06T22:23:08.444683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.09617","last_updated":"2023-05-16T17:11:29Z","snapshot_observed_at":"2026-08-16T14:30:19.311365Z","submitted_at":"2023-05-16T17:11:29Z","title":"Towards Expert-Level Medical Question Answering with Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.09617","snapshot_observed_at":"2026-08-06T22:23:08.526543Z","title":"Sara Mahdavi, Joelle K","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.526543Z"},"links":{"cited_paper":"/paper/2305.09617","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:f9efb89d758806c3d44e8d1a4a057e5e18deda1ba7c0ed59343f0131b5c048d2","observation_id":"c1ecfb16-997c-4e47-a17a-62abb0dc90ad","resolution":{"observed_at":"2026-08-06T22:23:08.526543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:08.609107Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.609107Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:a6a52e5ab8b3c526768e61b273dd8c66c1b0638008530b7077580078f45b5c5e","observation_id":"1a523f3c-1d79-4f06-bb7b-ed7fc809766f","resolution":{"observed_at":"2026-08-06T22:23:08.609107Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18796","last_updated":"2024-05-01T15:37:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-04-29T15:33:23Z","title":"Replacing Judges with Juries: Evaluating LLM Generations with a Panel of Diverse Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18796","snapshot_observed_at":"2026-08-06T22:23:08.743737Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.743737Z"},"links":{"cited_paper":"/paper/2404.18796","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:e79509ec7b9e85486bf02f3a344c026dcd82ea4a826d1a443c645e171008c63b","observation_id":"3879feb9-1cdd-4a58-b215-225226892bfa","resolution":{"observed_at":"2026-08-06T22:23:08.743737Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2109.03438","last_updated":"2022-02-22T04:51:20Z","snapshot_observed_at":"2026-08-16T17:58:10.518343Z","submitted_at":"2021-09-08T05:21:51Z","title":"ArchivalQA: A Large-scale Benchmark Dataset for Open Domain Question Answering over Historical News Collections","version":4},"cited_work":{"arxiv_id":"2109.03438","doi":null,"metadata_source":"pith","pith_arxiv_id":"2109.03438","snapshot_observed_at":"2026-08-06T22:23:10.420794Z","title":"ArchivalQA: A Large-scale Benchmark Dataset for Open Domain Question Answering over Historical News Collections","venue":"cs.CL","work_id":"e7fdbe0b-42fe-4b3d-aedd-818f33cbf9a7","year":2021},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.857310Z"},"links":{"cited_paper":"/paper/2109.03438","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:2efd3973752e3f9dee00fedcad34fe5de4473eae0df0a38678dc696faf5f0c59","observation_id":"74ba833e-ba53-4f45-876e-c757b26c1e89","resolution":{"observed_at":"2026-08-06T22:23:10.538312Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1093/bioinformatics/btac397","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":null,"venue":"Bioinformatics","work_id":"40d59324-6d68-422c-905d-82e8531bc3e9","year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:08.973656Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:2409904abb1b8838b796eb2f8c8e144a61d96360c50d5b290b6db3caf428c91e","observation_id":"b5de69e0-e00a-45ad-99a7-03196e1aeac0","resolution":{"observed_at":"2026-08-06T22:23:09.859003Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.00435","last_updated":"2023-06-01T08:22:21Z","snapshot_observed_at":"2026-08-16T15:27:37.572329Z","submitted_at":"2023-06-01T08:22:21Z","title":"How Many Answers Should I Give? An Empirical Study of Multi-Answer Reading Comprehension","version":1},"cited_work":{"arxiv_id":"2306.00435","doi":"10.48550/arxiv.2306.00435","metadata_source":"pith","pith_arxiv_id":"2306.00435","snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"How Many Answers Should I Give? An Empirical Study of Multi-Answer Reading Comprehension","venue":"cs.CL","work_id":"a9f5a12e-8d31-400e-aa40-8ed79f79c336","year":2023},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:09.002189Z"},"links":{"cited_paper":"/paper/2306.00435","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:3d6a17235df8e613c945c323e9834ca13ea1a68a9301bd29a32ed58a903ab3d8","observation_id":"de805122-6c09-4514-9169-dfe3dba66cd0","resolution":{"observed_at":"2026-08-06T22:23:09.642121Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-acl.696","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":null,"venue":null,"work_id":"470d4b05-0551-4dce-b7de-b24fb2d42f55","year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:09.050234Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:458fe3e9201a436ecf5154c1318ad06346b0d7e94d3084d654d93a35b1ef4959","observation_id":"a5e2c379-9126-4630-a075-52303ae80a5a","resolution":{"observed_at":"2026-08-06T22:23:09.451625Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:12.674659Z","title":null,"venue":null,"work_id":"602c36ae-4f70-413d-91f3-85ba45c921de","year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:09.095678Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:2fec81eec38370e0b28c14b0b712a1f40d9805cff4f8215042d34dd46e964684","observation_id":"2e84cd8f-b864-475e-a305-8fc72ca84a8c","resolution":{"observed_at":"2026-08-06T22:23:12.785804Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.14353","last_updated":"2022-11-15T17:30:07Z","snapshot_observed_at":"2026-08-16T16:21:09.342150Z","submitted_at":"2022-10-25T21:39:36Z","title":"RoMQA: A Benchmark for Robust, Multi-evidence, Multi-answer Question Answering","version":2},"cited_work":{"arxiv_id":"2210.14353","doi":"10.48550/arxiv.2210.14353","metadata_source":"pith","pith_arxiv_id":"2210.14353","snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"RoMQA: A Benchmark for Robust, Multi-evidence, Multi-answer Question Answering","venue":"cs.CL","work_id":"e164a0e9-5806-4ca4-bde9-c09fa2c6f1cc","year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:09.158810Z"},"links":{"cited_paper":"/paper/2210.14353","citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:92d8f64144815be57d35a1193c58f7b0bf782def4d0ba2948df9aa3804188855","observation_id":"35ff5a90-e615-44a0-be8d-1b1164a892ac","resolution":{"observed_at":"2026-08-06T22:23:09.271555Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:06.518752Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":273,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.518752Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:95f6dd23528e8d9268a07704d9d24ff695b46edb0353ce703c5bb2a0f297312f","observation_id":"af91475f-08fa-42f3-a49e-32ac71b518cf","resolution":{"observed_at":"2026-08-06T22:23:06.518752Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.acl-long.87","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":null,"venue":null,"work_id":"ee9e974c-9892-40ba-88bc-5ec7c4c8ce10","year":2024},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":1606,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:09.120557Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:bf1bbfaf27a6944157e0d78a5825546b2029dfcc3d28f6ee53db90b5f10d011b","observation_id":"baa7bb6e-20fb-4441-b7fa-d0efee54f0f6","resolution":{"observed_at":"2026-08-06T22:23:09.335411Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:23:06.359589Z","title":"Knowledge-Based Systems 251 (2022), 109134","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-06T22:23:06.359589Z"},"links":{"citing_paper":"/paper/2506.21783"},"observation_digest":"sha256:dddb77fb2f1ad6fc71f03d50ed2af80092a049e48d4427c9bdac0d283f225cae","observation_id":"64ae214d-e7f6-4ae8-882b-9854e44c617c","resolution":{"observed_at":"2026-08-06T22:23:06.359589Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.21783","last_updated":"2025-06-26T21:40:58Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-11T22:57:05.533237Z","submitted_at":"2025-06-26T21:40:58Z","title":"Evaluating List Construction and Temporal Understanding capabilities of Large Language Models"},"reference_resolution":{"displayed":47,"state_counts":{"malformed_identifier":4,"metadata_mismatch":3,"parse_uncertain":0,"unresolved":28,"verified_exact":11,"verified_fuzzy":1},"total_outbound_references":47},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 47 of 47 outbound references and 0 inbound Pith citation observations for arXiv:2506.21783."}