{"as_of":"2026-08-09T15:21:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b6963222d810978686f2614ec2fbc7a3c604d7d2d1398d04aaa61d59e2af2fd5","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":23,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":23,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":23,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":23,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T14:57:41.362935Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":3,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-09T14:57:41.362935Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01587","last_updated":"2025-02-03T18:20:10Z","snapshot_observed_at":"2026-08-09T14:51:07.507475Z","submitted_at":"2025-02-03T18:20:10Z","title":"Verbalized Bayesian Persuasion","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-09T14:57:41.362935Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2502.01587"},"observation_digest":"sha256:21ce5c0c94c46d0eae627f0b0eec62b332b71590beb6803695416fe9c9317d4e","observation_id":"d822c53b-83bc-4bc7-897f-ac9d279a87f7","resolution":{"observed_at":"2026-08-09T14:57:41.362935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-07T14:06:31.207570Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations.arXiv preprint arXiv:2402.12348,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20047","last_updated":"2025-05-26T14:34:04Z","snapshot_observed_at":"2026-08-07T23:31:54.437843Z","submitted_at":"2025-05-26T14:34:04Z","title":"Grammars of Formal Uncertainty: When to Trust LLMs in Automated Reasoning Tasks","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T14:06:31.207570Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2505.20047"},"observation_digest":"sha256:4e28de9b0f1a67b1864e5dba415cedd51a97d7cbe7239c1295192852d7aa2789","observation_id":"117a72f0-ac6a-4a25-9b3a-786bc51021ee","resolution":{"observed_at":"2026-08-07T14:06:31.207570Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2506.02387","last_updated":"2026-04-13T08:26:12Z","snapshot_observed_at":"2026-08-04T07:16:14.991803Z","submitted_at":"2025-06-03T02:57:38Z","title":"VS-Bench: Evaluating VLMs for Strategic Abilities in Multi-Agent Environments","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-19T11:57:08.314088Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.02387"},"observation_digest":"sha256:b7628dfe2d2da6ae1dbc1e99f36e26678e4a4c7856fa5cd5a2747feaf8aa7c8a","observation_id":"ccdb28e0-7a0d-4687-90bb-27dd9651a7fe","resolution":{"observed_at":"2026-05-19T11:57:16.225848Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-07T05:40:48.118071Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations.arXiv preprint arXiv:2402.12348, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07388","last_updated":"2025-06-09T03:24:01Z","snapshot_observed_at":"2026-08-09T01:50:48.491215Z","submitted_at":"2025-06-09T03:24:01Z","title":"Shapley-Coop: Credit Assignment for Emergent Cooperation in Self-Interested LLM Agents","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T05:40:48.118071Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.07388"},"observation_digest":"sha256:d56462348130ab17e0918111cf428f2f553f23d304c3cd3cda21b05b19f11375","observation_id":"10149eab-6e4a-41bf-af76-6a916325d0b8","resolution":{"observed_at":"2026-08-07T05:40:48.118071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-07T04:34:38.244643Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evalua- tions,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.10264","last_updated":"2025-06-12T01:16:34Z","snapshot_observed_at":"2026-08-08T15:38:22.407812Z","submitted_at":"2025-06-12T01:16:34Z","title":"WGSR-Bench: Wargame-based Game-theoretic Strategic Reasoning Benchmark for Large Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T04:34:38.244643Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.10264"},"observation_digest":"sha256:72e12c990f450d217c2b438ccbe103c676879513d7c13fea61c271dade4ab81b","observation_id":"d37a19b2-dfde-4166-84ee-8fc49b616ab4","resolution":{"observed_at":"2026-08-07T04:34:38.244643Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-07T06:00:19.537892Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.11102","last_updated":"2025-06-06T17:52:18Z","snapshot_observed_at":"2026-08-07T05:54:56.593167Z","submitted_at":"2025-06-06T17:52:18Z","title":"Evolutionary Perspectives on the Evaluation of LLM-Based AI Agents: A Comprehensive Survey","version":1},"reference_index":114,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:19.537892Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.11102"},"observation_digest":"sha256:f6b75bb58b90484febcd5f2fbe585adadebea2b7f98eb5667dff041b14449ffe","observation_id":"3ecc2b31-fb22-4aec-984e-ffa1694e5e8a","resolution":{"observed_at":"2026-08-07T06:00:19.537892Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-07T01:06:45.860130Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12012","last_updated":"2025-06-13T17:59:10Z","snapshot_observed_at":"2026-08-08T19:07:01.916168Z","submitted_at":"2025-06-13T17:59:10Z","title":"Tracing LLM Reasoning Processes with Strategic Games: A Framework for Planning, Revision, and Resource-Constrained Decision Making","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T01:06:45.860130Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.12012"},"observation_digest":"sha256:69b605a3d7419ef02c6845fd8da50154a817a26098aba56023a9531c25ca0f94","observation_id":"2794b4f1-eeb1-45e2-bcbb-2ac90dc8244e","resolution":{"observed_at":"2026-08-07T01:06:45.860130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-07T00:00:11.342661Z","title":"https://doi.org/10.48550/arXiv.2402.12348 arXiv:2402.12348","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.15624","last_updated":"2025-06-18T16:53:38Z","snapshot_observed_at":"2026-08-08T15:45:16.463188Z","submitted_at":"2025-06-18T16:53:38Z","title":"The Effect of State Representation on LLM Agent Behavior in Dynamic Routing Games","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T00:00:11.342661Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.15624"},"observation_digest":"sha256:5586c1d3e8bd1ca86f906c9d49f0d37cb32d15a3e7e00363e300a4dd558f22d4","observation_id":"88deb9cc-21ec-4eea-a3d7-eeea63417592","resolution":{"observed_at":"2026-08-07T00:00:11.342661Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-06T22:58:09.308919Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.20353","last_updated":"2025-06-25T12:04:53Z","snapshot_observed_at":"2026-08-07T23:12:17.952101Z","submitted_at":"2025-06-25T12:04:53Z","title":"DipSVD: Dual-importance Protected SVD for Efficient LLM Compression","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T22:58:09.308919Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.20353"},"observation_digest":"sha256:5d3f61ab19302fb5dd6563af0afc3ebc61c2e8bd9a35f8adb19befddc7f4be89","observation_id":"ade39c0f-f51e-4c82-8f16-222cdfb2ae79","resolution":{"observed_at":"2026-08-06T22:58:09.308919Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-06T23:20:18.435764Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21620","last_updated":"2025-06-23T08:54:32Z","snapshot_observed_at":"2026-08-06T23:13:56.347786Z","submitted_at":"2025-06-23T08:54:32Z","title":"How Large Language Models play humans in online conversations: a simulated study of the 2016 US politics on Reddit","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:20:18.435764Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2506.21620"},"observation_digest":"sha256:d4dd4aabf88e49eb232dd04c5920ce1641341e5edbd51f658b4c6504ce98cf50","observation_id":"025f81de-cd78-4010-a445-bc1cccddaa54","resolution":{"observed_at":"2026-08-06T23:20:18.435764Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-06T17:57:08.266513Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.09495","last_updated":"2025-07-13T05:02:43Z","snapshot_observed_at":"2026-08-06T23:14:57.681420Z","submitted_at":"2025-07-13T05:02:43Z","title":"GenAI-based Multi-Agent Reinforcement Learning towards Distributed Agent Intelligence: A Generative-RL Agent Perspective","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T17:57:08.266513Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2507.09495"},"observation_digest":"sha256:bc4a9bea2041e5a10cee22dbf930b2b9780a551c8e2d25f2c95a659522d17319","observation_id":"611fddd2-c4d0-4182-8f8c-7c17f856217b","resolution":{"observed_at":"2026-08-06T17:57:08.266513Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T22:07:47.740830Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.07485","last_updated":"2025-08-10T21:07:08Z","snapshot_observed_at":"2026-08-08T15:44:48.200452Z","submitted_at":"2025-08-10T21:07:08Z","title":"Democratizing Diplomacy: A Harness for Evaluating Any Large Language Model on Full-Press Diplomacy","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-05T22:07:47.740830Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2508.07485"},"observation_digest":"sha256:c97b23d428a8fded5a522dca40870e2500664b7b29047525a8ebfe3a9972d81c","observation_id":"30f59d22-4284-4e1c-a3ca-b9322be083f2","resolution":{"observed_at":"2026-08-05T22:07:47.740830Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2604.21282","last_updated":"2026-04-23T04:58:18Z","snapshot_observed_at":"2026-07-06T23:07:56.487654Z","submitted_at":"2026-04-23T04:58:18Z","title":"Strategic Heterogeneous Multi-Agent Architecture for Cost-Effective Code Vulnerability Detection","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-09T22:06:47.552043Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2604.21282"},"observation_digest":"sha256:618b01c2ba5bcaf7ea64d5dabffd8eb9d36c8fcc3af8da0730e7b921a7543f68","observation_id":"12bcbc80-d732-47ad-be20-fc098b3ac09b","resolution":{"observed_at":"2026-05-11T14:16:28.941171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2605.13875","last_updated":"2026-05-08T06:56:35Z","snapshot_observed_at":"2026-08-02T22:27:33.584660Z","submitted_at":"2026-05-08T06:56:35Z","title":"Common-agency Games for Multi-Objective Test-Time Alignment","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-05-15T06:14:53.685486Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2605.13875"},"observation_digest":"sha256:e494c01b3a9b0098ade438798eb020fac70f8c859aedb7ef45ddeb621613b516","observation_id":"9f2990fc-a0f6-4512-8b83-4b798bd08e2d","resolution":{"observed_at":"2026-05-15T06:15:06.376948Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2605.24539","last_updated":"2026-05-23T12:10:52Z","snapshot_observed_at":"2026-08-06T12:48:39.939001Z","submitted_at":"2026-05-23T12:10:52Z","title":"DemoEvolve: Overcoming Sparse Feedback in Agentic Harness Evolution with Demonstrations","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-30T13:12:46.927103Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2605.24539"},"observation_digest":"sha256:cfd663e10996363189aa67295743855d29c38688763e7830ac96ede7d4735042","observation_id":"2f7e3d0e-74a3-4145-8dd5-c89289876e4b","resolution":{"observed_at":"2026-06-30T13:14:40.656457Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2606.06556","last_updated":"2026-06-04T10:43:14Z","snapshot_observed_at":"2026-08-09T06:57:25.072034Z","submitted_at":"2026-06-04T10:43:14Z","title":"Robots Need More than VLA and World Models","version":1},"reference_index":143,"source":"arxiv_source","source_observed_at":"2026-06-28T01:01:33.530167Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2606.06556"},"observation_digest":"sha256:65fb430564ed6d9e750aac672a51da71514aadc414ce45d2cda94e84c4d67ecc","observation_id":"aa4058ae-4165-46e4-a750-1b83aee8f181","resolution":{"observed_at":"2026-06-28T01:11:28.888729Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2606.12191","last_updated":"2026-06-10T15:15:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-06-10T15:15:01Z","title":"Agentic Environment Engineering for Large Language Models: A Survey of Environment Modeling, Synthesis, Evaluation, and Application","version":1},"reference_index":120,"source":"pdf_text","source_observed_at":"2026-06-27T09:46:30.702256Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2606.12191"},"observation_digest":"sha256:0e87b393a9a6b70647e2617411aa7238e179c83eda1eff99853d1db120bc4db7","observation_id":"d50edb7b-404d-45ca-b0c9-ac88b352d22b","resolution":{"observed_at":"2026-07-03T10:58:02.988053Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2606.19338","last_updated":"2026-06-17T17:59:34Z","snapshot_observed_at":"2026-07-06T23:54:37.596257Z","submitted_at":"2026-06-17T17:59:34Z","title":"Beyond the Current Observation: Evaluating Multimodal Large Language Models in Controllable Non-Markov Games","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-26T21:17:02.332687Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2606.19338"},"observation_digest":"sha256:68b4c096ab50ff4fd8de0d694fd1a0260b58f7608e3340b1d89b732de5d9ed04","observation_id":"1c59a98e-472a-4f27-a2aa-27ca3b39248e","resolution":{"observed_at":"2026-07-04T00:19:13.715084Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2606.24162","last_updated":"2026-06-23T05:30:54Z","snapshot_observed_at":"2026-08-06T14:03:04.558539Z","submitted_at":"2026-06-23T05:30:54Z","title":"BehaviorBench: Benchmarking Foundation Models for Behavioral Science Tasks","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-06-26T00:38:14.922388Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2606.24162"},"observation_digest":"sha256:fe54f688c0fe9f1c6a37da66cd5792c7ec569f009c0dd786c3dae5b259ceb8dc","observation_id":"b292097a-1f13-4384-a22d-90153f7ca041","resolution":{"observed_at":"2026-06-26T00:38:42.745298Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2606.24391","last_updated":"2026-06-23T10:25:31Z","snapshot_observed_at":"2026-08-08T22:58:51.830314Z","submitted_at":"2026-06-23T10:25:31Z","title":"Age of LLM: A Strategic 1v1 Benchmark for Reasoning, Diplomacy and Reliability of Large Language Models under Fog of War","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-25T23:49:25.799611Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2606.24391"},"observation_digest":"sha256:4a116503fd9baab4e55d9abbac11aea5e6fb43af4689fb9788e80f00291bba78","observation_id":"45aeb99a-2a1f-4e2d-8263-e1c954e94bb1","resolution":{"observed_at":"2026-07-04T17:20:00.182045Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":"2402.12348","doi":"10.48550/arxiv.2402.12348","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":"arXiv (Cornell University)","work_id":"693a1523-382d-4f21-acff-e36e876aa530","year":2024},"citing_paper":{"arxiv_id":"2606.27757","last_updated":"2026-06-26T06:24:33Z","snapshot_observed_at":"2026-07-07T00:01:54.432346Z","submitted_at":"2026-06-26T06:24:33Z","title":"Towards Reliable and Robust LLM Planning: Symbolic Feedback-Driven Iterative Self-Refinement Framework","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-29T04:53:17.854375Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2606.27757"},"observation_digest":"sha256:05fd007955c674123d27c0774dda72159e9aaf0d896d7f3e83bef901e7ba06b0","observation_id":"71dc7915-1b89-4660-a538-38354c769283","resolution":{"observed_at":"2026-06-29T19:13:52.970237Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-07-14T16:35:42.294582Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09743","last_updated":"2026-07-03T10:23:31Z","snapshot_observed_at":"2026-08-02T08:35:04.901722Z","submitted_at":"2026-07-03T10:23:31Z","title":"Scaffolding the Strategist: Architecture-Dependent Reasoning Interventions in Hotelling Spatial Markets","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-14T16:35:42.294582Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2607.09743"},"observation_digest":"sha256:2bdcf8d2b842f95f5e1d49c0f0571865086e095d56cd0eef9d296020d7890adc","observation_id":"d655b8b7-d299-4544-a174-9373c020718b","resolution":{"observed_at":"2026-07-14T16:35:42.294582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-01T06:09:13.779994Z","title":"Gtbench: Uncovering the strategic reasoning limitations of llms via game-theoretic evaluations","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.27536","last_updated":"2026-07-30T00:10:21Z","snapshot_observed_at":"2026-08-08T08:52:04.663177Z","submitted_at":"2026-07-30T00:10:21Z","title":"Strategy, Not Payoffs: A Behavioural Embedding of Normal-Form Games","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-01T06:09:13.779994Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2607.27536"},"observation_digest":"sha256:ed64dc92628f0fca298f548d673e4235d1cea53ad1f0cf7648cda4cdedfcb8a1","observation_id":"99ee0253-bde2-48e4-acba-bb511d2da326","resolution":{"observed_at":"2026-08-01T06:09:13.779994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2402.12348/citation-record","integrity":"/paper/2402.12348/integrity","json":"/paper/2402.12348/citation-record.json","paper":"/paper/2402.12348"},"outbound":[],"paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-08T15:45:02.939196Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 23 inbound Pith citation observations for arXiv:2402.12348."}