{"as_of":"2026-08-15T01:48:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bf9aef1dc7c4e8834118ccd08e8145c1c6eb51059185e14b219c06f72cea88d1","coverage":[{"denominator":68,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":68,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T23:05:34.799690Z","state":"measured"},{"denominator":68,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":68,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.09128/citation-record","integrity":"/paper/2608.09128/integrity","json":"/paper/2608.09128/citation-record.json","paper":"/paper/2608.09128"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.524828Z","title":"Acta Universitatis Sapientiae, Informatica , volume=","venue":null,"work_id":"3ad59943-91c4-438f-8c8f-31da4a54dfd8","year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.540529Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:bd068cf8cc50f68881c967fd10ee525b7aa9f30283f816f816093f481eb4da86","observation_id":"66099bef-a43e-4607-8aed-e7f68deef70a","resolution":{"observed_at":"2026-08-11T23:05:35.528855Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.512328Z","title":"NAACL , year=","venue":null,"work_id":"9c191617-8322-4b0f-8ef7-9a2d7500fe0e","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.546631Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:d8276ce2a4b180dfd5823b1e383b0ab2220cb78d0b4b5fe23129231802a1dff3","observation_id":"3daf168c-8bbb-42d9-aeae-cfefd8631438","resolution":{"observed_at":"2026-08-11T23:05:35.516410Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.498691Z","title":"2026 , booktitle=","venue":null,"work_id":"26eb5c15-8ed7-4baa-9d7c-cfa8d90ff36d","year":2026},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.550779Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:19e429db7c2c938edbf3bdcf19c4a71606f42192214e549e0cce147f1bae7b1c","observation_id":"ce18d556-0ae4-4e64-9da2-f90757690548","resolution":{"observed_at":"2026-08-11T23:05:35.502952Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.484778Z","title":"Findings of ACL , year=","venue":null,"work_id":"6979b374-c96c-4232-9c0c-6da0622ad772","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.554824Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:09cdd1d87119613f7451cce1526265ac2a3d6cbc2099d14c86e1af3b51bf46c0","observation_id":"29ab7034-16d0-42d2-8406-7e74c1caee57","resolution":{"observed_at":"2026-08-11T23:05:35.489237Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.473471Z","title":"2024 , booktitle=","venue":null,"work_id":"89fafae2-7c58-42ca-aa76-c428af6668b4","year":2024},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.559522Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:91b2e41054269a755fd4a6f078e0f146c41f58b7be54c2b7513e38fa812d9f11","observation_id":"594e5ffc-4f00-47f3-bef2-f6cf10b0011c","resolution":{"observed_at":"2026-08-11T23:05:35.477232Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.461379Z","title":"2025 , eprint=","venue":null,"work_id":"f4032f4a-e3bd-483c-a879-4868d1b92555","year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.563464Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:703a4cdb82da0d62ed307dc7a33d49abaa3de0bc32127c64bbb19fd877ee77ac","observation_id":"14d77b26-a516-41de-a5a5-469d56fcc0f1","resolution":{"observed_at":"2026-08-11T23:05:35.465794Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.567592Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.567592Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:4a3fbb2fbe7281c6c24492fd088facd017f0785519ec6cdb9c86d3fb8c698763","observation_id":"d74d4f7d-84bb-45ba-b997-44c7be788b18","resolution":{"observed_at":"2026-08-11T23:05:34.567592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.441785Z","title":"2025 , month = may, note =","venue":null,"work_id":"b50070c8-9b2a-42d3-89d7-f861b648d976","year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.571679Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:cf8090c3dc9301a99a63e6a0a24eea9bfa971244c941c930b247736b0a932dfc","observation_id":"6e832812-f636-4569-b8b4-254b609b2c2f","resolution":{"observed_at":"2026-08-11T23:05:35.445925Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.575440Z","title":"Proceedings of the 2023 Conference on Empirical Methods in Natural Language Processing , pages=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.575440Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:e6dacfd60910f0624d9b7a66178e7aa3001f217a0898cf8561b69ad51e808aca","observation_id":"89396d3a-c612-4fdf-a82c-907dfb2ad00e","resolution":{"observed_at":"2026-08-11T23:05:34.575440Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00069","last_updated":"2026-06-23T16:02:17Z","snapshot_observed_at":"2026-08-07T17:42:58.380771Z","submitted_at":"2025-02-27T13:26:07Z","title":"Societal Alignment Frameworks Can Improve LLM Alignment","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.00069","snapshot_observed_at":"2026-08-11T23:05:34.579870Z","title":"arXiv preprint arXiv:2503.00069 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.579870Z"},"links":{"cited_paper":"/paper/2503.00069","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:38474732ded36203169ced8aad3463181d8d6919024c635280af6eb0c77882d9","observation_id":"52c6cd20-333f-404f-b795-6f83819e7cf6","resolution":{"observed_at":"2026-08-11T23:05:34.579870Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.584416Z","title":"2005 , publisher=","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.584416Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:9117481e341dbea5b3bdaf6c749a348429e2c19a93780391560979eef5c1fd88","observation_id":"3675d94c-a18b-443c-9102-a62a273b4660","resolution":{"observed_at":"2026-08-11T23:05:34.584416Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1606.06565","last_updated":"2016-07-25T17:23:29Z","snapshot_observed_at":"2026-07-06T05:00:46.434335Z","submitted_at":"2016-06-21T13:37:05Z","title":"Concrete Problems in AI Safety","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.06565","snapshot_observed_at":"2026-08-11T23:05:34.588577Z","title":"arXiv preprint arXiv:1606.06565 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.588577Z"},"links":{"cited_paper":"/paper/1606.06565","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:ebde10eb1120bd49e27643aa65717a168b7673772a5361c90c4e3dc8b7bcce2d","observation_id":"8b79025d-e39c-446b-94b7-b8c4220a931c","resolution":{"observed_at":"2026-08-11T23:05:34.588577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2201.03544","last_updated":"2022-02-14T09:05:38Z","snapshot_observed_at":"2026-08-13T04:40:05.019883Z","submitted_at":"2022-01-10T18:58:52Z","title":"The Effects of Reward Misspecification: Mapping and Mitigating Misaligned Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.03544","snapshot_observed_at":"2026-08-11T23:05:34.592692Z","title":"arXiv preprint arXiv:2201.03544 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.592692Z"},"links":{"cited_paper":"/paper/2201.03544","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:cf073616ec4e3826865acd2cfddcfa1f7dc91d0b6a34d6122c7279ef952e4e62","observation_id":"d894df8e-ee2f-42df-9b5e-1ddcdea9fabb","resolution":{"observed_at":"2026-08-11T23:05:34.592692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.597498Z","title":"Artificial intelligence , volume=","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.597498Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:7bbc41450fbe4db91bb0315edc756ca6452df22c3ccb388845929b86c65b08cf","observation_id":"e58cea66-1fc2-4454-862b-38e8fe86ffcc","resolution":{"observed_at":"2026-08-11T23:05:34.597498Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-13T15:58:13.809876Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-11T23:05:34.601259Z","title":"arXiv preprint arXiv:2501.12948 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.601259Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:20f5dbd34145292b5db86aaed4a8fd7f94fff656220234ef87f21f588ae8f501","observation_id":"45bf8575-057e-4261-a378-331508fe4b8b","resolution":{"observed_at":"2026-08-11T23:05:34.601259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-11T23:05:34.605581Z","title":"arXiv preprint arXiv:2407.21783 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.605581Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:1297c22c46bebcd7684ab1bfe6002de513b1352b7c442726014e6e5fcf9a75d3","observation_id":"da6f06b8-be26-4c03-ac42-45b7dc6c5ee9","resolution":{"observed_at":"2026-08-11T23:05:34.605581Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.609259Z","title":"2024 , eprint=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.609259Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:7cde02aade54478682bf4f596f747bba72aa47f1bba29be4e37ff9611c322625","observation_id":"e9625301-19d0-46f4-a001-fee5262eba8c","resolution":{"observed_at":"2026-08-11T23:05:34.609259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.394617Z","title":"2025 , eprint=","venue":null,"work_id":"22d91e4c-9dc4-48f7-86d9-e40dc18f0f2d","year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.613177Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:65a27f44d7b3dc77475c3d732b38a85ef1c442fccfe6c0040fb05f481d797420","observation_id":"b0f12945-cad7-4cf5-8ef4-1249fd64fc6a","resolution":{"observed_at":"2026-08-11T23:05:35.398675Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.380589Z","title":"2025 , eprint=","venue":null,"work_id":"eb83f2c3-fc31-41ff-a43a-07f0d447c296","year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.616839Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:69eb0fbdad21ec71ff1da6a3a61b0812227173e30b24becac9f8da5853e119c3","observation_id":"f7dac24d-6f34-45eb-a0b2-19621337f3a0","resolution":{"observed_at":"2026-08-11T23:05:35.385597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.367597Z","title":"2026 , eprint=","venue":null,"work_id":"5901766b-3781-40dc-8cf1-aa268016e82b","year":2026},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.620698Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:adc2811b48443a6bf7d1ae1bf10acac0df2659abbabadffb3e455378208066aa","observation_id":"954e4e6b-ef1d-49a7-b228-086e3cca04e2","resolution":{"observed_at":"2026-08-11T23:05:35.371792Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.624472Z","title":"2025 , eprint=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.624472Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:f80c8e9045c5e326f0dda791e0202da81dfa7992ce33f2cf7ee641393f08a998","observation_id":"3d011c37-78ee-4b60-bb73-3d9b00f803b9","resolution":{"observed_at":"2026-08-11T23:05:34.624472Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.628227Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.628227Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:93cef7bd0466a1959c40efeac609f708c0b87821e7e5fc00000d2076e035c968","observation_id":"84d06b87-c67f-491f-b547-53ee684b4f78","resolution":{"observed_at":"2026-08-11T23:05:34.628227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.631953Z","title":"Proceedings of the 36th annual acm symposium on user interface software and technology , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.631953Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:cd4b86c003708e4bdcfd5555c11cb303f77fcb78e4c312101ae0ef9a213640aa","observation_id":"773ede4d-d5d0-4f6a-ae28-e37296867b06","resolution":{"observed_at":"2026-08-11T23:05:34.631953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.635924Z","title":"Proceedings of the National Academy of Sciences , volume=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.635924Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:5dfaaeee45557b489ebd19adbf00f5ea6212bee23ce871687c209cebec4832be","observation_id":"c8645620-ea0e-4777-b1cc-ccc40dd6a6ee","resolution":{"observed_at":"2026-08-11T23:05:34.635924Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.325603Z","title":"Science , volume=","venue":null,"work_id":"1d83f779-01d9-492c-aa49-225448d8bd5d","year":2022},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.639702Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:6e08076d24c56db964f3e36eda629251abbd520ed0c690385701e9ca473dae24","observation_id":"2894105d-3f0e-4a29-aa5b-eddbeedf496f","resolution":{"observed_at":"2026-08-11T23:05:35.329977Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04658","last_updated":"2024-05-11T07:08:16Z","snapshot_observed_at":"2026-08-13T10:17:34.745714Z","submitted_at":"2023-09-09T01:56:40Z","title":"Exploring Large Language Models for Communication Games: An Empirical Study on Werewolf","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.04658","snapshot_observed_at":"2026-08-11T23:05:34.643711Z","title":"arXiv preprint arXiv:2309.04658 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.643711Z"},"links":{"cited_paper":"/paper/2309.04658","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:9d7237391e5c7d27749fa0e29a5febb1c3414460b21d5c35ba76273118fb6611","observation_id":"bbe5a7ab-7a84-44d4-a768-0abee4856344","resolution":{"observed_at":"2026-08-11T23:05:34.643711Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.647924Z","title":"Nature Human Behaviour , volume=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.647924Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:8ab2830da840e4b7d3361a308da7ecdca14f4602908d9db4de6a4d31e79e1ac8","observation_id":"975706fe-9bd4-4fc4-81a7-b86d4ce42419","resolution":{"observed_at":"2026-08-11T23:05:34.647924Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.651661Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.651661Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:9eb515382f959cccd1d6b3a8401ef534371190ba3712e288f8719710fa40dde7","observation_id":"4454e18a-3ab3-4ea1-98bf-c82f13f29c35","resolution":{"observed_at":"2026-08-11T23:05:34.651661Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.655444Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.655444Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:9554e1c1d7ad887452bdd49b2ddb08e75c84f36548d183f4783599c6b051d7f3","observation_id":"bdf964a1-19f4-4463-a7e5-9bb85bd572c8","resolution":{"observed_at":"2026-08-11T23:05:34.655444Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.659162Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.659162Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:8321aab8b4abd8544f04e2e99c4d3d474a7229585ae62192df7461e2233df601","observation_id":"77601210-e2c9-4ae2-a34a-c04fb698737a","resolution":{"observed_at":"2026-08-11T23:05:34.659162Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.14985","last_updated":"2024-10-13T22:09:00Z","snapshot_observed_at":"2026-08-13T05:43:03.846551Z","submitted_at":"2023-10-23T14:35:26Z","title":"LLM-Based Agent Society Investigation: Collaboration and Confrontation in Avalon Gameplay","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.14985","snapshot_observed_at":"2026-08-11T23:05:34.662836Z","title":"arXiv preprint arXiv:2310.14985 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.662836Z"},"links":{"cited_paper":"/paper/2310.14985","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:4824e64903000393fd086f7e0028eae59af6cc03497aee100614b183b2383fff","observation_id":"17b30183-258b-43a1-985e-8288a569860e","resolution":{"observed_at":"2026-08-11T23:05:34.662836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.666916Z","title":"Proceedings of the 62nd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.666916Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:3c155bc342174e8539ec990f5c8b99ef2c257ac2d5137a3e39f97c12fd9f0928","observation_id":"127bc410-80ff-4c93-a5df-4d0663cca398","resolution":{"observed_at":"2026-08-11T23:05:34.666916Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.278031Z","title":"1978 , publisher=","venue":null,"work_id":"9b82ec9e-5a4f-48dc-b428-b3fe7b8ecb51","year":1978},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.670736Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:e8c4e96b24aff4140df3b4b11bf3ed54c56083e6c57759867c8508b5de2123f6","observation_id":"496b2712-ea19-4b53-a284-9e8520e591f6","resolution":{"observed_at":"2026-08-11T23:05:35.281839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.265918Z","title":"International Conference on Machine Learning , pages=","venue":null,"work_id":"f1f4a229-3e66-4fb6-a9b9-bb5e86a4529b","year":2024},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.674770Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:f95849c9bd8336e8e99fe951b9fcedf203849fcf6777a4d85dad05e2adb54faf","observation_id":"a19a1aeb-11ce-43c4-a3f4-6b1ea2129004","resolution":{"observed_at":"2026-08-11T23:05:35.270011Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.678632Z","title":"the method of paired comparisons , author=","venue":null,"work_id":null,"year":1952},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.678632Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:3e291ad1b2ff5544800dc6c10556fabcd6b0cf13c050aa2a1f0a758321698587","observation_id":"8dbaa87e-8e9a-4f7d-8174-a8fb22c37e6e","resolution":{"observed_at":"2026-08-11T23:05:34.678632Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.247198Z","title":"Harper's Magazine , volume=","venue":null,"work_id":"3325c885-c60b-42d7-84ec-57c5a5e672b8","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.682484Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:fc44571dbc3e1a0d17c70e8f065715dc11c6c97af4cbbb0074e0841f5d091541","observation_id":"bfa86fc2-5cce-4729-a821-5265ed3d9a80","resolution":{"observed_at":"2026-08-11T23:05:35.251302Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.686119Z","title":"2011 , publisher=","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.686119Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:74ef209c42e726a339b798b2ea343da9dc05cdc0df7a2b66187f1ad877725783","observation_id":"5e5c370e-7aab-42ea-b054-d79e08c65fed","resolution":{"observed_at":"2026-08-11T23:05:34.686119Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.228069Z","title":"science , volume=","venue":null,"work_id":"aa5197b8-edd5-4aa7-8ac2-8181aa2c56cb","year":1981},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.689919Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:37bb7c3d8f245cd7b271626668a09fa0ac96abb7bf6d1ab8042f150c0b59289b","observation_id":"e78b11c9-eff8-4e9c-a444-ab8939b2b668","resolution":{"observed_at":"2026-08-11T23:05:35.231837Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.215377Z","title":"1980 , publisher=","venue":null,"work_id":"424d646e-84e5-4eea-9ced-d0d254108b1d","year":1980},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.693614Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:f3f0fa4038da4484533351053e063ebc5ec0c335bb540c2f7958a3c20ef54e7d","observation_id":"70a90aee-3bb6-4f08-9859-8440827bb1d2","resolution":{"observed_at":"2026-08-11T23:05:35.219880Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.697208Z","title":"1990 , publisher=","venue":null,"work_id":null,"year":1990},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.697208Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:d0475fc946ffab7f736b4e61621509791848a16b2993ad33607500b831dee94d","observation_id":"9923f44d-4655-4f10-92c0-bff98c3bd522","resolution":{"observed_at":"2026-08-11T23:05:34.697208Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.195148Z","title":"American Economic Review , volume=","venue":null,"work_id":"e9cef004-1c79-4f08-815e-781136a8dafb","year":2000},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.700971Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:541a3b42412e8b98ba88d139c36e170a5d5f72e933ebc87ffe794eae5d756381","observation_id":"72eece5e-e793-4ef5-ba53-1c448cd309d5","resolution":{"observed_at":"2026-08-11T23:05:35.199317Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.182748Z","title":"Journal of Economic theory , volume=","venue":null,"work_id":"d53cbfef-1c9b-4fda-9e81-f99c691177cf","year":1981},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.704798Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:184c2ee38c65b0cf85ac6199d4d96c7ee1fe9e5c87449544b61f2550fdb30572","observation_id":"5e426b3f-8ced-4adf-9d53-be49d5a8976c","resolution":{"observed_at":"2026-08-11T23:05:35.186977Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.169801Z","title":"Econometrica: Journal of the Econometric Society , pages=","venue":null,"work_id":"f045906f-8cfb-4c75-81c8-3f11bd4e7095","year":1982},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.708341Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:c4257929e716dbf92a57ead242e0931827902bd59796f7834ebef4a04eaa02b4","observation_id":"d61c3931-0e43-4924-b939-4beba99641db","resolution":{"observed_at":"2026-08-11T23:05:35.174422Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.156685Z","title":", author=","venue":null,"work_id":"c68d70de-2fb0-467a-8d07-5436e5b40115","year":2003},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.711994Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:8d7992872ecccaf14ee06d8c6519073ced14fb18d4b82196d466ab5b8fdc9604","observation_id":"3b7977d3-78c1-4106-970a-0b4540b5b125","resolution":{"observed_at":"2026-08-11T23:05:35.160685Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.715483Z","title":"John thinks that Mary thinks that…","venue":null,"work_id":null,"year":1985},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.715483Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:952d60da8dbfc2c4d711b56eb850d47372f1a824f7ee4dc036e5f79a74a927a4","observation_id":"8fe176d4-dc5e-4082-8cc1-b7d4313b7eb5","resolution":{"observed_at":"2026-08-11T23:05:34.715483Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.143880Z","title":"Evolutionary Anthropology: Issues, News, and Reviews , volume=","venue":null,"work_id":"2a66f09d-4d0f-472b-a00f-f47c9c6a297b","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.721296Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:942c0d1614c7a8cd211ab91a545e8b1ba67d1420484bb9a10388279d2a1244b5","observation_id":"156be659-4703-4394-bce9-251668320b51","resolution":{"observed_at":"2026-08-11T23:05:35.147997Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.132691Z","title":"1984 , publisher=","venue":null,"work_id":"97b7d74d-d558-4f62-9c98-6aee5905feec","year":1984},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.725068Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:122d0932aafa441e0a5ead64c8a732de53b0739364cd47b9f3840e0a64a0a57c","observation_id":"191024f9-c656-491a-a870-4531d69f5959","resolution":{"observed_at":"2026-08-11T23:05:35.136750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.120787Z","title":"Cognition , volume=","venue":null,"work_id":"7aeea2aa-7414-4497-a623-4dc52d200277","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.728658Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:05ca78c0a488534ef9219cc32778089f1843a6b3496d9b9bd917031b46fd05d0","observation_id":"c96f0ba7-6dc2-493b-beb6-1dc1959b545d","resolution":{"observed_at":"2026-08-11T23:05:35.125061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.107848Z","title":"Cognition , volume=","venue":null,"work_id":"97f98523-c7c9-42e9-8b1a-2991c0445bd6","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.731838Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:ee2df96c3d8568f9cb0fc768d20a5a5d1142eec4437db078143fc3655ac745ee","observation_id":"dbe2ea54-809a-4c1a-919d-03136eda64ba","resolution":{"observed_at":"2026-08-11T23:05:35.111809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.095782Z","title":"1988 , publisher=","venue":null,"work_id":"b9a0a4a2-e4d2-4fce-a630-78630659631a","year":1988},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.735336Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:7d27f200a849089857ad69c90ac04a5e5fa210e26dede9618b10259a2265cb35","observation_id":"519e20e8-bb82-4bca-b07b-8c09befafff7","resolution":{"observed_at":"2026-08-11T23:05:35.099667Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.738866Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.738866Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:7f12d28a60aebcf9fb9358c26986f410d16b37a110631e346e6c7a211a68f582","observation_id":"2a5a11b2-97c1-415a-9ec8-739f58cfd8bf","resolution":{"observed_at":"2026-08-11T23:05:34.738866Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.742161Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.742161Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:90f9281275cd9aec10df1918158217beb46476fc42f1c3ff989acef119ff7403","observation_id":"8efb1f2c-664e-4b59-b3d8-904f78ce0a79","resolution":{"observed_at":"2026-08-11T23:05:34.742161Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.069224Z","title":"Proceedings of the 2022 conference on empirical methods in natural language processing , pages=","venue":null,"work_id":"cd3837af-51ba-4937-b734-c22344044588","year":2022},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.745589Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:281399667546f52210b20bc1b445f11138b2d8eaecc271337eb0e036a2434d74","observation_id":"9b046e81-6669-4544-82ea-8f22a48688d7","resolution":{"observed_at":"2026-08-11T23:05:35.073627Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.02083","last_updated":"2024-11-04T19:51:53Z","snapshot_observed_at":"2026-08-14T23:20:10.618972Z","submitted_at":"2023-02-04T03:50:01Z","title":"Evaluating Large Language Models in Theory of Mind Tasks","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.02083","snapshot_observed_at":"2026-08-11T23:05:34.748777Z","title":"arXiv preprint arXiv:2302.02083 , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.748777Z"},"links":{"cited_paper":"/paper/2302.02083","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:5db1d12d169ccd9c10a5fe1f036f9d73b46ae961164f50d47c52717b81200eba","observation_id":"0bdd7741-8875-4777-b65a-8942b19cddd3","resolution":{"observed_at":"2026-08-11T23:05:34.748777Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.08399","last_updated":"2023-03-14T13:47:26Z","snapshot_observed_at":"2026-08-13T12:43:18.044776Z","submitted_at":"2023-02-16T16:18:03Z","title":"Large Language Models Fail on Trivial Alterations to Theory-of-Mind Tasks","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.08399","snapshot_observed_at":"2026-08-11T23:05:34.752241Z","title":"arXiv preprint arXiv:2302.08399 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.752241Z"},"links":{"cited_paper":"/paper/2302.08399","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:2577dc0d4eaf6ae4beeb7fef6b1f1b6dd564435b0bf4be4e4dd444fbda489356","observation_id":"bfd5295b-c9a8-40a6-878a-bc1bde0975ae","resolution":{"observed_at":"2026-08-11T23:05:34.752241Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.05036","last_updated":"2023-11-08T16:01:32Z","snapshot_observed_at":"2026-08-13T05:54:39.576650Z","submitted_at":"2023-10-08T06:37:08Z","title":"AvalonBench: Evaluating LLMs Playing the Game of Avalon","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.05036","snapshot_observed_at":"2026-08-11T23:05:34.755550Z","title":"arXiv preprint arXiv:2310.05036 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.755550Z"},"links":{"cited_paper":"/paper/2310.05036","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:4e5c79ed0232e89de61dac330305d76aea2573dd4c12d4546ca132f4de363137","observation_id":"f800c627-c57e-4d71-9de6-51a68e340350","resolution":{"observed_at":"2026-08-11T23:05:34.755550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12348","last_updated":"2024-06-10T17:14:09Z","snapshot_observed_at":"2026-08-13T18:47:34.471203Z","submitted_at":"2024-02-19T18:23:36Z","title":"GTBench: Uncovering the Strategic Reasoning Limitations of LLMs via Game-Theoretic Evaluations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12348","snapshot_observed_at":"2026-08-11T23:05:34.759492Z","title":"arXiv , author=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.759492Z"},"links":{"cited_paper":"/paper/2402.12348","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:1148b0ea2d117a427ba6ac1254fab5d06b5acbf00e295e9638a4640d08e0276e","observation_id":"6f7e96a1-d04a-46b3-8b66-362b221ffcdd","resolution":{"observed_at":"2026-08-11T23:05:34.759492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.055080Z","title":"International Conference on Learning Representations , volume=","venue":null,"work_id":"3c2dded7-f0ee-43ae-90e7-af399af55026","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.764350Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:4c6cc149307984cf1b7c6aa30b9eada34e8f27b614211e58e13c4f70711dfee0","observation_id":"a9219dcf-52df-49df-b870-ed3972b466cb","resolution":{"observed_at":"2026-08-11T23:05:35.061007Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.16291","last_updated":"2023-10-19T16:27:03Z","snapshot_observed_at":"2026-08-13T16:55:46.498156Z","submitted_at":"2023-05-25T17:46:38Z","title":"Voyager: An Open-Ended Embodied Agent with Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.16291","snapshot_observed_at":"2026-08-11T23:05:34.767877Z","title":"arXiv preprint arXiv:2305.16291 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.767877Z"},"links":{"cited_paper":"/paper/2305.16291","citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:29d5b809a8bf99bda4819407e8d4e9ec7eb6c226c490b996128480323be9862e","observation_id":"ffbaccbe-f001-4bae-aca2-42e1d6af22eb","resolution":{"observed_at":"2026-08-11T23:05:34.767877Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.771496Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.771496Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:a1d1f770fb2383ef060a0abc87dd2513f262e2d40a90e757850728eb93dae043","observation_id":"0d998186-d093-463f-b125-9c66efc035c4","resolution":{"observed_at":"2026-08-11T23:05:34.771496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.034525Z","title":"Handbook of Intelligence , editor=","venue":null,"work_id":"659513ca-fe9e-40b2-b3c3-ce0ea4bf91d5","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.774899Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:dd397d91bf0647a67f97950e09f22a6e910816bee3979e840d7f5c3854ef9bbd","observation_id":"3fda412d-e11f-4999-9d0e-50d33c78b3ed","resolution":{"observed_at":"2026-08-11T23:05:35.039264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.778497Z","title":"2025 , eprint=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.778497Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:47c785849cf0ceb73d9af8e3584f3e00f20d8bcb0efdff8f9b0e39d40e43f92e","observation_id":"81347811-66ea-4fea-af28-c0484b9fe772","resolution":{"observed_at":"2026-08-11T23:05:34.778497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.781939Z","title":"2023 , eprint=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.781939Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:c188c010a2824bbf5108fb4a9e96b9df8df73070a49d8a76df01ae6b357c0d40","observation_id":"e596d2f7-e9ad-4f6c-ac81-4781c8d7da5d","resolution":{"observed_at":"2026-08-11T23:05:34.781939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:35.006591Z","title":"Proceedings of the 41st International Conference on Machine Learning , pages=","venue":null,"work_id":"ce87d6c7-1644-4714-9aa5-96dae08227a7","year":null},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.785370Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:9c85d5e25a81ff7f4cb6d83b78d0d35d10fd4cfdb0730dfef47bb3905b3a6e70","observation_id":"8898c9df-a433-4473-89cc-42766f86ecae","resolution":{"observed_at":"2026-08-11T23:05:35.011999Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.788809Z","title":"Proceedings of the 2023 Conference on Empirical Methods in Natural Language Processing , pages=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.788809Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:6e1abe31e8fa2aea2159193a016d99cd71ff8a8c6fa54c3acd0436859e6e4768","observation_id":"5c155570-937a-4238-9a6d-c70946b3ad1f","resolution":{"observed_at":"2026-08-11T23:05:34.788809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.792530Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.792530Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:50152ceada0de36246f3caecec1b6ef864aada9c9f0012dede4296db6c4624b8","observation_id":"2cbbefed-68ef-4c45-9eea-87278b528cb0","resolution":{"observed_at":"2026-08-11T23:05:34.792530Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.987712Z","title":"2024 , publisher=","venue":null,"work_id":"de793dfe-b2a4-4394-ba91-df647aa3524e","year":2024},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.796172Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:1b0139172308c9bc7d1ea7d554d70f77068bd99fe47a9c8e03d11dbea0e0290f","observation_id":"77bc36d6-7d6d-482c-a4b2-f9903d3755ec","resolution":{"observed_at":"2026-08-11T23:05:34.991239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:05:34.973738Z","title":"Journal of Consumer Research , volume=","venue":null,"work_id":"4b9067bb-c2e6-47d7-bce9-e2857a2b1e5a","year":2026},"citing_paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-11T23:05:34.799690Z"},"links":{"citing_paper":"/paper/2608.09128"},"observation_digest":"sha256:26757a715ac9320df86ae1d2f06ded7ac4a0dd33d9a32a634237811ac9e57d02","observation_id":"89df69cd-4405-40d8-a20b-7102f040f3d0","resolution":{"observed_at":"2026-08-11T23:05:34.979116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2608.09128","last_updated":"2026-08-10T05:12:56Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-14T22:26:09.075446Z","submitted_at":"2026-08-10T05:12:56Z","title":"Social Gym and SPaRTan: Benchmarking and Improving LLM Social Reasoning via Multi-Agent Game Tournaments"},"reference_resolution":{"displayed":68,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":37,"verified_exact":0,"verified_fuzzy":31},"total_outbound_references":68},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 15 August 2026, this Paper Citation Record lists 68 of 68 outbound references and 0 inbound Pith citation observations for arXiv:2608.09128."}