{"as_of":"2026-08-09T19:50:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:95f0c6920d2a8564870de2864dc224ee125e36bcdac5c05c33b96e9ba8fd5e77","coverage":[{"denominator":64,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":64,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T15:22:45.112343Z","state":"measured"},{"denominator":67,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":67,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:02:22.194525Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T22:37:25.772308Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06470","snapshot_observed_at":"2026-08-07T14:02:22.194525Z","title":"A survey of theory of mind in large lang uage models: Evaluations, representations, and safety risks","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20120","last_updated":"2025-05-26T15:22:04Z","snapshot_observed_at":"2026-08-09T05:33:18.875561Z","submitted_at":"2025-05-26T15:22:04Z","title":"Agents Require Metacognitive and Strategic Reasoning to Succeed in the Coming Labor Markets","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T14:02:22.194525Z"},"links":{"cited_paper":"/paper/2502.06470","citing_paper":"/paper/2505.20120"},"observation_digest":"sha256:4b464dfd9a496836369970a252e1f548b7a5714235f4d73128dd86eeda5871b6","observation_id":"822582f7-0c6d-429d-9f80-a40485da0b2b","resolution":{"observed_at":"2026-08-07T14:02:22.194525Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"cited_work":{"arxiv_id":"2502.06470","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06470","snapshot_observed_at":"2026-07-02T22:37:25.772308Z","title":"arXiv preprint arXiv:2502.06470 , year=","venue":null,"work_id":"9a9bb270-7516-45b1-a2ae-fcc716a98654","year":null},"citing_paper":{"arxiv_id":"2605.15205","last_updated":"2026-04-28T15:38:31Z","snapshot_observed_at":"2026-08-08T12:32:02.265510Z","submitted_at":"2026-04-28T15:38:31Z","title":"Does Theory of Mind Improvement Really Benefit Human-AI Interactions? Empirical Findings from Interactive Evaluations","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-19T17:55:04.458343Z"},"links":{"cited_paper":"/paper/2502.06470","citing_paper":"/paper/2605.15205"},"observation_digest":"sha256:795a77734e5a4d8354c70fb0f550aec6a7d9bb475ce91ea087ee4740626da75b","observation_id":"841517d8-08d0-4bd8-9042-0a5d2ad0f541","resolution":{"observed_at":"2026-05-19T17:57:42.627226Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"cited_work":{"arxiv_id":"2502.06470","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06470","snapshot_observed_at":"2026-07-02T22:37:25.772308Z","title":"arXiv preprint arXiv:2502.06470 , year=","venue":null,"work_id":"9a9bb270-7516-45b1-a2ae-fcc716a98654","year":null},"citing_paper":{"arxiv_id":"2606.08274","last_updated":"2026-07-25T20:16:49Z","snapshot_observed_at":"2026-08-07T16:30:40.027384Z","submitted_at":"2026-06-06T17:40:21Z","title":"Toward Human-Centered Multi-Agent Systems: Integrating Cognition, Culture, Values, and Cooperation in AI Agents","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-06-27T18:47:36.189582Z"},"links":{"cited_paper":"/paper/2502.06470","citing_paper":"/paper/2606.08274"},"observation_digest":"sha256:bfbfdf706be0de5fe797f4983aee1d5fd4285d9928e1e6ddb97d5ff6fe3e8da0","observation_id":"50bbfd55-d1c8-4153-99a4-f826df319961","resolution":{"observed_at":"2026-07-02T22:37:25.774226Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.06470/citation-record","integrity":"/paper/2502.06470/integrity","json":"/paper/2502.06470/citation-record.json","paper":"/paper/2502.06470"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:44.769888Z","title":", \" * write output.state after.block = add.period write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.769888Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:9bd2cbfa903bdd8090a710665b10d0b45af11afa10d1526c4520748129dd0fb6","observation_id":"6cd02e11-537a-4cc7-98e6-8f2b1124796a","resolution":{"observed_at":"2026-08-08T15:22:44.769888Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:44.776121Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.776121Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:342df44293337ca048d2fd285b8d439cb66be5df5858a4d5b1beede647478a4f","observation_id":"206b6ca1-d8df-48a8-8663-cab022113c61","resolution":{"observed_at":"2026-08-08T15:22:44.776121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1610.01644","last_updated":"2018-11-22T23:40:00Z","snapshot_observed_at":"2026-07-06T05:13:30.860932Z","submitted_at":"2016-10-05T20:59:01Z","title":"Understanding intermediate layers using linear classifier probes","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1610.01644","snapshot_observed_at":"2026-08-08T15:22:44.782384Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.782384Z"},"links":{"cited_paper":"/paper/1610.01644","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:777833449e4152128f65207066aa0a64be792feceb17e09f1bd5c399d8698b0c","observation_id":"a877984e-bb44-4910-8304-89fc8beed9f0","resolution":{"observed_at":"2026-08-08T15:22:44.782384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.01781","last_updated":"2024-07-03T11:20:43Z","snapshot_observed_at":"2026-07-31T17:52:54.507082Z","submitted_at":"2024-02-01T19:12:25Z","title":"When Benchmarks are Targets: Revealing the Sensitivity of Large Language Model Leaderboards","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.01781","snapshot_observed_at":"2026-08-08T15:22:44.788604Z","title":"A.; Alnumay, Y.; Alrashed, S.; Alsubaie, S.; Almushaykeh, Y.; Mirza, F.; Alotaibi, N.; Altwairesh, N.; Alowisheq, A.; Bari, M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.788604Z"},"links":{"cited_paper":"/paper/2402.01781","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:692aa7e60ca52ebf1087605cf5d5fa1711145fb562d337988b43af73876ea18f","observation_id":"cc8883a8-ff25-4446-a322-8650f92e5d4f","resolution":{"observed_at":"2026-08-08T15:22:44.788604Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.204493Z","title":"S.; Jenner, E.; Casper, S.; Sourbut, O.; Edelman, B","venue":null,"work_id":"785e1ed0-5500-44e8-850e-3a25b329f223","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.794498Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:62b0dcf942ebb2cc0d03f4ae213f21b9f1bf253973fb75bccf0e5d09ba9edd10","observation_id":"0b887c14-4be0-4c01-838d-5625f8d50d92","resolution":{"observed_at":"2026-08-08T15:22:46.210568Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.185712Z","title":"theory of mind","venue":null,"work_id":"151e7bd7-733e-47fc-a7fd-59fe1b0f26c2","year":2012},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.800330Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:e31fafd064821816a45d0266112fafeabf561f1e28bcb5db7c5a1be4eee58f31","observation_id":"0746af7b-cddc-4883-ad73-40ed2558eedd","resolution":{"observed_at":"2026-08-08T15:22:46.191345Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.167551Z","title":null,"venue":null,"work_id":"6f27db28-000a-42be-a271-38c526b413a6","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.805659Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:082152f0951f64f4350608f44122a2622c8f2574e3889da13dfe603252e864a0","observation_id":"de61f0dc-12f4-466a-97a9-87f51cb6a01f","resolution":{"observed_at":"2026-08-08T15:22:46.173063Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05030","last_updated":"2025-07-29T09:44:14Z","snapshot_observed_at":"2026-08-07T17:43:00.958834Z","submitted_at":"2024-03-08T04:22:48Z","title":"Defending Against Unforeseen Failure Modes with Latent Adversarial Training","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05030","snapshot_observed_at":"2026-08-08T15:22:44.811321Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.811321Z"},"links":{"cited_paper":"/paper/2403.05030","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:52ce3089747027b75c0a93e7879968b78755d27137461bf14c99605be2c32c25","observation_id":"29c2bc8d-cc02-46d0-95ea-e758b4a3f63c","resolution":{"observed_at":"2026-08-08T15:22:44.811321Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07882","last_updated":"2024-10-14T17:46:28Z","snapshot_observed_at":"2026-07-06T18:29:21.709395Z","submitted_at":"2024-06-12T05:20:16Z","title":"Designing a Dashboard for Transparency and Control of Conversational AI","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07882","snapshot_observed_at":"2026-08-08T15:22:44.817749Z","title":"C.; Patel, O.; Riecke, J.; Raval, S.; Seow, O.; Wattenberg, M.; and Viégas, F","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.817749Z"},"links":{"cited_paper":"/paper/2406.07882","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:75eaa09e8fde3ed76b876c2ff547bc323c8d8b6ae155c3c59d396c9b04d996a9","observation_id":"dcf5387b-fd12-4489-aeac-4e22d6570019","resolution":{"observed_at":"2026-08-08T15:22:44.817749Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.149075Z","title":null,"venue":null,"work_id":"2fdada03-28e5-43f9-a308-0a931584522d","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.824077Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:74cd150368ca2cbf2cccb61df058dcb47c8120f93117c2af27f99e80fc43263d","observation_id":"2efbbbad-a974-46bf-96ad-17a5026d30f3","resolution":{"observed_at":"2026-08-08T15:22:46.154807Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.07413","last_updated":"2023-12-12T16:34:19Z","snapshot_observed_at":"2026-07-06T17:00:33.386675Z","submitted_at":"2023-12-12T16:34:19Z","title":"AI capabilities can be significantly improved without expensive retraining","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.07413","snapshot_observed_at":"2026-08-08T15:22:44.829284Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.829284Z"},"links":{"cited_paper":"/paper/2312.07413","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:e65fa18e3460957ba7d0c1d87a72d978d9d9c4d0eeae30918d50c33c9f6f83bf","observation_id":"31dc4491-51cd-4743-b9b5-81d2460fdd41","resolution":{"observed_at":"2026-08-08T15:22:44.829284Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.132073Z","title":null,"venue":null,"work_id":"3fa17c7d-8c60-414f-968e-54f0789c5841","year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.834477Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:cedd64b192bf95510edc77d2e3c2ffbc70a87de7e6e08b05f5d9f5d76860603a","observation_id":"573e6e4f-86e2-48e0-b037-5afba48391e0","resolution":{"observed_at":"2026-08-08T15:22:46.137282Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.114770Z","title":null,"venue":null,"work_id":"623d1248-60f1-4e81-af2f-556e9c2a9491","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.839605Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:c0511d6819e89e58aa184c04a17d44a537e214c17f2d8cee41e030cbbebd66db","observation_id":"9b9d80f4-7702-4146-9f78-c8b104a6801c","resolution":{"observed_at":"2026-08-08T15:22:46.120361Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:44.845044Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.845044Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:791527f452002ea489eb2d7ed4fb90261789fb5d0e8abbeb5d1b7b2fcbaf25e6","observation_id":"7ed83bd0-703b-4e47-aaf9-67b0cddebfc8","resolution":{"observed_at":"2026-08-08T15:22:44.845044Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06677","last_updated":"2023-12-04T07:51:58Z","snapshot_observed_at":"2026-07-06T17:00:00.321552Z","submitted_at":"2023-12-04T07:51:58Z","title":"Intelligent Virtual Assistants with LLM-based Process Automation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06677","snapshot_observed_at":"2026-08-08T15:22:44.850403Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.850403Z"},"links":{"cited_paper":"/paper/2312.06677","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:758d9db28a2df2303e87ec04c5ef78c5af65ae9ce8e866a8e9ffed1e2e9746f7","observation_id":"30945236-e2b6-4f7c-b67e-affae368ed34","resolution":{"observed_at":"2026-08-08T15:22:44.850403Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.087533Z","title":null,"venue":null,"work_id":"82a3850a-2481-432e-843c-9d31d5af6a06","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.855893Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:26bce526d61a825cb7d0a09888234fb804dccc5affcea056bc5b4eea7625f972","observation_id":"8f902a8a-de8b-4d30-a549-24b9c9e8a546","resolution":{"observed_at":"2026-08-08T15:22:46.092685Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1906.01820","last_updated":"2021-12-01T11:22:52Z","snapshot_observed_at":"2026-08-01T23:32:03.638122Z","submitted_at":"2019-06-05T04:43:25Z","title":"Risks from Learned Optimization in Advanced Machine Learning Systems","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1906.01820","snapshot_observed_at":"2026-08-08T15:22:44.860935Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.860935Z"},"links":{"cited_paper":"/paper/1906.01820","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:d08e0e434c33a720f62dbcd11e116e34995a003c036792e809103284ab100406","observation_id":"10da55c9-07aa-4dec-9455-7a4d5f6f9e2c","resolution":{"observed_at":"2026-08-08T15:22:44.860935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1805.00899","last_updated":"2018-10-22T17:36:07Z","snapshot_observed_at":"2026-08-02T15:33:17.783178Z","submitted_at":"2018-05-02T16:27:32Z","title":"AI safety via debate","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.00899","snapshot_observed_at":"2026-08-08T15:22:44.871075Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.871075Z"},"links":{"cited_paper":"/paper/1805.00899","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:c528d663b728818e336a71494bb30ca16daf009226248131613290c507c1b995","observation_id":"2b8e680f-b342-4078-9380-d66d987ad91d","resolution":{"observed_at":"2026-08-08T15:22:44.871075Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.01660","last_updated":"2023-09-04T15:26:15Z","snapshot_observed_at":"2026-08-07T16:57:20.985569Z","submitted_at":"2023-09-04T15:26:15Z","title":"Unveiling Theory of Mind in Large Language Models: A Parallel to Single Neurons in the Human Brain","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.01660","snapshot_observed_at":"2026-08-08T15:22:44.876548Z","title":"M.; and Cai, J","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.876548Z"},"links":{"cited_paper":"/paper/2309.01660","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:6debcc271c4d20666571b87fe1aaec20e435edba27f2fdafca3d7abeec2b4357","observation_id":"933ebd2d-c29a-4a63-af67-3f2f0d5b513c","resolution":{"observed_at":"2026-08-08T15:22:44.876548Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.01576","last_updated":"2024-04-25T17:29:53Z","snapshot_observed_at":"2026-07-06T18:09:00.688840Z","submitted_at":"2024-04-25T17:29:53Z","title":"Uncovering Deceptive Tendencies in Language Models: A Simulated Company AI Assistant","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.01576","snapshot_observed_at":"2026-08-08T15:22:44.881881Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.881881Z"},"links":{"cited_paper":"/paper/2405.01576","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:4b64761fcb84e8d955d6404cceced86379e83e78ed2004cce31987c601f1089c","observation_id":"f7702c32-0ac5-45b9-8bd8-3d9076b32dac","resolution":{"observed_at":"2026-08-08T15:22:44.881881Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.069388Z","title":"Y.; Kramar, J.; Brown-Cohen, J.; Albanie, S.; Bulian, J.; Agarwal, R.; Lindner, D.; Tang, Y.; Goodman, N.; and Shah, R","venue":null,"work_id":"5377f2ff-94f5-48a2-b4ea-aa8a09fa19fb","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.887246Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:c2dfe4c7cac3dc7bc20aa8d3eefa62b791c5ae9d6ca8721dec240adea73a038f","observation_id":"16b7ce84-6f18-43fd-9187-040d03b6a75d","resolution":{"observed_at":"2026-08-08T15:22:46.075378Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.052211Z","title":null,"venue":null,"work_id":"aacda28d-16c5-4e06-8044-f3c29d3b65d1","year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.892203Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:2647e8f5713e7aa6889d7748371f2be9ef8cff4f6a4219fe0bf806374452d7ae","observation_id":"80c60722-ab96-48ed-864c-c318fbbdf98a","resolution":{"observed_at":"2026-08-08T15:22:46.057915Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.035731Z","title":"M.; Kundu, A.; Jawhar, S.; Park, J.; and Jurewicz, M","venue":null,"work_id":"7e33d7c3-702e-4e97-a272-53d0f97d1abc","year":2025},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.897457Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:c7775f4288824ed2c0b432740f7017d52af3eff16966ee1ca7ab1bb265313fe4","observation_id":"3d3ad469-529f-4724-9e18-127e5e95c28b","resolution":{"observed_at":"2026-08-08T15:22:46.040914Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.020051Z","title":null,"venue":null,"work_id":"8996663a-83c0-4051-94b6-992aa8545d2b","year":2021},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.902260Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:44ab27b4638fad14f58e666d3d7cdb9901271c5240f8d4292ae52fd5f7e5e479","observation_id":"fdf842a5-b312-4806-a4d4-f53831836d5e","resolution":{"observed_at":"2026-08-08T15:22:46.024838Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:46.004096Z","title":"Q.; Stepputtis, S.; Campbell, J.; Hughes, D.; Lewis, C","venue":null,"work_id":"5debd7c7-421d-4d19-8e50-6878e179e0c6","year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.907051Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:a9d332fb73b37afeca8846b516bfd211508a317e29531ae503ae4719be42c335","observation_id":"8f2bd0b3-d6aa-4505-9e26-d1f0ae003aec","resolution":{"observed_at":"2026-08-08T15:22:46.009315Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.987193Z","title":"D.; Dombrowski, A.-K.; Goel, S.; Mukobi, G.; Helm-Burger, N.; Lababidi, R.; Justen, L.; Liu, A","venue":null,"work_id":"242a521b-ce2b-44de-a2d2-225d60d06113","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.911934Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:4d8ae3b0e499d40e356e78e119f38234465afa6d231497b5050116dd784a1c79","observation_id":"89d5da1c-959a-405f-bed3-532545066a16","resolution":{"observed_at":"2026-08-08T15:22:45.992958Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.970897Z","title":null,"venue":null,"work_id":"cbe5e442-2345-4561-b941-ef4eb1586042","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.917280Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:dcf78f908b24a6601188e56daff0b01575cbbd484a4772be8cda0db2959d9744","observation_id":"9774ac63-f2b6-4993-8d80-81d6ad9d400a","resolution":{"observed_at":"2026-08-08T15:22:45.976152Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08787","last_updated":"2024-12-06T21:39:49Z","snapshot_observed_at":"2026-08-09T11:53:57.275236Z","submitted_at":"2024-02-13T20:51:58Z","title":"Rethinking Machine Unlearning for Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08787","snapshot_observed_at":"2026-08-08T15:22:44.922410Z","title":"Y.; Xu, X.; Li, H.; Varshney, K","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.922410Z"},"links":{"cited_paper":"/paper/2402.08787","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:59b107dfcb219b346b1a9cde124d69d834df16f981fe728fae233e696a09e70b","observation_id":"4b8c5de1-8909-4e67-8d58-6f9a875dd5e6","resolution":{"observed_at":"2026-08-08T15:22:44.922410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.954886Z","title":"S.; Cope, D.; and Schoots, N","venue":null,"work_id":"ba6d1d03-e369-4cd9-b91e-b73a2e5d76e4","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.927713Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:2c7b08cbaadf4e5f093571e8809f55357192d6cc45548af506331ef7be7862da","observation_id":"57490763-0d01-4943-8e3b-de84b2446e4c","resolution":{"observed_at":"2026-08-08T15:22:45.959915Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.937389Z","title":"R.; Baranchuk, M.; Strohmeier, M.; Bolina, V.; Torr, P.; Hammond, L.; and de Witt, C","venue":null,"work_id":"a418e358-8d0d-4f4c-a767-a15996fc77e1","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.932771Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:337b9129a0b39dc639bad18fd511276b085a891d23f45fe7e02036778deb4679","observation_id":"893bfeee-1300-4124-874d-4ebec08404da","resolution":{"observed_at":"2026-08-08T15:22:45.943348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08901","last_updated":"2023-10-13T07:15:32Z","snapshot_observed_at":"2026-08-04T14:01:14.844864Z","submitted_at":"2023-10-13T07:15:32Z","title":"Welfare Diplomacy: Benchmarking Language Model Cooperation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08901","snapshot_observed_at":"2026-08-08T15:22:44.937791Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.937791Z"},"links":{"cited_paper":"/paper/2310.08901","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:e05a0522228f88ed852201f117755cd399728420d06ce9d7a540cc5330d1d338","observation_id":"4e1c8841-c02f-4d93-8722-3e329ad24936","resolution":{"observed_at":"2026-08-08T15:22:44.937791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.921303Z","title":null,"venue":null,"work_id":"7a435775-f2fe-4cd2-b1a6-e4b547d95a80","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.943052Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:c8ee09f14a429397133963d6dc8e3ad7884d53a83d4cd72536faafdf2eea2808","observation_id":"3c97cda4-e145-48cc-a58b-ba06e4c45106","resolution":{"observed_at":"2026-08-08T15:22:45.926410Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-08T15:22:44.948112Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.948112Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:173bf2ebf77b9463411d3548f91f8804856ed224e1cda44ceebb38f301754054","observation_id":"c3de1ec6-b36a-475a-b61e-2c634e000f78","resolution":{"observed_at":"2026-08-08T15:22:44.948112Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.905313Z","title":null,"venue":null,"work_id":"ce7a27bd-a1f3-4c5e-b9a3-3b8e9e782839","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.953589Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:c85a558823e59cb5eb5ac9a9ac589db9de24da20392efa8d259ee2f8563b6dc6","observation_id":"1d6a2552-dec3-405b-a988-6d2b7e4b2826","resolution":{"observed_at":"2026-08-08T15:22:45.910445Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.03442","last_updated":"2023-08-06T00:21:19Z","snapshot_observed_at":"2026-07-06T15:13:13.695535Z","submitted_at":"2023-04-07T01:55:19Z","title":"Generative Agents: Interactive Simulacra of Human Behavior","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.03442","snapshot_observed_at":"2026-08-08T15:22:44.958547Z","title":"S.; O'Brien, J","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.958547Z"},"links":{"cited_paper":"/paper/2304.03442","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:ad36ea32b69a7836e43e06c9a48236e997d6da5705fed68e1eeba4588a663265","observation_id":"36d6f1b0-29c1-450d-baf5-214319ce0742","resolution":{"observed_at":"2026-08-08T15:22:44.958547Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.886969Z","title":"S.; Goldstein, S.; O’Gara, A.; Chen, M.; and Hendrycks, D","venue":null,"work_id":"2bc1c3ef-f63d-4f11-a5ce-4426e6615026","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.964464Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:77938f710a56213cbcf419a04d9640792d57393ccb7598fc1af7de815de585b8","observation_id":"9484155d-7282-4476-bcba-696e0bc11699","resolution":{"observed_at":"2026-08-08T15:22:45.892387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.870956Z","title":null,"venue":null,"work_id":"07958d4a-9fba-43c0-920c-73e9ee714256","year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.969455Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:ba7d7b938d66a857d88562e1a95cfd7084b22e4eaa31a3f1c41c66dc3478bb4a","observation_id":"d071f3cc-b5e0-4a43-b38e-4e0cb6c9ffa1","resolution":{"observed_at":"2026-08-08T15:22:45.876076Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.853612Z","title":null,"venue":null,"work_id":"dd1378d7-7555-472f-b497-7d2e14663dec","year":2022},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.974926Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:4445a5b6681d2e913d7f36317d3ef68896523e9c446e793880b93f791c06087e","observation_id":"21f19ec5-6798-4ca7-b008-49f6f07a591b","resolution":{"observed_at":"2026-08-08T15:22:45.858963Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.836150Z","title":null,"venue":null,"work_id":"90860f49-8e14-4d5b-8ec9-96ccbf10a70f","year":1978},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.980417Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:6de942adbb56da638284a2e81f11073de68a05c6481f4ed23ba7b812213fee7b","observation_id":"de7d4bf4-d26a-4ecd-84c7-e22222b2a28e","resolution":{"observed_at":"2026-08-08T15:22:45.841531Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.815313Z","title":null,"venue":null,"work_id":"7748f35f-89cd-4eac-b618-13a531e9cf28","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.985768Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:9b7e15fe235b4593c7d342ff989f793b73a5be2ddeee0e2a1deaceb00cc7a482","observation_id":"e7ac3f25-1cb4-40d0-a48f-ab403d9f84fe","resolution":{"observed_at":"2026-08-08T15:22:45.821466Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.797082Z","title":null,"venue":null,"work_id":"e0ba84bd-719e-4140-80a0-0ce1838bf54d","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.991158Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:0216143464b5c14151a4eb3cc2b93ee26bbf1f728408ab3fd1f8da3e5af0ac56","observation_id":"5f9dde6a-183b-4b27-85fd-b14beac53a30","resolution":{"observed_at":"2026-08-08T15:22:45.802716Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.15943","last_updated":"2025-02-04T03:38:57Z","snapshot_observed_at":"2026-07-06T18:19:38.561528Z","submitted_at":"2024-05-24T21:14:10Z","title":"Transformers represent belief state geometry in their residual stream","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.15943","snapshot_observed_at":"2026-08-08T15:22:44.996501Z","title":"S.; Marzen, S","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:44.996501Z"},"links":{"cited_paper":"/paper/2405.15943","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:128fc0152acbf0b5087150a9dfab422b770964515352b2ff022a36e55d8bd498","observation_id":"0ee4f544-be7d-4c64-8f5d-0b8459439d51","resolution":{"observed_at":"2026-08-08T15:22:44.996501Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.779373Z","title":"H.; Zhou, X.; Choi, Y.; Goldberg, Y.; Sap, M.; and Shwartz, V","venue":null,"work_id":"4a84faea-8f4a-4837-99fc-ed812fce6c43","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.001999Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:2b846e9fef820c12242a05124e3cd5f374ed6410f06202ff8c38360caaa7925b","observation_id":"25876f44-6dc4-4f82-8e67-aa834d41b1c2","resolution":{"observed_at":"2026-08-08T15:22:45.785287Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.03314","last_updated":"2024-08-06T17:35:05Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-06T17:35:05Z","title":"Scaling LLM Test-Time Compute Optimally can be More Effective than Scaling Model Parameters","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.03314","snapshot_observed_at":"2026-08-08T15:22:45.007633Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.007633Z"},"links":{"cited_paper":"/paper/2408.03314","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:901b5e31a9cb097bd6a074703c510502f7c7d2472cca48de437a621d7c52ccfb","observation_id":"24f2bfbf-c2ba-4a96-8ef6-bba8fa773dfd","resolution":{"observed_at":"2026-08-08T15:22:45.007633Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.762983Z","title":null,"venue":null,"work_id":"5df232b0-7ace-437f-87de-f63e13f272b4","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.014937Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:71d92c201e9d20529b2f8f51ad2f5b36810762f022bf4e2a791361938fc808ac","observation_id":"eb44456a-1b3b-4e55-8dd2-7a80243ecfe0","resolution":{"observed_at":"2026-08-08T15:22:45.768040Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.746377Z","title":null,"venue":null,"work_id":"d82365e4-c9df-4e8d-89f6-c3dce244655c","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.019830Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:80e019e5785640955207ec66a4896d1aac2dcb1944086e586c87300ac4c145ac","observation_id":"50395c04-7f3c-4874-a8e1-c12984b57b65","resolution":{"observed_at":"2026-08-08T15:22:45.751506Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.08154","last_updated":"2024-05-13T19:52:16Z","snapshot_observed_at":"2026-08-09T16:27:46.187791Z","submitted_at":"2024-05-13T19:52:16Z","title":"LLM Theory of Mind and Alignment: Opportunities and Risks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.08154","snapshot_observed_at":"2026-08-08T15:22:45.024982Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.024982Z"},"links":{"cited_paper":"/paper/2405.08154","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:b8ac884c4a0fdad8b6ea80a0318c1a06628312374a5191fdf2b785fdadfd6164","observation_id":"3c8fa8f0-da92-4b99-8b1a-dc740b0cf524","resolution":{"observed_at":"2026-08-08T15:22:45.024982Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.18870","last_updated":"2024-05-31T12:45:50Z","snapshot_observed_at":"2026-07-06T18:21:48.356478Z","submitted_at":"2024-05-29T08:31:16Z","title":"LLMs achieve adult human performance on higher-order theory of mind tasks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.18870","snapshot_observed_at":"2026-08-08T15:22:45.030149Z","title":"O.; Keeling, G.; Baranes, A.; Barnett, B.; McKibben, M.; Kanyere, T.; Lentz, A.; y Arcas, B","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.030149Z"},"links":{"cited_paper":"/paper/2405.18870","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:8c43e1ee36303d1a749228c70df853607697895f0ebc5a4fc930f8008dcb72b0","observation_id":"dba2078f-4897-407b-926d-62dfa5b805c7","resolution":{"observed_at":"2026-08-08T15:22:45.030149Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.730645Z","title":null,"venue":null,"work_id":"50e524ab-19c4-4b20-a6a5-3f1d97576aa8","year":2019},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.035504Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:aabd5f861d13919ca1072fc9805119d33f455d07ab57a3399ed253c1d07e4054","observation_id":"61779350-4ab4-4f45-aaff-8a8e8e3a7d48","resolution":{"observed_at":"2026-08-08T15:22:45.735490Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.714740Z","title":null,"venue":null,"work_id":"7842d011-ebc0-451b-ade1-abe65485893d","year":2020},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.040599Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:6de31173d989f04e8a75fddfe5d736e092c3f617cae9327ae65b87168b76057f","observation_id":"1ed7a4d9-058e-4b99-94e5-7960f1058790","resolution":{"observed_at":"2026-08-08T15:22:45.719878Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.04360","last_updated":"2025-07-04T03:07:07Z","snapshot_observed_at":"2026-07-06T19:28:30.731752Z","submitted_at":"2024-10-06T05:02:23Z","title":"GenSim: A General Social Simulation Platform with Large Language Model based Agents","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.04360","snapshot_observed_at":"2026-08-08T15:22:45.045448Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.045448Z"},"links":{"cited_paper":"/paper/2410.04360","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:95908c287270e1f0a7b5ac33ec10a51db18ba50bb75738c4aa91495283a69ec1","observation_id":"7ceda1a8-ad88-497e-b6b1-651955a74066","resolution":{"observed_at":"2026-08-08T15:22:45.045448Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10248","last_updated":"2024-10-10T13:20:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-20T12:21:05Z","title":"Steering Language Models With Activation Engineering","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10248","snapshot_observed_at":"2026-08-08T15:22:45.050879Z","title":"M.; Thiergart, L.; Leech, G.; Udell, D.; Vazquez, J","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.050879Z"},"links":{"cited_paper":"/paper/2308.10248","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:ecf560eba37caf2842c9723ef3e6108d866e2cf1a5a55d089395c4f793972460","observation_id":"6855c7d3-7d6d-4126-952b-77ceacf3eeb9","resolution":{"observed_at":"2026-08-08T15:22:45.050879Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.08399","last_updated":"2023-03-14T13:47:26Z","snapshot_observed_at":"2026-08-09T17:51:01.952144Z","submitted_at":"2023-02-16T16:18:03Z","title":"Large Language Models Fail on Trivial Alterations to Theory-of-Mind Tasks","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.08399","snapshot_observed_at":"2026-08-08T15:22:45.056255Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.056255Z"},"links":{"cited_paper":"/paper/2302.08399","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:471ee5c7df56a1ab7129c67207ffa3d93dcf65b552a351d65ff81b82896050ae","observation_id":"1649c9ec-82ab-4eb1-a22c-e240d09e33e8","resolution":{"observed_at":"2026-08-08T15:22:45.056255Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07358","last_updated":"2025-02-06T20:58:43Z","snapshot_observed_at":"2026-08-06T05:53:39.626163Z","submitted_at":"2024-06-11T15:26:57Z","title":"AI Sandbagging: Language Models can Strategically Underperform on Evaluations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07358","snapshot_observed_at":"2026-08-08T15:22:45.061524Z","title":"F.; and Ward, F","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.061524Z"},"links":{"cited_paper":"/paper/2406.07358","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:55ef58f99837480dd10c2fa14a7898c9b859a1edae8a55d8191d7416e42dc0f1","observation_id":"715b3a84-e7c1-4a58-8a95-fe7eb84c2449","resolution":{"observed_at":"2026-08-08T15:22:45.061524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.698118Z","title":null,"venue":null,"work_id":"c180a82b-79d6-423c-be69-5486f9760fb1","year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.067050Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:a44a345873d40cca5c95089be39d03cd0665c1b84288f0cd630fbc300f7b197b","observation_id":"66fc2730-c3a7-4b5c-8c4c-5e0cbf9f2459","resolution":{"observed_at":"2026-08-08T15:22:45.703545Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.681853Z","title":"P.; and Morency, L.-P","venue":null,"work_id":"014366d4-0747-4f9b-aec1-8af7c3b23b40","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.071851Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:875a6c060c4ddee3c0439363317c3f38b6fc513391ae9d998aa01ca177c24452","observation_id":"da5ba737-5f0e-48bc-95d7-8df83b7aaaca","resolution":{"observed_at":"2026-08-08T15:22:45.686701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.662960Z","title":null,"venue":null,"work_id":"0c7a5cfb-aef4-421b-9b16-fcb3a4a7ca88","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.077014Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:bc55d8a58f9f1435beeff7267ae978706f709d1137ca939431969deff44cea9b","observation_id":"62454549-30fa-4401-b93a-ebd7638e1291","resolution":{"observed_at":"2026-08-08T15:22:45.669506Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.644157Z","title":null,"venue":null,"work_id":"fdacd8ac-3b8d-47f7-ac02-468c636f4814","year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.082337Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:26c914bed008643c4817735c3d0da7479bb2bf3bb7a9d0548c926ed59729388e","observation_id":"a8c3addc-f454-4b1d-981b-340b4fffbb23","resolution":{"observed_at":"2026-08-08T15:22:45.649870Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.625615Z","title":null,"venue":null,"work_id":"7a09beb7-a368-4905-9ef8-9878a5a952be","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.087249Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:7e6250c209b49db7b89ebbbd6f39f6a9d1d107ee516f33534267c064952c6452","observation_id":"a0d907c1-d251-469c-b824-fdd4f6acdd2d","resolution":{"observed_at":"2026-08-08T15:22:45.630885Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.608891Z","title":null,"venue":null,"work_id":"d3338d72-758d-4b06-8467-6ab579227dd9","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.092277Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:1a40af27e93576a5c99f33115f528f04ed48085ba0c87d9cec6fd6a47ba2d6b7","observation_id":"98d547f2-6be6-438e-84a3-e244994410cb","resolution":{"observed_at":"2026-08-08T15:22:45.613924Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.00332","last_updated":"2024-11-22T22:27:49Z","snapshot_observed_at":"2026-07-06T18:08:04.815730Z","submitted_at":"2024-05-01T05:52:05Z","title":"A Careful Examination of Large Language Model Performance on Grade School Arithmetic","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.00332","snapshot_observed_at":"2026-08-08T15:22:45.097354Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.097354Z"},"links":{"cited_paper":"/paper/2405.00332","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:1f48053c9b65b7ae8d19580dcdab165c9fa2ca96e63fdf30407276c05191bfad","observation_id":"0b58a3d6-b9b1-42c3-896e-a3f41e1f58e0","resolution":{"observed_at":"2026-08-08T15:22:45.097354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.592067Z","title":null,"venue":null,"work_id":"2b04f5e0-6f94-4fca-9ef1-acb3a816d517","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.102646Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:dd6fe2092e0373934b6b4bc18f650e96b1e9c46fbaeeb547abe97cbe1ddb7a52","observation_id":"658e47e3-2348-49f2-95c5-1b1e74bab22a","resolution":{"observed_at":"2026-08-08T15:22:45.597345Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T15:22:45.573772Z","title":null,"venue":null,"work_id":"270dafd7-2b56-4b3e-bd62-5a15e5b629ce","year":2024},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.107446Z"},"links":{"citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:ceea2fcc582a0b12ac2e82bfe86a983ae7a72332f6152242acf4a0b577ebfc2c","observation_id":"0bd53ef2-e478-4077-8ee6-18c833023895","resolution":{"observed_at":"2026-08-08T15:22:45.580230Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01405","last_updated":"2025-03-03T06:14:14Z","snapshot_observed_at":"2026-07-06T16:26:38.284922Z","submitted_at":"2023-10-02T17:59:07Z","title":"Representation Engineering: A Top-Down Approach to AI Transparency","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01405","snapshot_observed_at":"2026-08-08T15:22:45.112343Z","title":"J.; Wang, Z.; Mallen, A.; Basart, S.; Koyejo, S.; Song, D.; Fredrikson, M.; Kolter, J","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-08T15:22:45.112343Z"},"links":{"cited_paper":"/paper/2310.01405","citing_paper":"/paper/2502.06470"},"observation_digest":"sha256:0877db7f64dc0aa8ff3737604b7845306106e6a6e054d33eb41208f096675435","observation_id":"fea45c52-340f-4f29-92e7-1f2acde4c355","resolution":{"observed_at":"2026-08-08T15:22:45.112343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.06470","last_updated":"2025-02-10T13:50:25Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T05:33:06.269788Z","submitted_at":"2025-02-10T13:50:25Z","title":"A Survey of Theory of Mind in Large Language Models: Evaluations, Representations, and Safety Risks"},"reference_resolution":{"displayed":64,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":53,"verified_exact":0,"verified_fuzzy":11},"total_outbound_references":64},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 64 of 64 outbound references and 3 inbound Pith citation observations for arXiv:2502.06470."}