{"as_of":"2026-08-10T08:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fff1d90dc57d0aa7b1e8d30fb662c17278b84e0c5b714abd798df4d247ca4f77","coverage":[{"denominator":68,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":68,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T21:43:42.163889Z","state":"measured"},{"denominator":68,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":68,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2507.05266/citation-record","integrity":"/paper/2507.05266/integrity","json":"/paper/2507.05266/citation-record.json","paper":"/paper/2507.05266"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:31.780206Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:31.780206Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:b1876ea0bedaf2ac2fe79fa82770315d09c29f02a43f659d08f32f19655e5c87","observation_id":"302b938b-9e5d-4e70-ab47-d0fb66818a15","resolution":{"observed_at":"2026-08-06T21:43:31.780206Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-06T21:43:31.875312Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:31.875312Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:a54435b1fca7bff7217e3365a48a5a996e44902872461433ed56df93d12448ab","observation_id":"5a840d97-3a01-4e97-833e-d55072cc25eb","resolution":{"observed_at":"2026-08-06T21:43:31.875312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:31.999784Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:31.999784Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:fdd3221171e36e8ac542825692510d3a374ac2cd7d4d82e5ea7142844c5fe9f3","observation_id":"57efcae8-9538-4c16-a50f-102830ffb2a3","resolution":{"observed_at":"2026-08-06T21:43:31.999784Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:32.114826Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:32.114826Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:19bc650bad83c8a3fbca94fb4b48ce2f0424fff2d5f89a6107a4ad4650ca9fa4","observation_id":"2dc54533-f717-4608-b31b-3dba1fde6e29","resolution":{"observed_at":"2026-08-06T21:43:32.114826Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:52.296754Z","title":null,"venue":null,"work_id":"38b56e00-cce0-4b1c-a6e5-0047dfe0e627","year":2005},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:32.258580Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:0bbd02890676927edbd1e52fc9e9148421ab36229b4b2305c6f38346d7e7abbc","observation_id":"10879fc8-cf60-4d2e-a67f-b9dcbfe83cac","resolution":{"observed_at":"2026-08-06T21:43:52.458598Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:51.973614Z","title":null,"venue":null,"work_id":"a3f0822f-b7ca-46ce-b033-6f3e798a00c1","year":2006},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:32.376986Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:382011297932de0705bfe7fe095a0c89df0dbbdfbf2022b65d2df98777774010","observation_id":"8974b10e-f782-4c6d-9f3f-1cb23bccc5fd","resolution":{"observed_at":"2026-08-06T21:43:52.151476Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:32.573252Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:32.573252Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:c61a70319fc6cc2030f784b5c0078c1186f1853e793d48cbf16927cc934c50fb","observation_id":"6e57049b-4de5-4a73-b7ed-a5192107b069","resolution":{"observed_at":"2026-08-06T21:43:32.573252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:51.699585Z","title":null,"venue":null,"work_id":"2caacbb1-f5a6-472d-8e0f-95641a2ee26e","year":2017},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:32.738990Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:1f3e8c301495cb8c3d3465284a8909aa67ce794684bb2dae219af00ec397657e","observation_id":"01c74ad5-effb-4355-9775-7a999a9e6d33","resolution":{"observed_at":"2026-08-06T21:43:51.843815Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:32.886441Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:32.886441Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:d6468370801d0c5feb8b1012dd7d63d8de9199a5144766eb576addf7355de668","observation_id":"d8d53ded-173b-47a0-a3e6-520eb5fb0c35","resolution":{"observed_at":"2026-08-06T21:43:32.886441Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:51.416262Z","title":null,"venue":null,"work_id":"d998fad8-9db7-4d4b-9261-0e6b8669184e","year":2010},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:33.047126Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:1de43361d7837a50a1d07a86a65b29a6342503abb2e8e3a07da08604864fa3fa","observation_id":"3f64f436-4568-43ac-aecc-20b0dc93b1a0","resolution":{"observed_at":"2026-08-06T21:43:51.533764Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:33.188732Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:33.188732Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:8d2ee395c124d4dd2442a7b3cedb5edf81f71dc4a9fbfd19ada3ebb7b3f0ef6b","observation_id":"66a0cc6b-f77d-4826-9a2e-7b760a2a1bd6","resolution":{"observed_at":"2026-08-06T21:43:33.188732Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:51.115765Z","title":null,"venue":null,"work_id":"453410d7-7eba-4ea5-a92c-e2ca560f0c95","year":2010},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:33.390981Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:c13eac02f952535b512408e4c77fe4a550e84a83ddbf6aef3744ecb891286df0","observation_id":"1eae7162-38e5-47b4-9009-0d5023475ea6","resolution":{"observed_at":"2026-08-06T21:43:51.250174Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.17161","last_updated":"2025-05-26T17:16:45Z","snapshot_observed_at":"2026-08-09T18:26:12.869738Z","submitted_at":"2025-01-28T18:59:44Z","title":"SFT Memorizes, RL Generalizes: A Comparative Study of Foundation Model Post-training","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.17161","snapshot_observed_at":"2026-08-06T21:43:33.539108Z","title":"Le, Sergey Levine, and Yi Ma","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:33.539108Z"},"links":{"cited_paper":"/paper/2501.17161","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:50a7504d99f1c338330a5ff96c35533ad758c1f89f4242217ebc31800c6f0440","observation_id":"fa79763b-22ed-41b7-a961-63e28d4b1669","resolution":{"observed_at":"2026-08-06T21:43:33.539108Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-07T01:45:38.840969Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-06T21:43:33.704744Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:33.704744Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:3f5b3673d53c9e782db320b0f167f565d7db15d253ba9495cc0d04ca1dc472ed","observation_id":"bbe85ef2-aba5-4247-83d3-15b00f701a1b","resolution":{"observed_at":"2026-08-06T21:43:33.704744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:50.819634Z","title":null,"venue":null,"work_id":"ebf4ac0c-f5a1-48af-94fe-2f8c41e70ef5","year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:33.830889Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:ee1e912bceb637e3457d8237360bc5b4a11658eaaac3b352ce7304a57b324166","observation_id":"705695e9-9f3a-4a10-aeff-888991ad8aa7","resolution":{"observed_at":"2026-08-06T21:43:50.952985Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:50.533520Z","title":null,"venue":null,"work_id":"520b95a5-cb7f-4a1a-87cf-9d0827678652","year":2000},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:33.956895Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:82fb17d41b6490c14f205f8d0c3eb1a97d38d21fd43b7c9a85e199d5cc8db721","observation_id":"95ad60b4-0b66-483d-9a5e-4501433dc73b","resolution":{"observed_at":"2026-08-06T21:43:50.673809Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:50.337954Z","title":null,"venue":null,"work_id":"ab308944-8617-4f66-88df-77a8069b1bbf","year":2006},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:34.096827Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:e8f52d12f52926e601fb69122da379bef856ded7d93a9f7da207509f99b25de0","observation_id":"3b0cb5c1-c708-4a80-b7be-469ea929e3ee","resolution":{"observed_at":"2026-08-06T21:43:50.419361Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-06T21:43:34.302368Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:34.302368Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:be4992a47bd70eda41529d25b6885aa550fae473afc54e7aed0e84dfbc8d97db","observation_id":"807f6f38-74f8-48b6-9c23-fbd04f1143d1","resolution":{"observed_at":"2026-08-06T21:43:34.302368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:50.141929Z","title":null,"venue":null,"work_id":"87c002cf-8b45-48f2-a410-6e51cddf5eb7","year":2006},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:34.460806Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:05a3afff005f7c567104b02e6fe58838f73ae6fcffdf8a3fc69c82f844c05ff9","observation_id":"4c2fe7a5-9d42-48a0-a81b-7f7164fb809e","resolution":{"observed_at":"2026-08-06T21:43:50.242081Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:49.866805Z","title":null,"venue":null,"work_id":"2a869e40-ac21-41f6-bd19-cae85bf50cfe","year":2016},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:34.645032Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:bc20507a5634f2a02a3e075ae817b2f6dc9f9abd3f1e2a26e8ccb7ecc82c0f8e","observation_id":"c3602873-a393-4ebd-9cca-f3786a0a1db6","resolution":{"observed_at":"2026-08-06T21:43:50.049174Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:49.569672Z","title":null,"venue":null,"work_id":"1fd13edd-6f2c-47c1-9473-7c5bdd71bacf","year":2020},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:34.815490Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:3a682ffe191444f7fb661ac566eb640b8ed430321a19b4b091c9d19782297f90","observation_id":"678f20e6-1941-4621-8df4-ba5e34ad7027","resolution":{"observed_at":"2026-08-06T21:43:49.698061Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.19736","last_updated":"2023-11-25T17:35:12Z","snapshot_observed_at":"2026-07-06T16:40:35.590974Z","submitted_at":"2023-10-30T17:00:52Z","title":"Evaluating Large Language Models: A Comprehensive Survey","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.19736","snapshot_observed_at":"2026-08-06T21:43:34.962660Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:34.962660Z"},"links":{"cited_paper":"/paper/2310.19736","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:6487becd4a23d14a6cb46fbadf8e6d8669a195560f7cdf4d00aa046389df7d2b","observation_id":"528253e7-3801-47df-a54b-7e0035840e8d","resolution":{"observed_at":"2026-08-06T21:43:34.962660Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:49.250802Z","title":null,"venue":null,"work_id":"e59e47bd-09a6-4bf6-a2eb-a54cea05f406","year":2006},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:35.108577Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:21327b682b0ea195d690f63b3d5e1d9d8d73a5628f4d1cc519592344167b3038","observation_id":"1c48d83d-abd6-424c-8494-15464844d20c","resolution":{"observed_at":"2026-08-06T21:43:49.432484Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:49.043599Z","title":null,"venue":null,"work_id":"6aa4a8e4-1890-4287-bbdb-08cce6c22c69","year":2015},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:35.348888Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:6f08cb44af5af8f71d249e577f4859ed646beac949245b5d86a31121fe3820a5","observation_id":"cbcdbf1d-310b-4a80-a839-33bc7136cb65","resolution":{"observed_at":"2026-08-06T21:43:49.145384Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:35.518142Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:35.518142Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:d8cba89a39e09b40fb7e83050c64d30ea0d49425d31778d64d95cd76da91bc80","observation_id":"fc879cd2-5b81-46d5-99b3-f9ba5aaac410","resolution":{"observed_at":"2026-08-06T21:43:35.518142Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-08-09T10:28:06.906299Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-08-06T21:43:35.760263Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:35.760263Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:de7f950686a6aa59aeb730f2108f0df178fa60dbc4a156483e689828e29b7fa6","observation_id":"84fe2f46-16c0-4b57-8df9-f2cfeea8c4ab","resolution":{"observed_at":"2026-08-06T21:43:35.760263Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:48.716166Z","title":null,"venue":null,"work_id":"f4208ea6-365e-412b-b3db-7dc11ccf97c2","year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:35.962034Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:135f62256a21ac8f00ca26e71666b98dd51cb7e1cae89e5c3cff86876b08a912","observation_id":"f1e1386a-ead7-4b1d-8ad4-0dbe18ca923b","resolution":{"observed_at":"2026-08-06T21:43:48.855108Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:36.088502Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:36.088502Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:e543f28b20beeeafdc6944fde9cfc4cea982f1f23a24ab9414c9b1baa5afb602","observation_id":"77459335-f923-4f6a-8c3a-38edff5cbd33","resolution":{"observed_at":"2026-08-06T21:43:36.088502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:48.477240Z","title":null,"venue":null,"work_id":"815d383e-b0d0-407e-8546-404b77124480","year":2002},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:36.226337Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:2aa317f8c6e26db3bfc9afd4506f7ad4f4fa3050d42beba17b2740c2013e7853","observation_id":"d5079bbc-8f53-448a-adf6-47548034f98f","resolution":{"observed_at":"2026-08-06T21:43:48.554427Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:48.210359Z","title":null,"venue":null,"work_id":"4f6e787c-90a3-4ea8-9305-c1189c28d773","year":2022},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:36.406874Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:c1ef2586edfe3c6e888f580a71cca06974fa60228500d4d3a3a84f99f91af56e","observation_id":"2587c13a-aa98-4a9e-9db5-acafcc25eea8","resolution":{"observed_at":"2026-08-06T21:43:48.313627Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.07870","last_updated":"2023-11-07T16:28:33Z","snapshot_observed_at":"2026-07-06T15:54:25.070719Z","submitted_at":"2023-07-15T19:04:33Z","title":"Large Language Models as Superpositions of Cultural Perspectives","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.07870","snapshot_observed_at":"2026-08-06T21:43:36.579640Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:36.579640Z"},"links":{"cited_paper":"/paper/2307.07870","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:4a3236ef04e2105d9866e9734f90497adf6218e85b28667cb1f43f1b91cc4173","observation_id":"7bdf0408-6c0d-48bf-977e-d323094bd47b","resolution":{"observed_at":"2026-08-06T21:43:36.579640Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:47.997518Z","title":null,"venue":null,"work_id":"b2a8f2f1-5b94-476f-a029-93c13ec146e7","year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:36.780484Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:d83cb90ff8cfd4abf4ae36072b5ae3cb9eedcc7d6c0f95282834d84ca94f5db4","observation_id":"80cc9615-efe8-448c-95d2-4e274449483e","resolution":{"observed_at":"2026-08-06T21:43:48.069470Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:47.769743Z","title":null,"venue":null,"work_id":"12f788ec-1476-46ef-a282-e86def8c8762","year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:36.953376Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:cacbbc38ca0d4e87ea8eac8e507ba40bf6926bc9c03f110383c07da400834948","observation_id":"11eeba2d-8f91-4911-a173-b4a508117b1a","resolution":{"observed_at":"2026-08-06T21:43:47.902732Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.10199","last_updated":"2024-08-20T06:53:45Z","snapshot_observed_at":"2026-08-09T18:07:38.663665Z","submitted_at":"2024-04-16T00:50:43Z","title":"CULTURE-GEN: Revealing Global Cultural Perception in Language Models through Natural Language Prompting","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.10199","snapshot_observed_at":"2026-08-06T21:43:37.167104Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:37.167104Z"},"links":{"cited_paper":"/paper/2404.10199","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:736f426f502ec7ad2e098b919a21b7a3b957e97c31306bd37b33d1264d44f2ef","observation_id":"465c50b7-af95-4c84-9083-f0f3c6ad6568","resolution":{"observed_at":"2026-08-06T21:43:37.167104Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:47.497963Z","title":null,"venue":null,"work_id":"f2cefcdb-4525-4212-8532-a44daf01ab47","year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:37.358895Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:c013b60ebc790aeb203c0120b4571f5c7ca72cc7ffc2a0ccb4c6c810e299bd18","observation_id":"f1d3606a-65c6-4009-a9f1-367353ba9ead","resolution":{"observed_at":"2026-08-06T21:43:47.615443Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.10149","last_updated":"2023-10-27T09:11:50Z","snapshot_observed_at":"2026-07-06T15:17:48.587665Z","submitted_at":"2023-04-20T08:16:07Z","title":"Is ChatGPT a Good Recommender? A Preliminary Study","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.10149","snapshot_observed_at":"2026-08-06T21:43:37.525608Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:37.525608Z"},"links":{"cited_paper":"/paper/2304.10149","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:67c613fe43d42a89d21a5302a8c1d6562ddea8628d87d354c52202484d588bdb","observation_id":"6a3eae84-0260-4fd0-821a-3623203ddb66","resolution":{"observed_at":"2026-08-06T21:43:37.525608Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:47.290557Z","title":null,"venue":null,"work_id":"141e4da3-6b78-4a41-a52b-6ba8dfffb9ca","year":2019},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:37.644108Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:5f0c22e4a5a9ac26196dfe5fbcead917fc4df65c78323ebb5a0eb4a3e29a7eeb","observation_id":"d7b35eae-0f60-46cc-a997-1c1aae56d2f0","resolution":{"observed_at":"2026-08-06T21:43:47.369504Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.emnlp-main.884","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":null,"venue":null,"work_id":"f64a9b15-3b80-4dc1-9ea9-914e8540ecdf","year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:37.764807Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:667d299ddc929506f569dd6d9d7e08f2f737cb3d78438a20d5c70e43ae5777c5","observation_id":"e1e1076a-28ec-41cb-ba97-26630c9cbda7","resolution":{"observed_at":"2026-08-06T21:43:42.403995Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:37.901202Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:37.901202Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:e1b8787b57ccd5421a150c6099e7e11a69d8ea46c17413a71f70cc3066004cbf","observation_id":"ee09ef14-54c4-46d8-889d-3c3f51247fca","resolution":{"observed_at":"2026-08-06T21:43:37.901202Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:47.031021Z","title":null,"venue":null,"work_id":"b89e6aa0-d337-4da9-87e0-72fa7828084f","year":2015},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:38.022415Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:c97cf665454585242d0bd19f4492dd9ff1b68ab8f0a3949950f1e2b0647460eb","observation_id":"dd3faa59-e19b-491d-8af2-570c6bc63660","resolution":{"observed_at":"2026-08-06T21:43:47.152472Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:38.257084Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:38.257084Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:bb2a43e72678fc540b0cfcebbe816222095d899b3eb232b7b56ded6260f8378d","observation_id":"4bda4121-31bf-4c21-ba78-9b4c6e38d718","resolution":{"observed_at":"2026-08-06T21:43:38.257084Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:46.762709Z","title":null,"venue":null,"work_id":"f2ae5f60-bec7-4470-8a2c-4b542db467b4","year":2014},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:38.402880Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:05ba7cb0a49a6d4f2c71ee0078c015cfd04947ea1814713ff57e43f0ef15ea6b","observation_id":"483a487d-46e1-4255-b7b5-33fa0b3b4137","resolution":{"observed_at":"2026-08-06T21:43:46.876322Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:38.582482Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:38.582482Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:cffc754c23f6660723786ef072fc808020536beb9b6c5820cd56c62c5f889b3d","observation_id":"a924e22a-9008-4b74-82f2-cb3079ae35cf","resolution":{"observed_at":"2026-08-06T21:43:38.582482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:46.457483Z","title":null,"venue":null,"work_id":"adaf1e9b-74a9-4591-b63a-e851fa6e0aa8","year":2014},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:38.788471Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:e303855f3461b94ee612ad515972a3227b5c8e1bbfa363ce376a54b0b0dccfdc","observation_id":"174fdf78-60e7-45be-a6a3-2c8189526d83","resolution":{"observed_at":"2026-08-06T21:43:46.598042Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:46.205915Z","title":null,"venue":null,"work_id":"b45e3af5-f045-423c-9751-1ce10e96458e","year":2016},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:38.908575Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:c43c2175f6f2f5dfe35c3aa078933272e6b2424e26a2606136363094d8ec6470","observation_id":"dfabe30b-ac9d-41dc-821d-f97946597e2e","resolution":{"observed_at":"2026-08-06T21:43:46.269639Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:46.016394Z","title":null,"venue":null,"work_id":"08558afd-6717-4a9f-a349-356bc1568a64","year":1999},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.021122Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:431d9fe49696a9d6b39c0ae2f82ac498a5a146a6ef7b10373a3b539a0245e4ae","observation_id":"f49592ae-1fbe-42a4-9a30-27f8ccbe6ee4","resolution":{"observed_at":"2026-08-06T21:43:46.121174Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:45.677816Z","title":"Practices for governing agentic ai systems","venue":null,"work_id":"c9e1d7a8-d7aa-48be-9c77-0b7c33a4bc10","year":null},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.155956Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:02744717345b3e5061e6b3d1832faad3b8e1252f59c96aad7742359b04ff8ad6","observation_id":"dc354d82-956e-46b0-bd62-0867290ba786","resolution":{"observed_at":"2026-08-06T21:43:45.843799Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:45.501017Z","title":null,"venue":null,"work_id":"49bd4934-d262-4596-9ce1-c097c14d734d","year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.303135Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:9f539faae9c4130ee6e2f403bc46dfd48c838b4b6375182b8ccf8428e3c6cabe","observation_id":"9c23ad6c-3e8f-463d-a651-52d33676fa7e","resolution":{"observed_at":"2026-08-06T21:43:45.585026Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:45.182785Z","title":null,"venue":null,"work_id":"fa31bcce-be13-4f9e-84b2-c217cb8f0638","year":2010},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.468704Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:4306efc0920aae602e470d35ab73edf8d2924a5885d7e4b9a96ba32694fbf6cb","observation_id":"f4c85996-2a00-414a-a2b9-e99279b16178","resolution":{"observed_at":"2026-08-06T21:43:45.357904Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:44.913827Z","title":null,"venue":null,"work_id":"f032a9c6-12a5-4d1d-9432-4d1f0d32aba1","year":2018},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.557291Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:f1e4ae80ca2058487ddb2202b36113ef90c6c937db5ae0c485fa1d7be53a1297","observation_id":"945af8fe-5db2-4067-9638-7ae1ff1bdfc6","resolution":{"observed_at":"2026-08-06T21:43:45.052750Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.13356","last_updated":"2023-10-07T09:14:43Z","snapshot_observed_at":"2026-07-06T16:22:46.711142Z","submitted_at":"2023-09-23T12:17:10Z","title":"Probing the Moral Development of Large Language Models through Defining Issues Test","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.13356","snapshot_observed_at":"2026-08-06T21:43:39.663308Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.663308Z"},"links":{"cited_paper":"/paper/2309.13356","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:4597b343bff36877221d3268e338ae412929079ff1499118e74be391959de0f8","observation_id":"43a11978-c6c6-4a99-a9b8-e1fef89fd303","resolution":{"observed_at":"2026-08-06T21:43:39.663308Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:44.580801Z","title":null,"venue":null,"work_id":"703755e0-a815-4248-91d2-81c59a3e8758","year":2020},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.872495Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:e786114106e8f71bb07664379089bcc9ccb47a7144ca0097146953af1f2740ce","observation_id":"256934ef-e9f2-42ac-8350-ff31db532369","resolution":{"observed_at":"2026-08-06T21:43:44.713924Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:44.321046Z","title":null,"venue":null,"work_id":"192a79a8-3c58-453e-8f35-818a06234f88","year":1948},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:39.998506Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:56ff265230e1855cf8482310dcca7da975b98825b0c1e21b2eab375de841a077","observation_id":"53d6a2d8-8d93-4811-8cc3-c360f9b7d384","resolution":{"observed_at":"2026-08-06T21:43:44.471577Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:44.067147Z","title":null,"venue":null,"work_id":"8bce8b68-936c-4d99-86b6-abd6684b88cd","year":1950},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:40.093752Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:8ba79619172ecd9a59387f201092b821bd4cbae60cc466415d268b04256a2367","observation_id":"58bd3c7b-2ad9-4907-a3d2-a3f91ab94f56","resolution":{"observed_at":"2026-08-06T21:43:44.172884Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:43.907426Z","title":null,"venue":null,"work_id":"4d790cb9-c370-43c3-8449-17e1b764961c","year":2007},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:40.265620Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:f9dd601f0087ec328bafb1b22090f399ba3dc5c07c450baa4b2f9050fd02fde6","observation_id":"69e4078e-823e-4611-b3b9-b11b2b8e7e3e","resolution":{"observed_at":"2026-08-06T21:43:44.013805Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:43.712566Z","title":null,"venue":null,"work_id":"cd46a050-1251-4468-a049-8a7b38568767","year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:40.396998Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:e391881a3dec09c03b0f2f380aa34851168b2157dd813dff2a4d870cdee5662b","observation_id":"49f5bb56-d43a-49f8-906c-bec2b54ce04d","resolution":{"observed_at":"2026-08-06T21:43:43.833508Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.04325","last_updated":"2024-06-04T22:09:46Z","snapshot_observed_at":"2026-08-09T20:41:07.756820Z","submitted_at":"2022-10-26T00:28:40Z","title":"Will we run out of data? Limits of LLM scaling based on human-generated data","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.04325","snapshot_observed_at":"2026-08-06T21:43:40.536802Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:40.536802Z"},"links":{"cited_paper":"/paper/2211.04325","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:8b633fc66db7255b987fa32e9691ab655dea4f26a161ecd7d0df3f56ad156965","observation_id":"d0faab4c-2961-44a0-84f7-399bba3419db","resolution":{"observed_at":"2026-08-06T21:43:40.536802Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:40.701264Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:40.701264Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:08e138ce2d948de41c224dbae05f2e0ccf7df35b9238cc67e915951f6ad1bbd4","observation_id":"79fd0586-94f5-4ab4-b191-c0f5a0451bee","resolution":{"observed_at":"2026-08-06T21:43:40.701264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1804.07461","last_updated":"2019-02-22T23:53:34Z","snapshot_observed_at":"2026-07-06T06:34:26.609892Z","submitted_at":"2018-04-20T06:35:04Z","title":"GLUE: A Multi-Task Benchmark and Analysis Platform for Natural Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1804.07461","snapshot_observed_at":"2026-08-06T21:43:40.875956Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:40.875956Z"},"links":{"cited_paper":"/paper/1804.07461","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:5c483930696992939bc3549ce6e307de5817ec2554e61c18eb6af917584714fe","observation_id":"75fba999-5051-489a-9d51-99fa4ed0ba6f","resolution":{"observed_at":"2026-08-06T21:43:40.875956Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:43.512238Z","title":null,"venue":null,"work_id":"b83093e7-0b36-43ff-952c-773215ae2cd0","year":2011},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.011460Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:ca2c3c0dcafdd906977a08195e58b88582efe17406ef3b45267e62830d7930d2","observation_id":"692909f8-f2f6-44e4-873e-4b55deb4c7a7","resolution":{"observed_at":"2026-08-06T21:43:43.574933Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:43.246300Z","title":null,"venue":null,"work_id":"4f169081-1e6f-4ab5-8b9a-00b27ff253a9","year":1953},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.173567Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:5216d08d5ccc84e1dbe1bcef029deffff5aacf94b3d5b561b6951b63022a874a","observation_id":"a39066c8-bd37-441e-b1d0-e23805d7155a","resolution":{"observed_at":"2026-08-06T21:43:43.368445Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:41.308124Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.308124Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:610b1c97d69f55c4bc436cb4f403453f3db88ce43cf5ebc75066eb2c421bc2ae","observation_id":"d796a52b-3899-449a-95f2-40e0b4918faa","resolution":{"observed_at":"2026-08-06T21:43:41.308124Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.00027","last_updated":"2025-07-13T02:13:14Z","snapshot_observed_at":"2026-07-06T19:43:09.889618Z","submitted_at":"2024-10-29T04:01:11Z","title":"Personalization of Large Language Models: A Survey","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.00027","snapshot_observed_at":"2026-08-06T21:43:41.450167Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.450167Z"},"links":{"cited_paper":"/paper/2411.00027","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:20dd7ab25f721ec6053229672dbe1121ac8ed89845506a36e3204d8084467c6e","observation_id":"b44d338d-b7d4-4e3f-807f-f87cd0a4bbbe","resolution":{"observed_at":"2026-08-06T21:43:41.450167Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:42.845784Z","title":null,"venue":null,"work_id":"ff0eb6e3-57c1-4f88-bc17-496f03d03cda","year":2024},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.596384Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:8eb95a558cacda3379372c8cd89033641ce6aad0c907126d2a92a1737ff4144a","observation_id":"586ccf5a-1ff2-45a7-a63b-c815dfbbe265","resolution":{"observed_at":"2026-08-06T21:43:43.065499Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.01964","last_updated":"2023-11-03T14:59:54Z","snapshot_observed_at":"2026-08-06T10:56:05.075839Z","submitted_at":"2023-11-03T14:59:54Z","title":"Don't Make Your LLM an Evaluation Benchmark Cheater","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.01964","snapshot_observed_at":"2026-08-06T21:43:41.764235Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.764235Z"},"links":{"cited_paper":"/paper/2311.01964","citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:84d7ac2f8fae7ad47e176f3e738a66a3bd397aabd82de7f73933a06e352d1acc","observation_id":"5ecc2e53-7f59-435d-afc4-b5bb1194e4e6","resolution":{"observed_at":"2026-08-06T21:43:41.764235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:42.603487Z","title":null,"venue":null,"work_id":"2c897ed4-cf59-4917-b6ad-64813aebcf34","year":2007},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.856323Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:1657e3fb5ab0edad84be844d5ebd3ff322ade4ddda166375e5e4a69f742c232c","observation_id":"246ffb70-2d8d-4b04-824c-0f31d2261f69","resolution":{"observed_at":"2026-08-06T21:43:42.716986Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:41.969044Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:41.969044Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:775f6fd926932800476cff48370062c27d5f7ef0d80a6bca6405b429ce148117","observation_id":"289b7e1d-1326-4085-a0d8-c042bd2c8654","resolution":{"observed_at":"2026-08-06T21:43:41.969044Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:43:42.163889Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-06T21:43:42.163889Z"},"links":{"citing_paper":"/paper/2507.05266"},"observation_digest":"sha256:c0d97eb882b825780e5f3b75e54601ce5635425e3fbb78c865a2dce41b788f40","observation_id":"15a5a920-5113-4418-b910-a9df9eb18613","resolution":{"observed_at":"2026-08-06T21:43:42.163889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.05266","last_updated":"2025-06-30T06:14:32Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T18:08:17.943936Z","submitted_at":"2025-06-30T06:14:32Z","title":"User Behavior Prediction as a Generic, Robust, Scalable, and Low-Cost Evaluation Strategy for Estimating Generalization in LLMs"},"reference_resolution":{"displayed":68,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":66,"verified_exact":1,"verified_fuzzy":1},"total_outbound_references":68},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 68 of 68 outbound references and 0 inbound Pith citation observations for arXiv:2507.05266."}