{"as_of":"2026-08-17T07:36:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b02794f11dc69b83b5068085a85f9b31700c6e06f8c8b918f071d873f7cc333a","coverage":[{"denominator":41,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":41,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T10:53:10.815652Z","state":"measured"},{"denominator":42,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":42,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:03:57.161602Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T13:03:58.532424Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"cited_work":{"arxiv_id":"2412.16318","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.16318","snapshot_observed_at":"2026-08-07T13:03:58.532424Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","venue":"cs.LG","work_id":"9a7914cb-9cb1-430d-bb15-cae08751b2a5","year":2024},"citing_paper":{"arxiv_id":"2505.23124","last_updated":"2025-08-01T22:17:13Z","snapshot_observed_at":"2026-08-13T19:17:00.975743Z","submitted_at":"2025-05-29T05:46:01Z","title":"Learning to Incentivize in Repeated Principal-Agent Problems with Adversarial Agent Arrivals","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T13:03:57.161602Z"},"links":{"cited_paper":"/paper/2412.16318","citing_paper":"/paper/2505.23124"},"observation_digest":"sha256:65091455a82ca23fc14f56de49358756b1d27ed54c4e1e7c96fce22df8973315","observation_id":"af64217d-881b-43b3-a166-02718cfa7c29","resolution":{"observed_at":"2026-08-07T13:03:58.582546Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.16318/citation-record","integrity":"/paper/2412.16318/integrity","json":"/paper/2412.16318/citation-record.json","paper":"/paper/2412.16318"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.676025Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.676025Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:74882765978d05a38a088fe1fc766db51cd1a2ea35555995f7ead33aabf725aa","observation_id":"9cb9f57f-aa13-4243-a2f1-7dda7878e0ff","resolution":{"observed_at":"2026-08-11T10:53:10.676025Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.225414Z","title":"Toward a theory of discounted repeated games with imperfect monitoring","venue":null,"work_id":"d6b12aa8-8f5f-491d-b821-f763cb7be510","year":1990},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.680937Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:03cf521a5a56235417d25cf1c88c6fb443d0056cd95a8aa5d7a160d75f6ee479","observation_id":"4deca98b-b851-4643-8e4a-e4ba9e7cbd7f","resolution":{"observed_at":"2026-08-11T10:53:11.229931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.214775Z","title":"and Erev, I","venue":null,"work_id":"9dca315e-a920-499c-bf86-c2313ecac1c1","year":2003},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.685124Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:72bc2377480a7f8b6f2b80c93e702ba8b4b1b349f488d08ac3ce74bd7c995b02","observation_id":"efed639b-3301-4267-898c-000eb4156430","resolution":{"observed_at":"2026-08-11T10:53:11.218736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.688856Z","title":"Principal-agent reward shaping in mdps","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.688856Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:c6b6f31ba7512290c5b8f210bbfd4f6bfbcb990c29a30a0ceb5109ff7638cc40","observation_id":"09aa1a15-8294-4cb2-af91-0d97b1163f0e","resolution":{"observed_at":"2026-08-11T10:53:10.688856Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.692410Z","title":"and Dewatripont, M","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.692410Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:650cbc3f7f8a9b1451546de88d5c3febb872fc60805c6d65289d0c9a2b02a8da","observation_id":"f7c902c2-77fd-4bdb-8377-e3adc5a21682","resolution":{"observed_at":"2026-08-11T10:53:10.692410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.192196Z","title":"P., and Kakade, S","venue":null,"work_id":"da063561-0386-4e96-995b-96bbcf4b9fcf","year":2008},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.695970Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:ae1327f9074b84d36b94cc753562f6b3e0ee06ce48e5eb09c4d676088825c41c","observation_id":"75464755-8169-4fcb-827e-e4386aa6c20a","resolution":{"observed_at":"2026-08-11T10:53:11.195595Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.182070Z","title":"D., O'doherty, J","venue":null,"work_id":"21314eb4-4557-4e3d-b6c3-4c0200fa50c4","year":2006},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.699587Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:610e9b23380e3466b757d50597911fe4e9b569c33d60b683f86246223fd74f88","observation_id":"c172d0fa-e3df-4e64-91fa-aa3455c28083","resolution":{"observed_at":"2026-08-11T10:53:11.185652Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.06717","last_updated":"2023-08-13T08:12:01Z","snapshot_observed_at":"2026-08-16T15:08:51.830898Z","submitted_at":"2023-08-13T08:12:01Z","title":"Estimating and Incentivizing Imperfect-Knowledge Agents with Hidden Rewards","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.06717","snapshot_observed_at":"2026-08-11T10:53:10.703651Z","title":"M., and Aswani, A","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.703651Z"},"links":{"cited_paper":"/paper/2308.06717","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:c512bd16fe7a1219f5ef3e0c948f67e1d12df84907a147cde5099e9ab406cda7","observation_id":"3e003763-fcdc-467c-9cab-8a9fca62de84","resolution":{"observed_at":"2026-08-11T10:53:10.703651Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.07407","last_updated":"2023-05-07T19:30:01Z","snapshot_observed_at":"2026-08-16T15:40:18.499781Z","submitted_at":"2023-04-14T21:57:16Z","title":"Repeated Principal-Agent Games with Unobserved Agent Rewards and Perfect-Knowledge Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.07407","snapshot_observed_at":"2026-08-11T10:53:10.707423Z","title":"M., and Aswani, A","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.707423Z"},"links":{"cited_paper":"/paper/2304.07407","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:5c0978595b134d706a9ac2b86ac458fadab9f9677213537ba8665ed8dd0851a0","observation_id":"7ded677a-d13d-44ba-baab-4df13c707633","resolution":{"observed_at":"2026-08-11T10:53:10.707423Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.171804Z","title":"and Szentes, B","venue":null,"work_id":"d9d872e1-c7ce-4ff1-ba3b-fe070d189736","year":2017},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.711041Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:cf92294191e0d8fe2736b43ff2125a173564959a8401f93c758f8a5c0642a667","observation_id":"ab8da7fd-cc74-4154-b692-b93ba863cd71","resolution":{"observed_at":"2026-08-11T10:53:11.175345Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.160733Z","title":"Action elimination and stopping conditions for the multi-armed bandit and reinforcement learning problems","venue":null,"work_id":"c4405a21-ea00-4359-a15f-59976e1a9c3f","year":2006},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.714318Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:9303de431957d6f6a4a1551a4ad9e6578f68644938258b417b1faf82ab592d66","observation_id":"6bb27c9f-7a80-4be8-a6d1-82378caa5a78","resolution":{"observed_at":"2026-08-11T10:53:11.164531Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1803.04008","last_updated":"2019-03-01T22:03:13Z","snapshot_observed_at":"2026-08-14T19:37:23.056260Z","submitted_at":"2018-03-11T18:44:50Z","title":"Multi-Armed Bandits for Correlated Markovian Environments with Smoothed Reward Feedback","version":2},"cited_work":{"arxiv_id":"1803.04008","doi":null,"metadata_source":"pith","pith_arxiv_id":"1803.04008","snapshot_observed_at":"2026-08-11T10:53:10.890799Z","title":"Multi-Armed Bandits for Correlated Markovian Environments with Smoothed Reward Feedback","venue":"cs.LG","work_id":"bb7e9629-43aa-4437-b62a-9c0a1d9f8180","year":2018},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.717837Z"},"links":{"cited_paper":"/paper/1803.04008","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:241253dfc5dd40749ddb20fed3c1cf4ad578d6eb3962da1038e77bd83020b7f9","observation_id":"008b5109-ee6d-4c25-b46a-da988f20ae7e","resolution":{"observed_at":"2026-08-11T10:53:10.898079Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.150435Z","title":null,"venue":null,"work_id":"d51ace3a-5458-4208-828c-f466d4330729","year":2018},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.721547Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:e5a1031badd240b0a3400960dd13d82495ef9586dfe9911c38fba4728be998ab","observation_id":"a87eb957-7692-47bc-8573-95ac2325a162","resolution":{"observed_at":"2026-08-11T10:53:11.154078Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.139787Z","title":null,"venue":null,"work_id":"01382209-c5bd-4207-b54b-e469e7735f80","year":2018},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.724761Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:f95956b8cc2b54057a0d5f2fcb69cf2486f549a622dbd99707d289d8903b6767","observation_id":"7c018f40-b653-4eae-bab7-b7d722622dfd","resolution":{"observed_at":"2026-08-11T10:53:11.143474Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.129790Z","title":"and Moreira, H","venue":null,"work_id":"570acffe-9323-47f4-9825-8bf2ad28bbc8","year":2022},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.728215Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:52defe6fe425b802ca45d31fbffa3c8f563661667b3558132a010c16e56381bb","observation_id":"5983947f-143f-488d-aae1-2d1e82c66081","resolution":{"observed_at":"2026-08-11T10:53:11.133268Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.731668Z","title":null,"venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.731668Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:9239b11a4dd57c282c9cee1bd1a912ee6b3afa9fbc878fce9d6d4b1cd7900981","observation_id":"f8c86e4b-4b61-4533-9550-56e739ae15ce","resolution":{"observed_at":"2026-08-11T10:53:10.731668Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.113254Z","title":"Optimal contracts for experimentation","venue":null,"work_id":"61106cf3-d1a7-4dfb-9ba1-0f6b2b19b383","year":2016},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.734832Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:1b63cca28c4492c2be15ec5a0433441eb3fb8fde1f5c645f5452436e189b329c","observation_id":"9c2afcf9-8c5b-4bd1-9993-23131b2ead35","resolution":{"observed_at":"2026-08-11T10:53:11.117271Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.738015Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.738015Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:7993a98206d385caeed8e9c664e3dc992a2c9a682fdd31d6bdec6bd896db539c","observation_id":"8ac7fde2-9008-4747-9a20-2bb56bc4e464","resolution":{"observed_at":"2026-08-11T10:53:10.738015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.18074","last_updated":"2024-10-07T16:46:42Z","snapshot_observed_at":"2026-08-16T13:31:04.606847Z","submitted_at":"2024-07-25T14:28:58Z","title":"Principal-Agent Reinforcement Learning: Orchestrating AI Agents with Contracts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.18074","snapshot_observed_at":"2026-08-11T10:53:10.741502Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.741502Z"},"links":{"cited_paper":"/paper/2407.18074","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:a9b17acb3cd23edb766ce2597c5ffb3c981d3c8af20bfd55b1fa18aebd987bc6","observation_id":"f9d6d22a-a26a-4a25-8dfd-004ca432022b","resolution":{"observed_at":"2026-08-11T10:53:10.741502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.096045Z","title":"and Bowman, H","venue":null,"work_id":"b3e05b91-600d-4961-8d09-f90186a0666a","year":2007},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.746076Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:84a2b3d820a9c2ba63576f4116121e7e83ac61406d97befa9d6795e8badd7aa3","observation_id":"33442020-b8e0-42e0-8e8f-a01b1593563a","resolution":{"observed_at":"2026-08-11T10:53:11.100149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.749256Z","title":"and Wolfowitz, J","venue":null,"work_id":null,"year":1960},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.749256Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:863e7b212ec0533ff96cda5947fc1f670c0defa11fcbfe82856a86de35bccdc5","observation_id":"bf418276-3728-4607-8785-bd2c63ad3d45","resolution":{"observed_at":"2026-08-11T10:53:10.749256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.079766Z","title":"and Martimort, D","venue":null,"work_id":"1cc5b0e7-8a3c-4fe2-b825-91fec0e82846","year":2009},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.752421Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:dde82810806c28c31b639de426529fc9b064b667dbccbc77822c2d12b0b341ec","observation_id":"64a5e0e7-bcc1-4569-90fb-138b42bd9df8","resolution":{"observed_at":"2026-08-11T10:53:11.083233Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.755773Z","title":"and Szepesv \\'a ri, C","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.755773Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:801cc3a5a2b61f14764d546c58aa897336c0ea545011b775912018b9f5b839c2","observation_id":"cff73efe-c5d2-48ea-83ec-4fb6c8cb916c","resolution":{"observed_at":"2026-08-11T10:53:10.755773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.062637Z","title":"Linear multi-resource allocation with semi-bandit feedback","venue":null,"work_id":"18107528-4b51-441f-a71e-1cc1615cd6f6","year":2015},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.758973Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:95f3529ecc1e14cfee31cb30039a0bc04f7478c0120b869e15747932bd4ef150","observation_id":"4bc9284e-30fa-4f86-ae59-f0a2bd6590de","resolution":{"observed_at":"2026-08-11T10:53:11.066429Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.052079Z","title":"Learning with good feature representations in bandits and in rl with a generative model","venue":null,"work_id":"d611d936-772f-4779-b059-d241155417ba","year":2020},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.762145Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:81dab8fc68ac9340329ab49161e4372b253df2b95ecd9571055615965702f51d","observation_id":"87fdec8c-6001-45a7-b76c-5a7007c53629","resolution":{"observed_at":"2026-08-11T10:53:11.056022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.041504Z","title":"D., Zhang, S., Munro, M., and Steyvers, M","venue":null,"work_id":"bbafa458-7e61-48da-9cf2-0b61873a93b8","year":2011},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.765268Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:927bec7046d77812ffe2fe30e1f11a9136e28463f24be341179e8cb713aa6cb8","observation_id":"b72f5621-2e0c-414d-a4a5-a848d01c380b","resolution":{"observed_at":"2026-08-11T10:53:11.045348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.030791Z","title":null,"venue":null,"work_id":"99eee9dc-ad10-47c8-a088-6918802a0eba","year":2010},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.768354Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:ae338db1c78635ac6e90d26eb55d43e0c3e90773052d9596439a83e2a2db0976","observation_id":"bbf787ed-48bd-479f-afef-5a05739bd540","resolution":{"observed_at":"2026-08-11T10:53:11.034383Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.020458Z","title":"and Chen, Y","venue":null,"work_id":"2ecfa847-139a-42ec-9bb6-f84c6a2e35c7","year":2025},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.771793Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:4fae4d88add010372bb8f70c8ed752b69964f0cf6673bd1e91e89d36217a3228","observation_id":"95cac984-2cc3-444c-a07a-273169eeecf5","resolution":{"observed_at":"2026-08-11T10:53:11.023908Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:11.010052Z","title":"P., and Schneider, J","venue":null,"work_id":"6866275f-774c-4680-ac2d-475ce913f3f1","year":2021},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.774786Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:23ce4fbaef9c90cfe10adcf214ef0806e15a8f7d968b6ac2a1817cc179c15d14","observation_id":"737e0a4c-e431-4156-97fb-018d3302f866","resolution":{"observed_at":"2026-08-11T10:53:11.013708Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.998789Z","title":"and Lattimore, T","venue":null,"work_id":"36e09761-065f-4992-ad33-58f81c7fafa0","year":2020},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.777957Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:1330edf7b34e880f7fddc4f0af1d097eb4882901a5b28bd44941dd439d769c3c","observation_id":"d385d668-bc48-4b72-bce0-5ad418f54d1b","resolution":{"observed_at":"2026-08-11T10:53:11.002445Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.987525Z","title":"Monitoring cooperative agreements in a repeated principal-agent relationship","venue":null,"work_id":"d532f059-97cf-49bc-a2f5-a38a21887eb2","year":1981},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.781236Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:477f3e491a2bb60fa726983a71601984cbd91a840e14d02784119bec289a055c","observation_id":"f144d51e-d2be-4159-9c03-ee896725b7b0","resolution":{"observed_at":"2026-08-11T10:53:10.992406Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.977189Z","title":null,"venue":null,"work_id":"6844cf6a-6da4-4b03-ae96-b964df7316c4","year":2020},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.784731Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:bc79308f53cf782792a97ce08562b0bad49917e904b12fe0abcea1e25610b046","observation_id":"27f66512-ac4c-438a-96eb-02a9e72f6b66","resolution":{"observed_at":"2026-08-11T10:53:10.980954Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.966769Z","title":null,"venue":null,"work_id":"d2d45f2b-7310-4346-b3c4-95aa16f32b5f","year":1985},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.788055Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:9c221cc51c70785f651d44b3784efbba9b2dfc46f7e9bd5bedf539211898c2e3","observation_id":"4d6203dd-f99a-48c1-9bff-640d342022d8","resolution":{"observed_at":"2026-08-11T10:53:10.970132Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.956593Z","title":"Contracts: the theory of dynamic principal--agent relationships and the continuous-time approach","venue":null,"work_id":"479cc1ac-fabd-4a54-a41c-01142e0604d0","year":2013},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.791196Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:4e104bab8fc2ec8db03f41eef3e57742f3e2bb32a48228b5d67fcc6a40a9ae40","observation_id":"671e2c90-4437-40f4-8495-149620ae4924","resolution":{"observed_at":"2026-08-11T10:53:10.960139Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19824","last_updated":"2025-01-28T17:37:59Z","snapshot_observed_at":"2026-08-16T13:38:46.868402Z","submitted_at":"2024-06-28T11:00:53Z","title":"Learning to Mitigate Externalities: the Coase Theorem with Hindsight Rationality","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.19824","snapshot_observed_at":"2026-08-11T10:53:10.794350Z","title":"I., and Durmus, A","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.794350Z"},"links":{"cited_paper":"/paper/2406.19824","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:1a5e63ef4e9ec248454b1eb4773b40211701dfac7cbe173e6524290a1d832514","observation_id":"697705b9-fe7d-4ca2-a804-daa42b50778c","resolution":{"observed_at":"2026-08-11T10:53:10.794350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.946251Z","title":null,"venue":null,"work_id":"e4bca3a5-fd15-45bf-8884-cd34f8c6b1b6","year":2024},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.798208Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:cf8b5616a7ae4d970bfa703eec5829514ccc4ce917810568e1705ed697b162b0","observation_id":"1aed3558-faf0-4a00-84c5-24a6f07f4c42","resolution":{"observed_at":"2026-08-11T10:53:10.950107Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.935427Z","title":null,"venue":null,"work_id":"a5360c77-3b83-4550-9a5c-9dbbdbac7d00","year":2016},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.801557Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:2f350b7d4c17bbce54181086304e98188adfeefe6d07ba6c2962f46179a2208d","observation_id":"ba47e262-0d43-48d7-9f87-a023853a77ff","resolution":{"observed_at":"2026-08-11T10:53:10.939330Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T10:53:10.925169Z","title":"S., Bowden, J., and Wason, J","venue":null,"work_id":"98194961-449c-4a16-8e6a-9e20dadc3d98","year":2015},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.804857Z"},"links":{"citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:29bb6564a74f634ec56410b9835f25a73a1dc02f6ce415e8538794fc77c7772f","observation_id":"15392dd3-50ff-4aa3-adb1-410bdc41bb93","resolution":{"observed_at":"2026-08-11T10:53:10.928750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.01458","last_updated":"2024-07-02T15:17:50Z","snapshot_observed_at":"2026-08-16T13:38:02.779122Z","submitted_at":"2024-07-01T16:53:00Z","title":"Contractual Reinforcement Learning: Pulling Arms with Invisible Hands","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.01458","snapshot_observed_at":"2026-08-11T10:53:10.808207Z","title":"Contractual reinforcement learning: Pulling arms with invisible hands","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.808207Z"},"links":{"cited_paper":"/paper/2407.01458","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:1251fe5e136d08e17e43fcc70d8b7cc7094275d40fe8db3252b732026b84b866","observation_id":"e4590607-40c1-406b-8804-84ec0ad94f9b","resolution":{"observed_at":"2026-08-11T10:53:10.808207Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05732","last_updated":"2023-05-19T23:18:55Z","snapshot_observed_at":"2026-08-16T16:17:07.550675Z","submitted_at":"2022-11-10T17:59:42Z","title":"The Sample Complexity of Online Contract Design","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05732","snapshot_observed_at":"2026-08-11T10:53:10.811900Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.811900Z"},"links":{"cited_paper":"/paper/2211.05732","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:c801710d1c7f12c76aa3914d085db5bc995089d55cbe8ccc9ad80442c49be833","observation_id":"5e0e1328-219a-48b5-a024-74aeac7dd9e2","resolution":{"observed_at":"2026-08-11T10:53:10.811900Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11381","last_updated":"2023-05-19T01:58:13Z","snapshot_observed_at":"2026-08-16T15:31:50.228233Z","submitted_at":"2023-05-19T01:58:13Z","title":"Online Learning in a Creator Economy","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11381","snapshot_observed_at":"2026-08-11T10:53:10.815652Z","title":"P., Jiao, J., and Jordan, M","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-11T10:53:10.815652Z"},"links":{"cited_paper":"/paper/2305.11381","citing_paper":"/paper/2412.16318"},"observation_digest":"sha256:b0c89b7cf9e89c372ccb66f2b14d5fdd6075b55d5fb30cad12c88f7a7b29f429","observation_id":"abcac524-75f0-441a-bbeb-b4a056186cc8","resolution":{"observed_at":"2026-08-11T10:53:10.815652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.16318","last_updated":"2025-06-02T03:32:57Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-11T16:27:24.662951Z","submitted_at":"2024-12-20T20:04:50Z","title":"Principal-Agent Bandit Games with Self-Interested and Exploratory Learning Agents"},"reference_resolution":{"displayed":41,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":21,"verified_exact":1,"verified_fuzzy":19},"total_outbound_references":41},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 41 of 41 outbound references and 1 inbound Pith citation observation for arXiv:2412.16318."}