{"as_of":"2026-08-15T07:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:77813f105e4c531a62271b563ef8aae93d4148639f1934dafc2a2834fe74057b","coverage":[{"denominator":100,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T04:48:47.456192Z","state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-15T06:32:42.880941+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.28520/citation-record","integrity":"/paper/2607.28520/integrity","json":"/paper/2607.28520/citation-record.json","paper":"/paper/2607.28520"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.191281Z","title":"Bowling , title =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.191281Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:75bda030c38e4a0694f9debe764343434be80a3421f26513a5d47a089bfd518c","observation_id":"c8fbee1f-737b-425a-ad86-6c37619fe01b","resolution":{"observed_at":"2026-07-31T04:48:47.191281Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.194494Z","title":"Artificial Intelligence and Statistics , pages=","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.194494Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:773dca57b6f3ec0affe62ae72dfbf0031fbafd8664c66bdd17aac01968816183","observation_id":"2145736d-1962-4b09-af5a-8b557e5cb668","resolution":{"observed_at":"2026-07-31T04:48:47.194494Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.197341Z","title":"ACM Transactions on Economics and Computation (TEAC) , volume=","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.197341Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:6953121048368b2fb7185ceb17fbe0df027d7430088e53ebabadb9bbc4caf8fb","observation_id":"bd731ca9-ab3f-4a14-86fd-86d4850ef2d7","resolution":{"observed_at":"2026-07-31T04:48:47.197341Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.200068Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.200068Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:cbb0c8f6fc779427e126334c544d12a42f37c05169b7bb6ecc2b0737de15808f","observation_id":"9b0a1b1f-9fb7-4332-9360-2e0e91d20d94","resolution":{"observed_at":"2026-07-31T04:48:47.200068Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.206062Z","title":"Science , volume=","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.206062Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ccafa6c207501e878b27fb508328cced0c0a21e0c0c83333a86666293b52af18","observation_id":"62304136-5b85-414d-8045-0003d8ad6de4","resolution":{"observed_at":"2026-07-31T04:48:47.206062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.208931Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.208931Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:5ca58c2ed1b72d6209e5dca48d9e1b77b9b10a978549340cd4b4d3273a2f5b5a","observation_id":"02307224-e046-4fe3-908c-13a0167ce6a4","resolution":{"observed_at":"2026-07-31T04:48:47.208931Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.211562Z","title":"International Conference on Machine Learning , pages=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.211562Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:74ed2faf25fdd2fcb8f6dca9eb0fc6dc75d288754a7d536220a9ef3220375fe9","observation_id":"f35b3d1e-57df-4231-bcec-b3bd1debbfca","resolution":{"observed_at":"2026-07-31T04:48:47.211562Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.214270Z","title":"Proceedings of the 24th International Conference on Autonomous Agents and Multiagent Systems , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.214270Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:24d953ec9a02db1bf3569d2138e599294e100084298bffcad18a88438d063a74","observation_id":"c8013b70-776d-41f9-bf38-4beb10b7945a","resolution":{"observed_at":"2026-07-31T04:48:47.214270Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.217252Z","title":"Forty-first International Conference on Machine Learning , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.217252Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e56ba617cfbb793c48df28f9a01dbe86bc654c6b1737d4a48b19022ca596a165","observation_id":"6785c2f6-11e4-4b7a-875f-36b50233a49b","resolution":{"observed_at":"2026-07-31T04:48:47.217252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.220009Z","title":"Science , volume=","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.220009Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e30979d7e97711aa0024ec5d8b019bbec1ff81733c3de26e890625e2ca84f587","observation_id":"f57299c5-decc-4f89-9663-84beb93cf073","resolution":{"observed_at":"2026-07-31T04:48:47.220009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.222502Z","title":"Science , volume=","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.222502Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:3ae44ce8988eed3f417bef456fb5c0886b26846e54614b7655779dfe4ff81ca0","observation_id":"362d9309-1580-44e5-b153-b0d2c7301133","resolution":{"observed_at":"2026-07-31T04:48:47.222502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.225137Z","title":"Science , volume=","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.225137Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:82ca1a5d19438f0453645c18b53de274df4a65db78dfe8e152cbe982ac5ecb2c","observation_id":"67ae75b7-fbb2-48cc-944c-dd98fda1b9d1","resolution":{"observed_at":"2026-07-31T04:48:47.225137Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.227481Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.227481Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:62a66b207189b79012fe6459645abf2bdf55d7869c7d2f0304425e7e36dd3039","observation_id":"5ccb8e8a-207f-4898-96db-518d6021f182","resolution":{"observed_at":"2026-07-31T04:48:47.227481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.229923Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.229923Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e926bb9954e82037b2fdf2966dd1b85af4b5a912ce8a8d73d49902b8f1627e40","observation_id":"a1486383-e520-4bb3-8d31-e0748df522a2","resolution":{"observed_at":"2026-07-31T04:48:47.229923Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.232290Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.232290Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:77884a32821c0231d7f56a414d9127086bdecb8e8a3ca037084ab6fd16e33362","observation_id":"428bda40-ef63-4edf-8a60-66a02719f73a","resolution":{"observed_at":"2026-07-31T04:48:47.232290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.234761Z","title":"Proceedings of the AAAI conference on artificial intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.234761Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e6775e54e70634369bf3871631093a548de6c851a895fe770d3cb5ef46df9933","observation_id":"37b97009-bb19-48f1-b405-6ab69f6422c7","resolution":{"observed_at":"2026-07-31T04:48:47.234761Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.237440Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.237440Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:f9670adb21bf29aae3daf19462a080dd35a1215f283149a909a6daea8900847f","observation_id":"d14a0dbe-c1b4-41d6-9766-f1aba462be1c","resolution":{"observed_at":"2026-07-31T04:48:47.237440Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.239726Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.239726Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:aadc66f343044ccb13030b85dfafa5616ff5386a852aa7bc8019e929cf038b49","observation_id":"e45043cf-f62d-40f8-b8d4-212a15a9dc23","resolution":{"observed_at":"2026-07-31T04:48:47.239726Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.242214Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.242214Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:3bd03f18fdc261211eb59ec4b0054c309f707b5faf62ea9ace25cb18f645b73b","observation_id":"f3bb24ad-8208-44f8-a28b-ad3f9d28fac2","resolution":{"observed_at":"2026-07-31T04:48:47.242214Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.244547Z","title":"Proceedings of the Twenty-First Conference on Uncertainty in Artificial Intelligence , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.244547Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:9e8daf91edd8a040e4b0f5b3ebca3dfdbc2e1e09cfee22bb80b5769d17f52f16","observation_id":"53e1ba60-9f8e-4c1b-8610-efc8ab4557d8","resolution":{"observed_at":"2026-07-31T04:48:47.244547Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.246808Z","title":"The Annals of Statistics , volume=","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.246808Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:30b2af0591c942d5aafd077d79df50921c51f50537debfa82be82e52e956a84e","observation_id":"e15db01b-900b-4a20-943b-67523188fba3","resolution":{"observed_at":"2026-07-31T04:48:47.246808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.249155Z","title":"Journal of the Royal Statistical Society Series B: Statistical Methodology , volume=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.249155Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:7b88fa5e383e6e287281885e27046264c6d846838dd4ad65eb45f4cce52f89eb","observation_id":"e023aea7-1616-40d8-85e7-fada2b8a44e8","resolution":{"observed_at":"2026-07-31T04:48:47.249155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.256765Z","title":"CoRR , volume =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.256765Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:473964190a18dc898baf58256e3fd2f4163cd23de956f1bcb696d769ac2a75ac","observation_id":"20423ccb-0642-4200-bc96-65235570d0c2","resolution":{"observed_at":"2026-07-31T04:48:47.256765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.259014Z","title":"The International FLAIRS Conference Proceedings , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.259014Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:aeab17641e34efdb6f0c4ffd7762f67e70a3a3b849c4f0f8d825f0f346b114c2","observation_id":"51a16d91-a2b0-42b6-9dd8-0403a9e60647","resolution":{"observed_at":"2026-07-31T04:48:47.259014Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.261397Z","title":"AAAI Technical Report (2) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.261397Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:53343396855a7f98b1836506a303657657514c9cff20cac0e17468ecbcb13832","observation_id":"0456e291-f538-4392-8ecb-2e3c6dcfeb0a","resolution":{"observed_at":"2026-07-31T04:48:47.261397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.263661Z","title":"AAAI , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.263661Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:fe868d4941ead682e097fd0ba386f694d55c322ed589bbb6a3f05b286d0b49d7","observation_id":"3c69ef12-e3b2-4fdd-93f8-8cfb998b0232","resolution":{"observed_at":"2026-07-31T04:48:47.263661Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.266175Z","title":"Proceedings of the 3rd AAAI Conference on Interactive Decision Theory and Game Theory , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.266175Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:43be82c7461dc895924003b6ae8348e7b590512cca47549b092be4f288a2e6b1","observation_id":"1f9d38ea-fbdf-43e2-9cfd-6c347751d246","resolution":{"observed_at":"2026-07-31T04:48:47.266175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.268527Z","title":"Proceedings of the 41st International Conference on Machine Learning , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.268527Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:99e658f056028ce7725050419637bff15db124a43ea7072ca0a46f7a3a322d45","observation_id":"8fc4e63e-b87e-42de-a790-eacefd6577a3","resolution":{"observed_at":"2026-07-31T04:48:47.268527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.270806Z","title":"The Thirteenth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.270806Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:56d7bdea984d42a3181943cd500b2ecf74e80e8dfda951ad3ec1126707f7426e","observation_id":"7c4d4408-2f6d-4a37-b544-57ff779a07fa","resolution":{"observed_at":"2026-07-31T04:48:47.270806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.281175Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.281175Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b19e470f45c09d873a30c70e8c33ca756089f1ae6c41d748ded3d3a73faeea87","observation_id":"9ba0f7d9-892c-4a51-8578-4d3ff79a4116","resolution":{"observed_at":"2026-07-31T04:48:47.281175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.283487Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.283487Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:9b3eb168725e5af1038c6c590526b1856343366da8cb36df143c07b6533c7145","observation_id":"ee395685-1186-400f-b78b-2fa10ac078bf","resolution":{"observed_at":"2026-07-31T04:48:47.283487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.285735Z","title":"IEEE Transactions on Cybernetics , volume=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.285735Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:7a90d797a1e84f0752e13cb87c6ad66d948203a30213d0ae2a8b8f4c94ab5d8e","observation_id":"4e21a80e-dd8f-44ab-bbc3-aa0a09387974","resolution":{"observed_at":"2026-07-31T04:48:47.285735Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.288232Z","title":"Journal of Artificial Intelligence Research , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.288232Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:6b6adb758d3d9bae157e918a43ec4ae43431e0724fa9cb8ba8c710ea0975d466","observation_id":"918a4a66-ccb3-49f1-9971-33148f8303a8","resolution":{"observed_at":"2026-07-31T04:48:47.288232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.290482Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.290482Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:af13df8b3dfeac9ca364a58e6cea4b8e096d25e309fed763bafd91f220e085d2","observation_id":"029a4655-f316-4d29-9e85-79fc2a56d6c1","resolution":{"observed_at":"2026-07-31T04:48:47.290482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.292873Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.292873Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ee977da09e820b6f250d2a102ba7bf90aa1041319dbe13bf9357f13f32bb6859","observation_id":"42ac89b8-38f1-4559-80c0-11afd659abc8","resolution":{"observed_at":"2026-07-31T04:48:47.292873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.295160Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.295160Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:65a23f4b7c808ce31bc679737da544d8332f1adb4a0bf7678d2e4778aea487f6","observation_id":"be395faf-b533-4731-908f-f64cc2bb686d","resolution":{"observed_at":"2026-07-31T04:48:47.295160Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.297482Z","title":"International Conference on Machine Learning , pages=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.297482Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:35b2c53c632a30f049e1a0adef0b29014d916e6f1437b9aa1b81ecbdd0aabf65","observation_id":"e526642d-2566-4375-8cdf-f94d4a5aa6d5","resolution":{"observed_at":"2026-07-31T04:48:47.297482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.299690Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.299690Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:928b7faca05de23e601295bff0a4bbb2457adb1c4d6f6f61cc00e73dfd708cbe","observation_id":"a6fc2c57-292a-489f-985c-eac2cb1b41f8","resolution":{"observed_at":"2026-07-31T04:48:47.299690Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.301990Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.301990Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:37faf580d1a9e5986520c407737aaad746f859f4be9da99e2ac071b71d3f334c","observation_id":"4f548956-5396-4d64-b3d5-6e059cccb14e","resolution":{"observed_at":"2026-07-31T04:48:47.301990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.304302Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.304302Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:16d4fc995d61be5889906c35ee52b3edae257a26cf02bccf5fefc44128e8f90b","observation_id":"b217de98-af45-46eb-8576-3f20d746ac52","resolution":{"observed_at":"2026-07-31T04:48:47.304302Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.306782Z","title":"Expert Systems with Applications , volume=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.306782Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a1f278ea6b4c9a0e4b9095e247a8b8485c4c74aa0eb5fddf03bf68c2981284ee","observation_id":"efb569d8-74fc-4385-90fe-17d39d556fe2","resolution":{"observed_at":"2026-07-31T04:48:47.306782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.309069Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.309069Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:652ae7d1b4755ff1ea874f44a7b91e9765db1ca730ff8c946ed3e8e0815c7238","observation_id":"b3fdd216-3594-424e-a048-bad677e3f6c6","resolution":{"observed_at":"2026-07-31T04:48:47.309069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.312132Z","title":"The Thirty-ninth Annual Conference on Neural Information Processing Systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.312132Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:f2a5e8ec10b4151862a53c6740a1b2492c9ceca009749886519d841763eaf3c6","observation_id":"51160048-8a13-404e-8bfc-9ddf938d701f","resolution":{"observed_at":"2026-07-31T04:48:47.312132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.314571Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.314571Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e5947b69362abe4f66be237a65891bfc22cce11e2cf07d070320002fe1c40cca","observation_id":"e576e758-459c-462f-b370-19bb22c93c66","resolution":{"observed_at":"2026-07-31T04:48:47.314571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.316953Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.316953Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e1adfe3298c3dc016f17f6c125147556dcde8b8334ff30b5e8f8e5db5eb57c1c","observation_id":"c32e7f24-c8b6-4096-b506-5a3cdf311680","resolution":{"observed_at":"2026-07-31T04:48:47.316953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.324527Z","title":"2026 , eprint=","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.324527Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:2c96596d5ba80fa64507fe88f2f1bac3cae44b5a267c20cf11ff15c2222c8d4b","observation_id":"07816d28-fb5e-4ff3-b86d-760688954725","resolution":{"observed_at":"2026-07-31T04:48:47.324527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.327029Z","title":"Exploiting opponents under utility constraints in sequential games","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.327029Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b9d962be631c1a4b81f8467849c67c2b7b47254cb7f5862431a3d9e9e3472f53","observation_id":"bf9eca4a-bf34-4fc1-8e1b-802d62661485","resolution":{"observed_at":"2026-07-31T04:48:47.327029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.329452Z","title":"Heads-up limit hold'em poker is solved","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.329452Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:981d0b8e65d2c509170ccb8922c1c332617bf3ff7094d5a95b5d2f49308f33e0","observation_id":"6ab6f38b-0fa6-466d-a44f-699d93970a83","resolution":{"observed_at":"2026-07-31T04:48:47.329452Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.331971Z","title":"Regret-based pruning in extensive-form games","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.331971Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:4543afda147bdfcb1e93982fb4a05d81ceabfbf5efbc1b7af08f03c5f636df2d","observation_id":"4e7d203f-8794-4b9d-8c25-97291aceaf8a","resolution":{"observed_at":"2026-07-31T04:48:47.331971Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.334352Z","title":"Safe and nested subgame solving for imperfect-information games","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.334352Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:515662c2db154fd9a838b6388d38c7fc5c74d89cee89b857c164146d0aee0ff8","observation_id":"9f13f61f-2533-42ae-a566-0c343c97eb7b","resolution":{"observed_at":"2026-07-31T04:48:47.334352Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.336639Z","title":"Superhuman ai for heads-up no-limit poker: Libratus beats top professionals","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.336639Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0736621246b4abe43d510d0be6346c7baabcfe59d7b235730698e9b4ddf57ba5","observation_id":"a5cd1d36-fa7d-4dc9-ba8f-14a91cedccbb","resolution":{"observed_at":"2026-07-31T04:48:47.336639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.338989Z","title":"Solving imperfect-information games via discounted regret minimization","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.338989Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:67d8c80821ed168dd3ad348d7ad7adcbc510afd79a763556d5f6643760fd1925","observation_id":"d21ca11a-5ea9-443d-b96f-df986a320d57","resolution":{"observed_at":"2026-07-31T04:48:47.338989Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.341388Z","title":"Superhuman ai for multiplayer poker","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.341388Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:4696541f46860dfb38a45db394309797ddc548a7b396d84d5ce14ebde721eea4","observation_id":"1684e01c-9738-49e1-8f5c-a57893999323","resolution":{"observed_at":"2026-07-31T04:48:47.341388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.343696Z","title":"Dynamic thresholding and pruning for regret minimization","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.343696Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:93df87e8cf8550dabbdff4a77f1e8a79842b5507a7fa9355c0947818f2fb40af","observation_id":"5fbdcc35-499d-4085-a2be-12501ad5fbd7","resolution":{"observed_at":"2026-07-31T04:48:47.343696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.345938Z","title":"Depth-limited solving for imperfect-information games","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.345938Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1b6986ba27dd9f54b9c6022939673928d2a2dec9b24be8085e48ee032411e2e3","observation_id":"8d7a1f34-a9ff-4c3f-aa2d-f98b170f88b8","resolution":{"observed_at":"2026-07-31T04:48:47.345938Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.348215Z","title":"Deep counterfactual regret minimization","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.348215Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e9bf298f6c2fc17e14393e78c7149267103bd293c70dad65a77144084ce368c4","observation_id":"13f9e895-270f-44cc-a357-78ad9a6dbc50","resolution":{"observed_at":"2026-07-31T04:48:47.348215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.350681Z","title":"Combining deep reinforcement learning and search for imperfect-information games","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.350681Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:f8977052894a5a409a39d2fadeb7ab78066a341383fcc87a06d3d59e3ba4486a","observation_id":"e25a7da1-7264-4780-a207-564d029ca8ef","resolution":{"observed_at":"2026-07-31T04:48:47.350681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2604.25796","last_updated":"2026-04-28T16:03:14Z","snapshot_observed_at":"2026-08-15T01:13:23.899496Z","submitted_at":"2026-04-28T16:03:14Z","title":"StratFormer: Adaptive Opponent Modeling and Exploitation in Imperfect-Information Games","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2604.25796","snapshot_observed_at":"2026-07-31T04:48:47.352939Z","title":"Stratformer: Adaptive opponent modeling and exploitation in imperfect-information games","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.352939Z"},"links":{"cited_paper":"/paper/2604.25796","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:5bbc7af119839dfa24deb255bce95c992c423abdec1586e4a106bb9061ff45fd","observation_id":"60860922-b7b2-4aa0-951b-169d26ec1b3e","resolution":{"observed_at":"2026-07-31T04:48:47.352939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.355340Z","title":"Test-then-punish: A statistical approach to repeated games","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.355340Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:cbf64f0137fc2de18586e8184948b832ab7c15ddae86cf853c8da765cd35b22d","observation_id":"fb627b75-bb06-48b7-b737-14560d1066f6","resolution":{"observed_at":"2026-07-31T04:48:47.355340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.357580Z","title":"Faster game solving via predictive blackwell approachability: Connecting regret matching and mirror descent","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.357580Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:73094c88df2f7fd090082b8ccbc4a1138f8cd35cc67f382a0991248d567f184b","observation_id":"9a7f5ecc-47df-4264-8e7e-a84f97a047f6","resolution":{"observed_at":"2026-07-31T04:48:47.357580Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.359894Z","title":"Regret matching+:(in) stability and fast convergence in games","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.359894Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:bcbb320ce1bf5b524adc58f349732b35fafe0b792a56e7e11669919c6d859f6b","observation_id":"5909783a-a95a-4973-ade1-106ebb0a8bbb","resolution":{"observed_at":"2026-07-31T04:48:47.359894Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.362064Z","title":"Greedy when sure and conservative when uncertain about the opponents","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.362064Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e8e459a2824a5c8dfdfb3df5dae910accfac55bd2973fa0f0d86c50a4e00454d","observation_id":"44dc54f6-4101-4a0c-a91b-cb1d52494560","resolution":{"observed_at":"2026-07-31T04:48:47.362064Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.17671","last_updated":"2026-07-26T21:15:12Z","snapshot_observed_at":"2026-08-07T06:10:50.682271Z","submitted_at":"2025-08-25T05:08:49Z","title":"Consistent Opponent Modeling in Imperfect-Information Games","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.17671","snapshot_observed_at":"2026-07-31T04:48:47.364343Z","title":"Consistent opponent modeling of static opponents in imperfect-information games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.364343Z"},"links":{"cited_paper":"/paper/2508.17671","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:d51680320c57ccd71764b5ea7ca6a26e4f80b2af93dcbb1b41fdb3ebd08ca031","observation_id":"32afac23-d830-4e68-9d3f-0aa4f6c1692f","resolution":{"observed_at":"2026-07-31T04:48:47.364343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.366877Z","title":"Nonparametric strategy test","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.366877Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:53c68fde03bcddf72113691144dd26d76101514d4e754088a7b183e421cf4208","observation_id":"3055ca0f-38f2-4789-a549-2d6f37f7d6a6","resolution":{"observed_at":"2026-07-31T04:48:47.366877Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.369505Z","title":"Safe opponent exploitation","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.369505Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ffe343e2e50718ca007f430b3dd4a69488c71ae8ab42c79e978268dcf19b4f42","observation_id":"6126bee9-c13b-4d0f-bd5b-26e8540ecbab","resolution":{"observed_at":"2026-07-31T04:48:47.369505Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2601.05427","last_updated":"2026-05-22T16:37:56Z","snapshot_observed_at":"2026-07-31T05:26:20.127534Z","submitted_at":"2026-01-08T23:26:05Z","title":"Anytime Detection of Strategic Deviations in Multi-Agent Systems","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.05427","snapshot_observed_at":"2026-07-31T04:48:47.372071Z","title":"Betting on equilibrium: Monitoring strategic behavior in multi-agent systems","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.372071Z"},"links":{"cited_paper":"/paper/2601.05427","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:979ac13a691276ba6776bafd01168456699a7fb658dd68db17a2a1ac3ee2461a","observation_id":"509423cb-1d88-4611-9e23-675f6ef702ac","resolution":{"observed_at":"2026-07-31T04:48:47.372071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.374449Z","title":"Modeling rationality: Toward better performance against unknown agents in sequential games","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.374449Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:2043cd9dd11dedbdccae47086cf491ae01b83bfd8ace5c3b367194f456c0756b","observation_id":"16e03f78-fcd7-4305-9560-2460f4685c92","resolution":{"observed_at":"2026-07-31T04:48:47.374449Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.376811Z","title":"Efficient subgame refinement for extensive-form games","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.376811Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:887257ddc40c775669088748ba681df23ca78a8a83cc18425b6d726bf8714c3a","observation_id":"be6533dc-f740-4fee-a582-d6ed481f9e3d","resolution":{"observed_at":"2026-07-31T04:48:47.376811Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.379022Z","title":"Safe and robust subgame exploitation in imperfect information games","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.379022Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0f744cdf4bc7c3287fd3611e6cec64ee23f1805ec8b92de338ec97721c2a8f86","observation_id":"c4e12683-7496-4807-8c67-4193e47ef89c","resolution":{"observed_at":"2026-07-31T04:48:47.379022Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.381403Z","title":"Proposer of the vote of thanks to waudy-smith and ramdas and contribution to the discussion of `estimating means of bounded random variables by betting'","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.381403Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:cff9282a209f886fede357ad9f19f1f2ee075f174ca536b6ee4bbd87a1855269","observation_id":"67d7222f-ba3a-4cea-bb71-d285ad555004","resolution":{"observed_at":"2026-07-31T04:48:47.381403Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.383750Z","title":"Effective short-term opponent exploitation in simplified poker","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.383750Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:d8a2efcc2bcf00f9d5c95fc6adad6a2f075f1165817bf501430c8f4245b17774","observation_id":"559cedc7-0e83-4b73-a813-ea21725b0868","resolution":{"observed_at":"2026-07-31T04:48:47.383750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.386236Z","title":"Time-uniform, nonparametric, nonasymptotic confidence sequences","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.386236Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:567362429fe8bf392ab2d84c17fcca2a48773f264832e8f154f3d4857020671c","observation_id":"3a51915e-9bd6-4e5d-8720-da5cfdc9e1a6","resolution":{"observed_at":"2026-07-31T04:48:47.386236Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.388962Z","title":"Towards offline opponent modeling with in-context learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.388962Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:403290d0c99d0df0ff79cd539a84406dec37ff582372f7a0bda4533033c9a1a2","observation_id":"cd9039aa-61de-4b44-b3dd-b54516cbabc1","resolution":{"observed_at":"2026-07-31T04:48:47.388962Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.391267Z","title":"Opponent modeling with in-context search","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.391267Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:aa7cad92a77358200197b92bc4788d949d02dfb7ac1892eb2f98aa65cb7d4a7c","observation_id":"8bfc98e1-b0a4-4def-a576-b1d0c2be4cc9","resolution":{"observed_at":"2026-07-31T04:48:47.391267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.393533Z","title":"An open-ended learning framework for opponent modeling","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.393533Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:8d717098d3488e06cda2951da7b0950dd0d3deae60fef9c3718f67b28d929615","observation_id":"640c34b6-a159-44b7-a408-fdf24e3bbc46","resolution":{"observed_at":"2026-07-31T04:48:47.393533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.395918Z","title":"Data biased robust counter strategies","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.395918Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:141c8c0e1297b461c1e9b9803b2dd7b1765e8fd2151a41f7d1a3066efc10ad21","observation_id":"9a35002d-e7f6-4fe9-ab91-65eb4269029d","resolution":{"observed_at":"2026-07-31T04:48:47.395918Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.398332Z","title":null,"venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.398332Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:d73719705b2848bad19fded7b14389365244763cd4f4c337a2b5c7ede3bc2e3d","observation_id":"a4d838ec-50b9-4980-8756-824582a1de83","resolution":{"observed_at":"2026-07-31T04:48:47.398332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.400828Z","title":"Efficient online pruning and abstraction for imperfect information extensive-form games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.400828Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:00007aaded8efeba799033ecd57d5aaea95ac313e70160550a4c113e501fb3f2","observation_id":"59837463-2886-45da-8cb4-68a99c471fa9","resolution":{"observed_at":"2026-07-31T04:48:47.400828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.10900","last_updated":"2026-05-11T17:41:16Z","snapshot_observed_at":"2026-08-15T01:42:31.809481Z","submitted_at":"2026-05-11T17:41:16Z","title":"Effective, Efficient, and General Information Abstraction for Imperfect-Information Extensive-Form Games","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.10900","snapshot_observed_at":"2026-07-31T04:48:47.403159Z","title":"Effective, efficient, and general information abstraction for imperfect-information extensive-form games","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.403159Z"},"links":{"cited_paper":"/paper/2605.10900","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:9b70bc1cbacad30d1ef044ad330974e375c7e2f47b972adf3638d24cd5467cb8","observation_id":"6ef40b14-c882-43d0-8137-3e8f8aae0b9a","resolution":{"observed_at":"2026-07-31T04:48:47.403159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.19928","last_updated":"2026-05-19T14:49:30Z","snapshot_observed_at":"2026-08-13T08:59:15.236309Z","submitted_at":"2026-05-19T14:49:30Z","title":"Real-Time Parallel Counterfactual Regret Minimization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.19928","snapshot_observed_at":"2026-07-31T04:48:47.405402Z","title":"Real-time parallel counterfactual regret minimization","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.405402Z"},"links":{"cited_paper":"/paper/2605.19928","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:d5a2aab6cfd74a41ede94b7433b57dacbb10a0601f8491c8ee73d5529afeafbe","observation_id":"107caa0b-cfe2-40c3-9f78-325970cc94ec","resolution":{"observed_at":"2026-07-31T04:48:47.405402Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.407650Z","title":"Rl-cfr: improving action abstraction for imperfect information extensive-form games with reinforcement learning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.407650Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:65b5f6cd0d6cf6f866092dfd3c6d5ef9eb4e85471edc868345bc43e17eec2175","observation_id":"0c987494-cf00-42b2-a323-9532e723384b","resolution":{"observed_at":"2026-07-31T04:48:47.407650Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2607.27035","last_updated":"2026-07-29T15:28:36Z","snapshot_observed_at":"2026-08-08T19:05:43.842976Z","submitted_at":"2026-07-29T15:28:36Z","title":"Correlated Chance Sampling for Monte Carlo Counterfactual Regret Minimization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2607.27035","snapshot_observed_at":"2026-07-31T04:48:47.409963Z","title":"Correlated chance sampling for monte carlo counterfactual regret minimization, 2026 a","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.409963Z"},"links":{"cited_paper":"/paper/2607.27035","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a1504349199e8c2a60f67970bdc88040037aef8123261350ba7d59c59191273a","observation_id":"e04489e2-452a-430d-8bd1-fe7c8c8668dd","resolution":{"observed_at":"2026-07-31T04:48:47.409963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.30094","last_updated":"2026-05-28T15:38:33Z","snapshot_observed_at":"2026-08-12T03:34:37.663986Z","submitted_at":"2026-05-28T15:38:33Z","title":"PokerSkill: LLMs Can Play Expert-Level Poker without Training or Solvers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.30094","snapshot_observed_at":"2026-07-31T04:48:47.412479Z","title":"Pokerskill: Llms can play expert-level poker without training or solvers","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.412479Z"},"links":{"cited_paper":"/paper/2605.30094","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1dc09f17af5de9ede75034fc347c2029b6e0d970ca0f1d3af3f976f9fd2c11ab","observation_id":"8d7483f5-77fc-4ff0-92d7-25710fcdf6ea","resolution":{"observed_at":"2026-07-31T04:48:47.412479Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.414777Z","title":"Safe opponent-exploitation subgame refinement","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.414777Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:5473d4bc2ff0d83290d767e56dd99fcc77741ff37c2fed6622516c3e0c2d5a51","observation_id":"a9e2647a-dc66-4c99-b9a6-0af7618f4836","resolution":{"observed_at":"2026-07-31T04:48:47.414777Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.417439Z","title":"Opponent-limited online search for imperfect information games","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.417439Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1ddae4fcd70b67a12114a866201c5344a8a2faa766106238288bbbf4ea620e40","observation_id":"4a5b6f77-6b26-4d9c-aebf-053a2eaf1d0f","resolution":{"observed_at":"2026-07-31T04:48:47.417439Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.419742Z","title":"Safe strategies for agent modelling in games","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.419742Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b915b3cd61decc4d6d25276ac6375a92d6f2eeb192fcd611c9496a528f8e47d1","observation_id":"810e6e45-1e4b-40e7-9c74-0fb40e08d761","resolution":{"observed_at":"2026-07-31T04:48:47.419742Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.422670Z","title":"Efficient last-iterate convergence in solving extensive-form games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.422670Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b1f74cc7f3e78fe71730456c56fbca89136fe80db9a612cd6513c5f9b14a3c62","observation_id":"d8ac54b5-8b81-4fcb-9c7a-e2e8036ff334","resolution":{"observed_at":"2026-07-31T04:48:47.422670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.425081Z","title":"Faster game solving via asymmetry of step sizes","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.425081Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e983378b6d888234a3dbb1f0c4854ee5faa013bb91366f8e704e9e2c7f8adeb5","observation_id":"d828de44-0604-4d31-94cc-ed32bc6c1dd5","resolution":{"observed_at":"2026-07-31T04:48:47.425081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.427532Z","title":"Adapting beyond the depth limit: Counter strategies in large imperfect information games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.427532Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:2b06ecb73287e5d7b08dabbfa8f9af9bf03782de094b423ca0c3ddec4232885e","observation_id":"f512d5a2-7c16-43d1-b2f0-dd27ea2f21a1","resolution":{"observed_at":"2026-07-31T04:48:47.427532Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.429806Z","title":"Deepstack: Expert-level artificial intelligence in heads-up no-limit poker","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.429806Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:f7a86664c42bc1ac24fe45d815279c6e47e759c43fb9292813e67ef5280ba904","observation_id":"b7f0a2bb-3fe6-42e5-8735-4d6273787a85","resolution":{"observed_at":"2026-07-31T04:48:47.429806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.09150","last_updated":"2026-05-09T20:26:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-09T20:26:35Z","title":"AlphaExploitem: Going Beyond the Nash Equilibrium in Poker by Learning to Exploit Suboptimal Play","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.09150","snapshot_observed_at":"2026-07-31T04:48:47.432065Z","title":"Alphaexploitem: Going beyond the nash equilibrium in poker by learning to exploit suboptimal play","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.432065Z"},"links":{"cited_paper":"/paper/2605.09150","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c3883bf23dab765e082dbf50253f97c417d5cb9c346ce7fa5b245b540f83ddae","observation_id":"edab7af6-664e-45a1-9e30-6ea0d1086392","resolution":{"observed_at":"2026-07-31T04:48:47.432065Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.434218Z","title":"A survey of opponent modeling in adversarial domains","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.434218Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:dd6488af1d2f1809c5be5ea98a02924bef292911df76732d1ae47a70eab69672","observation_id":"09c48c32-efe3-4b33-b1e2-617230f6affc","resolution":{"observed_at":"2026-07-31T04:48:47.434218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.439864Z","title":"Mcrnr: fast computing of restricted nash responses by means of sampling","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":101,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.439864Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:2f4cb77fb6f9be37912020432ccc7d630bb6df7123ccc8c1dede8572fc20d0f9","observation_id":"7a046576-bd18-4296-8cf7-25024d858675","resolution":{"observed_at":"2026-07-31T04:48:47.439864Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.442292Z","title":"Bayes' bluff: opponent modelling in poker","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":102,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.442292Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:52204e31cef31177ab5c5f3182e66b345210c361246216dff7d1b4a4fa0b4c15","observation_id":"0d18b556-dd6b-4e35-af1d-c69fbfce0f8b","resolution":{"observed_at":"2026-07-31T04:48:47.442292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.444636Z","title":"Learning not to regret","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":103,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.444636Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:766b4f3139c6bf28c06db9a1fdc0f334a93c98837513bec5046771ea00c5816f","observation_id":"94963f48-b4a2-4389-9d1a-8a0bb24c26c4","resolution":{"observed_at":"2026-07-31T04:48:47.444636Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1407.5042","last_updated":"2014-07-18T15:41:28Z","snapshot_observed_at":"2026-08-14T23:26:13.593408Z","submitted_at":"2014-07-18T15:41:28Z","title":"Solving Large Imperfect Information Games Using CFR+","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1407.5042","snapshot_observed_at":"2026-07-31T04:48:47.446905Z","title":"Solving large imperfect information games using cfr+","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":104,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.446905Z"},"links":{"cited_paper":"/paper/1407.5042","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:42de1ec9a4fef124efa866f52d9d4d5efe2e348f9f95deae58bcf7c3d9aa1b30","observation_id":"9a177a92-e23d-4ab5-8f17-720b4d78d21e","resolution":{"observed_at":"2026-07-31T04:48:47.446905Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.449196Z","title":"Horse-cfr: Hierarchical opponent reasoning for safe exploitation counterfactual regret minimization","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":105,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.449196Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:aad0169248c62f73572139f2bb068c92c10fe12a5e7c37034c8730f93081e632","observation_id":"2aa92267-0347-4619-afe2-1fcd46db5804","resolution":{"observed_at":"2026-07-31T04:48:47.449196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.451676Z","title":"Dynamic discounted counterfactual regret minimization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":106,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.451676Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:5c624a3dffe55679dff39e40788aa8fe16746288ca7b1ff10a26aaff65cfff8b","observation_id":"2f52a83a-c96d-400c-9519-9c37b0cf30ac","resolution":{"observed_at":"2026-07-31T04:48:47.451676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.453889Z","title":"Deep (predictive) discounted counterfactual regret minimization","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":107,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.453889Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0ef3ac30e10cff9e3db9c81abe570fa29525aeb4174c9bc1389451d6a3e2d541","observation_id":"13962ccc-08b9-4604-b269-c4817a0a3490","resolution":{"observed_at":"2026-07-31T04:48:47.453889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.456192Z","title":"Regret minimization in games with incomplete information","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":108,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.456192Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:9d7a9e5d73a39338ad0d6d6085fe260f9fe3cba526a44fc162af4c8151159c79","observation_id":"3e8b75cc-0ebe-43d7-912f-2e57e9028432","resolution":{"observed_at":"2026-07-31T04:48:47.456192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","latest_version":1,"primary_category":"cs.GT","snapshot_observed_at":"2026-08-09T00:45:19.552928Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":100,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":100},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"thesis":"As of 15 August 2026, this Paper Citation Record lists 100 of 100 outbound references and 0 inbound Pith citation observations for arXiv:2607.28520."}