{"as_of":"2026-08-19T01:34:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:530acf5e0240db11b19541c735c99b9d64103c9a4526f3aa077cd66d88a5a8c0","coverage":[{"denominator":29,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":29,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T20:27:04.535925Z","state":"measured"},{"denominator":29,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":29,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2508.10804/citation-record","integrity":"/paper/2508.10804/integrity","json":"/paper/2508.10804/citation-record.json","paper":"/paper/2508.10804"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.985356Z","title":"Whittle’s index policy for a multi-class queueing system with convex holding costs","venue":null,"work_id":"b4bdcc96-a551-422b-a6b2-ed7e99227fa7","year":2003},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:01.997873Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:86bb7499946befe9644e71fb2543fbd166b7996eb8bae8078194f6352aaaf858","observation_id":"e219ae47-4f7b-4d24-bdd4-15a4dfab993d","resolution":{"observed_at":"2026-08-05T20:27:10.088833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.716010Z","title":"Dynamic allocation indices for restless projects and queueing admission control: a polyhedral approach","venue":null,"work_id":"6fba17ef-dfe8-4ed5-8cd3-a82ba5381d7b","year":2002},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.074608Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:17b518e6ce3236564b40466cab96896a70bd11a951f6da9ee4ed1f067077685c","observation_id":"57191e67-757e-47ab-a81c-20927c5d9c1f","resolution":{"observed_at":"2026-08-05T20:27:09.843742Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.479525Z","title":"Field study in deploying restless multi- armed bandits: Assisting non-profits in improving maternal and child health","venue":null,"work_id":"f723d386-1170-46a4-afd6-386604a6c0ec","year":2022},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.136881Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:5806bb0273a4113b210d94d1022f69000f4beca1dcf2955dcd8680013a762ff7","observation_id":"61adbc05-dd6d-496b-83aa-e821ae97042b","resolution":{"observed_at":"2026-08-05T20:27:09.581931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.270612Z","title":"Collapsing bandits and their application to public health intervention","venue":null,"work_id":"26660420-3e32-4369-89dd-681c07f0d4b0","year":2020},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.205380Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:646e8b0167fb892f081bc35645a7610e26458bfd90a7b6e4b9fc7c3179f208f9","observation_id":"5634b63f-55e1-4272-afdb-a302923553bc","resolution":{"observed_at":"2026-08-05T20:27:09.368472Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.018491Z","title":"Distributed optimal relay selection in wireless cooperative networks with finite-state markov channels","venue":null,"work_id":"36744545-133b-4338-b702-98965850ff3e","year":2010},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.269087Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:b7153f965c033c3a94cf2f9f29bb2d35b85985f3cec273fac4de2e42b527712b","observation_id":"36e713a6-2481-4cb5-b13c-0b44269523c4","resolution":{"observed_at":"2026-08-05T20:27:09.143708Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.836333Z","title":"Cell association with user behavior awareness in heterogeneous cellular networks","venue":null,"work_id":"da1b66b4-b1f7-4963-a827-7395a57c7b3a","year":2018},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.338129Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:d7b85d82dc72c67f01e15b1322bbbeb9b7ad25ab9eeb9d1cda774c8d0952d774","observation_id":"6531ab99-6dca-4f7b-bccf-825758c92033","resolution":{"observed_at":"2026-08-05T20:27:08.917036Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.623770Z","title":"Optimistic whittle index policy: Online learning for restless bandits","venue":null,"work_id":"452b05e6-8ddd-418f-bcc7-e413f63be622","year":2023},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.408278Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:0349155562423f56721aec3022ef8f70a25df42274ee5dd15a472a11bd155172","observation_id":"a863df58-7f4a-49c3-a485-c81a1ac39673","resolution":{"observed_at":"2026-08-05T20:27:08.723592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.320710Z","title":"Reinforcement learning for non- stationary markov decision processes: The blessing of (more) optimism","venue":null,"work_id":"f1a69b91-eb31-494d-abca-249f17f2b2ac","year":2020},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.473859Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:e383a5d2407a35a040da26fc5c7941961b38fba0136171875ad6e24b7eb3c5b7","observation_id":"489a499a-f7a1-41d8-a302-59a33b08cc3f","resolution":{"observed_at":"2026-08-05T20:27:08.424887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.121769Z","title":"Regret bounds for thompson sampling in episodic restless bandit problems","venue":null,"work_id":"f795ad79-e64e-47b1-8d2b-eacf348bb136","year":2019},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.566224Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:80609f9935b617993e10fb85d624584134a377b31adfd759a57fc1a488fd40b1","observation_id":"1ef18fa3-6690-455b-8141-7c7e02bd0c39","resolution":{"observed_at":"2026-08-05T20:27:08.210494Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1910.05654","last_updated":"2019-10-12T22:30:24Z","snapshot_observed_at":"2026-08-18T20:57:43.177251Z","submitted_at":"2019-10-12T22:30:24Z","title":"Thompson Sampling in Non-Episodic Restless Bandits","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.05654","snapshot_observed_at":"2026-08-05T20:27:02.648338Z","title":"Thompson sampling in non-episodic restless bandits","venue":null,"work_id":null,"year":1910},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.648338Z"},"links":{"cited_paper":"/paper/1910.05654","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:f044bf76d3a5d8e5669dcb349a513f1d568f839b528c830751e6479527ff916e","observation_id":"51d51be1-e07d-41b9-9b0f-d96e31b4ed55","resolution":{"observed_at":"2026-08-05T20:27:02.648338Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:02.738561Z","title":"Restless bandits: Activity allocation in a changing world","venue":null,"work_id":null,"year":1988},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.738561Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:838adafff486fd6c4beb1fc68acb596c2afe136b3615ea48657e7f61a36decad","observation_id":"95c17267-f2e2-44d0-9843-112aafd5f792","resolution":{"observed_at":"2026-08-05T20:27:02.738561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.939720Z","title":"Introduction to multi-armed bandits","venue":null,"work_id":"e9a6dc67-c62d-4f9e-886b-d727b443a5d0","year":2019},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.810158Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:d18c0c783f067cb3045a9147d97c5c8f47b493c25168ede8de426d216a2627a9","observation_id":"1a471b5c-7ae8-408f-aa8e-197778d1fcc4","resolution":{"observed_at":"2026-08-05T20:27:08.064518Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1402.6028","last_updated":"2014-02-25T01:34:43Z","snapshot_observed_at":"2026-08-17T00:30:50.886580Z","submitted_at":"2014-02-25T01:34:43Z","title":"Algorithms for multi-armed bandit problems","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1402.6028","snapshot_observed_at":"2026-08-05T20:27:02.916543Z","title":"Algorithms for multi-armed bandit problems.arXiv preprint arXiv:1402.6028, 2014","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.916543Z"},"links":{"cited_paper":"/paper/1402.6028","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:3ee3293364a0a60f3984d40be147a1f803a7be5f52a24acc5e14e718239b2e54","observation_id":"579c20fe-caee-4318-9d21-6dc70c8324f4","resolution":{"observed_at":"2026-08-05T20:27:02.916543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:02.998565Z","title":"The complexity of optimal queuing network control","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.998565Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:3b6e850f6b9ec0f93a8de0ec3bc113533adfd966e5607936aad7a19083d6feb8","observation_id":"1d4e2415-cdb7-434e-a79d-f3ae2fa2cc28","resolution":{"observed_at":"2026-08-05T20:27:02.998565Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.722009Z","title":"On an index policy for restless bandits","venue":null,"work_id":"67ad0e29-f01b-4835-af1d-cd4d6470ca3f","year":1990},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.095698Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:0bc60f10476d648189e27134401c3a62e00c656a6a695646fc43f47d8ec1815c","observation_id":"41cbc4cd-3155-420d-bd28-8e757f3872a2","resolution":{"observed_at":"2026-08-05T20:27:07.814489Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.503997Z","title":"A restless bandit formulation of opportunistic access: Indexablity and index policy","venue":null,"work_id":"5a6ca861-579f-49b9-bf3c-ed235d998ccb","year":2008},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.149894Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:9c2f31baa183f6f522a3285bb2553ae681ec5b727e4492dbb7391fd3152b1ff2","observation_id":"fa63806b-c50f-40da-9175-ee81a4a4f7e7","resolution":{"observed_at":"2026-08-05T20:27:07.643212Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.231599Z","title":"Indexability of restless bandit problems and optimality of whittle index for dynamic multichannel access.IEEE Transactions on Information Theory, 56(11):5547– 5567, 2010","venue":null,"work_id":"62235149-0607-4221-8dfc-3ea272b371f2","year":2010},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.152865Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:5027278c180466117b524c708588be0bd027f856032e431ac3048cfa1433de72","observation_id":"04aad0be-57c6-4ceb-9d86-78f7a112e136","resolution":{"observed_at":"2026-08-05T20:27:07.327107Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.039487Z","title":"Qwi: Q-learning with whittle index","venue":null,"work_id":"4038ff46-9dee-4bbf-8a72-d75977216c82","year":2022},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.186474Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:f12025f62a533fd7df81481e03f3fbd59019584a373439210c5dd2882a4caf62","observation_id":"3771e6c2-68db-4240-9077-3604d2a8ee23","resolution":{"observed_at":"2026-08-05T20:27:07.145281Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2105.07965","last_updated":"2021-07-22T22:37:04Z","snapshot_observed_at":"2026-08-18T14:21:14.654350Z","submitted_at":"2021-05-17T15:44:55Z","title":"Learn to Intervene: An Adaptive Learning Policy for Restless Bandits in Application to Preventive Healthcare","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2105.07965","snapshot_observed_at":"2026-08-05T20:27:03.273298Z","title":"Learn to intervene: An adaptive learning policy for restless bandits in application to preventive healthcare","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.273298Z"},"links":{"cited_paper":"/paper/2105.07965","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:3d699ec286761903d21db977b67d6a1044740afe7c6a751de0a988afb5d2e114","observation_id":"f69f6a6f-b8de-421c-a133-d25ac6296dcc","resolution":{"observed_at":"2026-08-05T20:27:03.273298Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:03.350381Z","title":"Towards q-learning the whittle index for restless bandits","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.350381Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:b391aa5488dcaf9a8cd2e09fa9b20c409ba8bf26efb32e74f8b5266c49fb19fa","observation_id":"96ce33c7-4b11-4ccf-9b88-4ea9e493c82a","resolution":{"observed_at":"2026-08-05T20:27:03.350381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:06.792736Z","title":"Near-optimal regret bounds for reinforcement learning","venue":null,"work_id":"7228532c-b657-410b-a7f8-d9b5337406b3","year":2008},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.411313Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:9e520dc0bf1b28242ae9b24d902866f37fd9f3191d671dda5928d7889a783a11","observation_id":"93a38f69-b969-433e-97bb-cf564ef7329d","resolution":{"observed_at":"2026-08-05T20:27:06.923679Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:06.561096Z","title":"Logarithmic online regret bounds for undiscounted reinforcement learning","venue":null,"work_id":"8516cbda-a1e5-4887-8dc8-e706b661a832","year":2006},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.496746Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:8255c931768b4718525ad2c29999fd82a7b81aa81180e8bbc3d93cb7099f7db9","observation_id":"18741a19-acf4-45f8-aa35-a7167068442e","resolution":{"observed_at":"2026-08-05T20:27:06.669105Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1205.2661","last_updated":"2012-05-09T14:47:06Z","snapshot_observed_at":"2026-08-15T05:14:28.809312Z","submitted_at":"2012-05-09T14:47:06Z","title":"REGAL: A Regularization based Algorithm for Reinforcement Learning in Weakly Communicating MDPs","version":1},"cited_work":{"arxiv_id":"1205.2661","doi":null,"metadata_source":"pith","pith_arxiv_id":"1205.2661","snapshot_observed_at":"2026-08-05T20:27:05.152416Z","title":"REGAL: A Regularization based Algorithm for Reinforcement Learning in Weakly Communicating MDPs","venue":"cs.LG","work_id":"edba61d7-2d79-44ad-8daa-52c1358d4fac","year":2012},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.662892Z"},"links":{"cited_paper":"/paper/1205.2661","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:80c4b2bb604dbdf7b674a00d110d87647ef6d2283aab50cc776d39893c78be40","observation_id":"c88f8d83-b79f-41fa-9a46-20deb7bd4c54","resolution":{"observed_at":"2026-08-05T20:27:05.256092Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:06.337524Z","title":"Non-stationary reinforcement learning without prior knowledge: An optimal black-box approach","venue":null,"work_id":"20eba94e-2d9a-489d-bf14-26f0879b7158","year":2021},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.815378Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:c4acc84913ba224e8dbc890eed6ed6e79779d6a2e865b4c38bb1cb592a362878","observation_id":"5325003f-4bfd-45f6-9c7d-5bfc62280e21","resolution":{"observed_at":"2026-08-05T20:27:06.481285Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.10066","last_updated":"2018-05-25T10:14:20Z","snapshot_observed_at":"2026-08-14T19:11:12.511581Z","submitted_at":"2018-05-25T10:14:20Z","title":"A Sliding-Window Algorithm for Markov Decision Processes with Arbitrarily Changing Rewards and Transitions","version":1},"cited_work":{"arxiv_id":"1805.10066","doi":null,"metadata_source":"pith","pith_arxiv_id":"1805.10066","snapshot_observed_at":"2026-08-05T20:27:04.973913Z","title":"A Sliding-Window Algorithm for Markov Decision Processes with Arbitrarily Changing Rewards and Transitions","venue":"cs.LG","work_id":"2a6b0248-4688-4ec8-8b5f-b0689f71196e","year":2018},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.035674Z"},"links":{"cited_paper":"/paper/1805.10066","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:47ec4d6340e995248b6aa22632c698f02681a66cf4adfe79601f2a352cd4f113","observation_id":"de96fa00-8325-4043-88b8-9c607e7228b5","resolution":{"observed_at":"2026-08-05T20:27:05.038455Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:05.980928Z","title":"Near-optimal model-free reinforcement learning in non-stationary episodic mdps","venue":null,"work_id":"212e4708-4749-4e12-bcfa-df19509b99e2","year":2021},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.203518Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:86c5005bcae2f71539476901118a9ad5f344a8eceeefa9142b383e706d14ca74","observation_id":"0de02645-306b-4cd3-b529-8919694cc5a7","resolution":{"observed_at":"2026-08-05T20:27:06.110026Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:05.655111Z","title":"Learning in a changing world: Restless multiarmed bandit with unknown dynamics","venue":null,"work_id":"5ed97db3-c355-43f3-8b19-b4783724de84","year":1902},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.329643Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:338377edb2ee7eb3ebf8607fd2feb5e0d430e50591a919c3568fb3b29fafd49a","observation_id":"c11553c4-59ca-472e-ab2e-dfb59c929753","resolution":{"observed_at":"2026-08-05T20:27:05.833904Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2002.05138","last_updated":"2021-05-27T22:32:32Z","snapshot_observed_at":"2026-08-11T06:34:08.209616Z","submitted_at":"2020-02-12T18:30:09Z","title":"Regret Bounds for Discounted MDPs","version":3},"cited_work":{"arxiv_id":"2002.05138","doi":null,"metadata_source":"pith","pith_arxiv_id":"2002.05138","snapshot_observed_at":"2026-08-05T20:27:04.675708Z","title":"Regret Bounds for Discounted MDPs","venue":"cs.LG","work_id":"2b4642f7-bd2e-4ffb-a3e2-295d7f237732","year":2020},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.441499Z"},"links":{"cited_paper":"/paper/2002.05138","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:f66f659d72930880d5a0e7f93377afcc19d1b5cbf95adbeb82dd17d9bbed092e","observation_id":"0ba162fe-673c-4ab2-8097-ecca8b779c80","resolution":{"observed_at":"2026-08-05T20:27:04.837991Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:05.428822Z","title":"X t∈N γt−1 R(st, at) − λ∗ t πt(st)⊤1 − K | s1 = s # (31) ≤ E(s,a)∼(P ,πt)","venue":null,"work_id":"de035725-9ffa-4c4b-a600-e0fe835ecf82","year":1982},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.535925Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:dcc13d7cd67ba10612d4aaa65f992aafab8c6758e9df4561b8c2ed9797b11f7f","observation_id":"d81409f8-2fdf-4991-bbee-084a78a9f3b2","resolution":{"observed_at":"2026-08-05T20:27:05.546835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee"},"reference_resolution":{"displayed":29,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":6,"verified_exact":3,"verified_fuzzy":20},"total_outbound_references":29},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 29 of 29 outbound references and 0 inbound Pith citation observations for arXiv:2508.10804."}