{"as_of":"2026-08-16T20:46:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:30c7450c37924ac5351f9120a77db144baa0ddd167fcb2a2a933923c5cf966a0","coverage":[{"denominator":90,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":90,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-14T12:58:10.702721Z","state":"measured"},{"denominator":91,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":91,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-07T12:39:22.493228Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":2,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"cited_work":{"arxiv_id":"1908.06158","doi":"10.48550/arxiv.1908.06158","metadata_source":"arxiv_reference","pith_arxiv_id":"1908.06158","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint(2019), 1–8","venue":"arXiv (Cornell University)","work_id":"e206b7d3-221d-4985-a1b9-52f56805f5cf","year":2019},"citing_paper":{"arxiv_id":"2604.26651","last_updated":"2026-04-29T13:18:48Z","snapshot_observed_at":"2026-07-06T23:12:15.279847Z","submitted_at":"2026-04-29T13:18:48Z","title":"The Bandit's Blind Spot: The Critical Role of User State Representation in Recommender Systems","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-07T12:39:22.493228Z"},"links":{"cited_paper":"/paper/1908.06158","citing_paper":"/paper/2604.26651"},"observation_digest":"sha256:cf116adaeec27da3c6ee5243279e2c06ac62782b4d255de57debd8db88e17576","observation_id":"a47198e2-a32f-4ef1-9011-26623dd1d837","resolution":{"observed_at":"2026-05-09T03:45:06.057153Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/1908.06158/citation-record","integrity":"/paper/1908.06158/integrity","json":"/paper/1908.06158/citation-record.json","paper":"/paper/1908.06158"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.298711Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.298711Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:8f51c7336ceaca046d4dd17c1c108661e4e5e45c951228377b6a8d90207be891","observation_id":"aba17654-437b-4555-80f1-36d01663fbd5","resolution":{"observed_at":"2026-08-14T12:58:10.298711Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.303254Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.303254Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:54cf2ed4218cb3ceb5d29762755b6113ea55436b85f0649db5ef43d45c867f0f","observation_id":"9f402ae9-b7a1-4764-ba31-ff8949dd655f","resolution":{"observed_at":"2026-08-14T12:58:10.303254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1612.06246","last_updated":"2017-06-06T03:21:09Z","snapshot_observed_at":"2026-08-16T07:36:12.130019Z","submitted_at":"2016-12-19T16:17:56Z","title":"Corralling a Band of Bandit Algorithms","version":3},"cited_work":{"arxiv_id":"1612.06246","doi":null,"metadata_source":"pith","pith_arxiv_id":"1612.06246","snapshot_observed_at":"2026-08-14T12:58:11.148120Z","title":"Corralling a Band of Bandit Algorithms","venue":"cs.LG","work_id":"60e3ec14-e098-4e65-87e0-c66a8f32a6a2","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.307201Z"},"links":{"cited_paper":"/paper/1612.06246","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:0a3d2f09fa696324c746930cf6691561ed24e6b4bec52b56b92dee1ed778dfcf","observation_id":"cf7b2ebc-cd8e-42d6-adaf-f81994f11451","resolution":{"observed_at":"2026-08-14T12:58:11.153305Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.312082Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.312082Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:1f2b4faf6315144a966ac6984b990e6f12e7f1299197efb5cd657c7ff6908abd","observation_id":"d40ce957-5b32-4ee5-9716-6720b4e91395","resolution":{"observed_at":"2026-08-14T12:58:10.312082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.317384Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.317384Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:7200a101d79bc4a8dd03a2d91ee19eb1df69937cdeb65518cf6457268f0db3df","observation_id":"67f81f28-b173-499b-9f40-0996d5022765","resolution":{"observed_at":"2026-08-14T12:58:10.317384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.321488Z","title":null,"venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.321488Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:ff6d75cfb2ee7e4d262f2d3570d79165e9155599644c8eafe503445636e6955f","observation_id":"e03fecca-9880-41f6-b5ad-ea5b7784bbf5","resolution":{"observed_at":"2026-08-14T12:58:10.321488Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.326189Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.326189Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:c77eebd0d6458ae21431bcc36a7390959c280f1abb9a6c6b413b27019032e85f","observation_id":"f7e1446c-bd47-4469-8289-9b803f463711","resolution":{"observed_at":"2026-08-14T12:58:10.326189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.329965Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.329965Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:0e8113d2109fde04c44aee1a762314d50fe70d40a4f93d2edeba234a6d20dbd0","observation_id":"de8320ea-b242-4778-9446-cfcdc56a7ac2","resolution":{"observed_at":"2026-08-14T12:58:10.329965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.333717Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.333717Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:d9f364e0b1ed0b5f7e80bd9a3a5b8b6cf68e8ee7d9c6d142e0535cd1c375bb12","observation_id":"7ec565d4-0227-4b05-a67d-683e4ce68d89","resolution":{"observed_at":"2026-08-14T12:58:10.333717Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1711.02317","last_updated":"2018-03-13T07:30:42Z","snapshot_observed_at":"2026-08-14T20:16:15.603878Z","submitted_at":"2017-11-07T07:10:47Z","title":"Multi-Player Bandits Revisited","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.02317","snapshot_observed_at":"2026-08-14T12:58:10.341524Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.341524Z"},"links":{"cited_paper":"/paper/1711.02317","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:470a616d5e0d55772bc19b2b5f169f26a0386354061891cdbf546d0416bd6db5","observation_id":"94174669-ffa7-4596-8043-94395e843f94","resolution":{"observed_at":"2026-08-14T12:58:10.341524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.345805Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.345805Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:de443c13c66cf961edcccf68dee4d3686ae255a34b411c0b29842b4e62bc45d8","observation_id":"ab86f3be-29e5-4bc9-acc8-4a2fd3b65d88","resolution":{"observed_at":"2026-08-14T12:58:10.345805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1809.08151","last_updated":"2019-11-19T09:51:23Z","snapshot_observed_at":"2026-08-14T18:25:18.446596Z","submitted_at":"2018-09-21T14:43:46Z","title":"SIC-MMAB: Synchronisation Involves Communication in Multiplayer Multi-Armed Bandits","version":4},"cited_work":{"arxiv_id":"1809.08151","doi":null,"metadata_source":"pith","pith_arxiv_id":"1809.08151","snapshot_observed_at":"2026-08-14T12:58:11.116387Z","title":"SIC-MMAB: Synchronisation Involves Communication in Multiplayer Multi-Armed Bandits","venue":"cs.LG","work_id":"224f197d-e1e0-4708-bf5e-04cf93640d73","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.349591Z"},"links":{"cited_paper":"/paper/1809.08151","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:4111b120d8acac0aa4c02389ae26f0741554b7e52ec7bdbd268ff798ea4d404a","observation_id":"fa106a21-feb7-4e18-a115-d5839e15e3c6","resolution":{"observed_at":"2026-08-14T12:58:11.121476Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.742644Z","title":null,"venue":null,"work_id":"63dd25fe-81ce-42a1-93ab-44e421cc3098","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.353619Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:02f8708742bae139a418ba4d404647d2a58c12a9bdb5d18f71ee2309562576dd","observation_id":"bab0c416-30a9-4a9d-81b7-9c783be72bdb","resolution":{"observed_at":"2026-08-14T12:58:11.746254Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.357683Z","title":null,"venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.357683Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:6e7510860ef537f30c4ffa670619322212d5ea4d3df64aef5f80c441a7df4080","observation_id":"983e3169-3e15-4f4d-ac0c-b2d243887c46","resolution":{"observed_at":"2026-08-14T12:58:10.357683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1802.03692","last_updated":"2019-01-24T06:13:01Z","snapshot_observed_at":"2026-08-14T19:47:08.194600Z","submitted_at":"2018-02-11T04:57:28Z","title":"Nearly Optimal Adaptive Procedure with Change Detection for Piecewise-Stationary Bandit","version":4},"cited_work":{"arxiv_id":"1802.03692","doi":null,"metadata_source":"pith","pith_arxiv_id":"1802.03692","snapshot_observed_at":"2026-08-14T12:58:11.097257Z","title":"Nearly Optimal Adaptive Procedure with Change Detection for Piecewise-Stationary Bandit","venue":"stat.ML","work_id":"1b488b98-1874-4d59-a30b-eb5c9a65b213","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.362024Z"},"links":{"cited_paper":"/paper/1802.03692","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:9c4805c8a96c6105b274faa778f57d76d81eaf53044b23005352198566993ca9","observation_id":"5184c622-1c77-47f9-b204-68a430d670da","resolution":{"observed_at":"2026-08-14T12:58:11.102725Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.722581Z","title":null,"venue":null,"work_id":"096aad5e-d9c0-4071-b0d8-d9ceb3eff876","year":2013},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.366254Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:a22680958e1bccb637c35a70a0c46461f09dad35b482ead02f44220751b5baea","observation_id":"787c3603-db2b-4f80-aff2-e3c6582c8fb7","resolution":{"observed_at":"2026-08-14T12:58:11.727057Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.711962Z","title":null,"venue":null,"work_id":"c026c087-b187-45bb-9dd2-4036d66b08d4","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.370195Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:d18b5ce79c5993ff1362be78152f1b8851ffe63c87f962a72e4d6e13e88d5f30","observation_id":"0f04e1e7-6fef-4e8c-b84b-6fd6a1dcf308","resolution":{"observed_at":"2026-08-14T12:58:11.715050Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.701368Z","title":null,"venue":null,"work_id":"8c69571c-c052-4384-b27e-22eac1e4dedf","year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.374414Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:dca2969d196062179f512f5fdceff1b4f1cab62ad3a5c2f06c1aa6cbb081f489","observation_id":"74573683-d389-448c-9649-b16e726ad6e5","resolution":{"observed_at":"2026-08-14T12:58:11.704737Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.690417Z","title":null,"venue":null,"work_id":"6c157506-1763-412b-8f39-1f574f5d00b2","year":2011},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.379029Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:38c7ac09b3e51992841fe390ddbcfb358b883322384af79d14c089936099f36b","observation_id":"f60efb19-094c-45f2-bd22-71e58b700e67","resolution":{"observed_at":"2026-08-14T12:58:11.694294Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.678375Z","title":null,"venue":null,"work_id":"346b97ac-5bc9-45d4-a0c5-fb6ac75d5474","year":2015},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.390401Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:ecb1f3071c88dbbfac0b9df1ef706f31c0e3e5d0ee8d1e5f48f88b4f554acacd","observation_id":"c2d7b1d7-ae03-450d-9502-c074109ca74a","resolution":{"observed_at":"2026-08-14T12:58:11.682686Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.395172Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.395172Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:26dc8393bb2736da0225eec560e5406c72bf1537c27f17847883212b7d4e0769","observation_id":"01df47e5-6f5f-44bb-80e4-24b272f43bd5","resolution":{"observed_at":"2026-08-14T12:58:10.395172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.649671Z","title":null,"venue":null,"work_id":"9d500727-dc17-4946-8127-694eb5ba0577","year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.404416Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:37fe3dd51269d065bb7e973d342d8c0948085d6fd023f8000a0818d127050b92","observation_id":"e606e182-ca84-4331-ba9a-021f0a943de7","resolution":{"observed_at":"2026-08-14T12:58:11.655135Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.637097Z","title":null,"venue":null,"work_id":"4c1941a4-4f6d-4fd5-a424-c5d0c7a47d2c","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.408628Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:889bfdcf759cbc164cde0403636a831bccdfc4bca7dd3774f3d1d75cc3357e3c","observation_id":"48a6d36b-cebf-49e9-8a37-6d24c9a10eab","resolution":{"observed_at":"2026-08-14T12:58:11.641176Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.413155Z","title":null,"venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.413155Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:965d01b3090f13ed38b7d7b713f5927affaf6c31329d7ff9bde3f4ac6b6f4f9c","observation_id":"84693e71-c8ed-40b6-9433-bdc5f4f1174a","resolution":{"observed_at":"2026-08-14T12:58:10.413155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.617701Z","title":null,"venue":null,"work_id":"6f8868cf-6def-4345-8d93-7f731648e626","year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.418079Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:af0557d59414403f56a05cf6fab41cb00b97c9329893991806288e9bc9345e3d","observation_id":"e654bc4b-e826-4dc0-9e2a-c668f79b35ad","resolution":{"observed_at":"2026-08-14T12:58:11.621617Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.606639Z","title":null,"venue":null,"work_id":"28eda030-cbf9-4595-9a87-f475e70b3333","year":2010},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.422464Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:7989d362c05598bbb739d02985e2ad3c15472d5cd6be0170ac5bfbd5930f81e7","observation_id":"0c29a500-222c-427a-a036-72c7ec59171c","resolution":{"observed_at":"2026-08-14T12:58:11.610300Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.595001Z","title":null,"venue":null,"work_id":"3b084318-c5a6-4fc2-bf4f-5da8fdfaa913","year":2006},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.426760Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:a4dac1f31746d5e394e0d920a67ea9ce2ef6c7ae8adcbdadbbe1a18c52848c37","observation_id":"9181f9a9-0bec-4a1b-97a4-d4869d4bade6","resolution":{"observed_at":"2026-08-14T12:58:11.598878Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.583072Z","title":null,"venue":null,"work_id":"6d9056d6-8bfc-4d88-ab8e-be9228f384a1","year":2011},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.430569Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:87dcd9698ecd6b966c0210e193696971a0de818d242c16a7555d126e595eead8","observation_id":"75ccac40-3174-4daa-a75a-f5356000dc67","resolution":{"observed_at":"2026-08-14T12:58:11.587222Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.05071","last_updated":"2022-07-01T10:12:30Z","snapshot_observed_at":"2026-08-16T14:47:24.419240Z","submitted_at":"2018-05-14T09:05:10Z","title":"KL-UCB-switch: optimal regret bounds for stochastic bandits from both a distribution-dependent and a distribution-free viewpoints","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.05071","snapshot_observed_at":"2026-08-14T12:58:10.434294Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.434294Z"},"links":{"cited_paper":"/paper/1805.05071","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:2511cfc29e94a59114800caa8da51287c63c39e6145c1da68072fd912604b749","observation_id":"35c38755-4f24-43ff-8b16-d53c0997f5dc","resolution":{"observed_at":"2026-08-14T12:58:10.434294Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.572161Z","title":null,"venue":null,"work_id":"a8a1c276-2c67-4a77-ac71-16aca20a8115","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.438801Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:b7048e5238e6bf250be7e9cac21096ba2046a505c7dfffd11c3695d518fa50b8","observation_id":"54de7c0e-6eb3-4903-9a95-63baf3bff53a","resolution":{"observed_at":"2026-08-14T12:58:11.575597Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.442971Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.442971Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:83c537992f629a134175fb8cc49582d1a2693227d266c2dffdcf879598e603db","observation_id":"84f5ab6d-bcd5-4ba4-b5c6-5fc4104aba82","resolution":{"observed_at":"2026-08-14T12:58:10.442971Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.555116Z","title":null,"venue":null,"work_id":"fbe99c36-a60e-4198-a7ce-0e40b8fae98d","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.447008Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:df1906728cea6363fbff4820d08cbadf49421320e6a61378e9c2ac49d10a3e6b","observation_id":"f64472d8-9e0b-465d-b148-4bf6b0745edd","resolution":{"observed_at":"2026-08-14T12:58:11.558555Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.544321Z","title":null,"venue":null,"work_id":"4f860faf-eef4-4656-ac36-3cb7bfca3053","year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.451225Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:1e7d1dae7e00c53eeab0fe0efcddee958e40bdbefdbc67fa7d54ee6a930d6636","observation_id":"af613d35-55e6-4588-9649-0566fbad91d7","resolution":{"observed_at":"2026-08-14T12:58:11.547707Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.455911Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.455911Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:5af39e47e83cf3f2a6da8494fa77eb76798d03c26319faec1ea750a6bb712369","observation_id":"ed73e1e0-9598-4bc4-91d2-e0d6db537bfb","resolution":{"observed_at":"2026-08-14T12:58:10.455911Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.527393Z","title":null,"venue":null,"work_id":"a57bcd0e-5516-4e53-9788-f72340f7d008","year":2015},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.460194Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:9c836c2f5ec6ce03848c791504e99b640d1928655c4b4cd4d28d9ef8045c02d4","observation_id":"9d3e6c4d-7139-4086-9fc2-a72d6bac1cce","resolution":{"observed_at":"2026-08-14T12:58:11.530631Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.517028Z","title":null,"venue":null,"work_id":"3f2d7dd2-91e7-4ed7-a0e9-6e96cddfd7c4","year":null},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.464282Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:39c7e93daa5619c8c16533675a7b24cd6c0d988d29ce99b426e13ad52c251ef5","observation_id":"21a6ff0d-095d-4fe0-8ee9-8c11ddd613ff","resolution":{"observed_at":"2026-08-14T12:58:11.520553Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.493843Z","title":null,"venue":null,"work_id":"948cf45a-2744-4aaf-b73e-a8d7d0e90532","year":2015},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.473480Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:fdb412af6857f25845375e0d36324e9bc6a345d0cda57183f2ad63c37434cd2c","observation_id":"829bdff1-96a0-4d79-a7ae-29199e25ca2f","resolution":{"observed_at":"2026-08-14T12:58:11.497756Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.477728Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.477728Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:f47590ab3514a52ba90cadcf6165e738e235fd65b99616c1771802db9a8dad57","observation_id":"2e6967c8-66be-4db7-a725-0f93e0d03706","resolution":{"observed_at":"2026-08-14T12:58:10.477728Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.476100Z","title":null,"venue":null,"work_id":"3d3a522c-0681-4767-92f6-edf2083d1706","year":2014},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.482061Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:e84dea5a4fbed773bdea4ea33aa4e08f30826d20025f62540d9199610f32d9e4","observation_id":"6029c14d-419d-4714-8fab-5846689732c6","resolution":{"observed_at":"2026-08-14T12:58:11.479766Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1902.00610","last_updated":"2019-12-10T07:38:51Z","snapshot_observed_at":"2026-08-14T17:22:05.052574Z","submitted_at":"2019-02-02T01:20:49Z","title":"On the Optimality of Perturbations in Stochastic and Adversarial Multi-armed Bandit Problems","version":4},"cited_work":{"arxiv_id":"1902.00610","doi":null,"metadata_source":"pith","pith_arxiv_id":"1902.00610","snapshot_observed_at":"2026-08-14T12:58:11.066971Z","title":"On the Optimality of Perturbations in Stochastic and Adversarial Multi-armed Bandit Problems","venue":"stat.ML","work_id":"c45407aa-40f8-4487-9dc2-47d8ed6ae63e","year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.486453Z"},"links":{"cited_paper":"/paper/1902.00610","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:7a7b895bc48aa7367859ad61596352020ef4f2ad2e9a133994551d526087cdd1","observation_id":"d308b056-e60a-4688-beb0-e82de0494917","resolution":{"observed_at":"2026-08-14T12:58:11.072063Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.465316Z","title":"Nakagawa H","venue":null,"work_id":"d8df7076-728f-41ba-b113-232961563d58","year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.491170Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:6a089e36c366dbdb70e5354892e780c7df783d669812965da3f76af5ac6100d1","observation_id":"da43aab4-ff9b-4edb-95ae-09262780cc51","resolution":{"observed_at":"2026-08-14T12:58:11.469296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.453751Z","title":null,"venue":null,"work_id":"12f1808f-98c1-47ab-8215-5d4a163102e8","year":2011},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.495448Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:edc63e29f0c802b7cb73cac8307558c8b48aedb32cba95cc06b666e814b5708c","observation_id":"256d99c0-72d3-49bd-a459-1d53dd50077b","resolution":{"observed_at":"2026-08-14T12:58:11.457963Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1902.10089","last_updated":"2019-11-05T18:27:37Z","snapshot_observed_at":"2026-08-14T17:10:51.116009Z","submitted_at":"2019-02-26T18:04:58Z","title":"Perturbed-History Exploration in Stochastic Multi-Armed Bandits","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1902.10089","snapshot_observed_at":"2026-08-14T12:58:10.499532Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.499532Z"},"links":{"cited_paper":"/paper/1902.10089","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:e9a42f1631f49afc43452aadf7b49121382d1280a671a113484f41e4479d39d4","observation_id":"b6fd56ab-8201-4c4e-b1e2-71c2c7125a4e","resolution":{"observed_at":"2026-08-14T12:58:10.499532Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1706.01383","last_updated":"2017-06-05T15:46:52Z","snapshot_observed_at":"2026-08-16T14:10:22.291672Z","submitted_at":"2017-06-05T15:46:52Z","title":"Sparse Stochastic Bandits","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1706.01383","snapshot_observed_at":"2026-08-14T12:58:10.503854Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.503854Z"},"links":{"cited_paper":"/paper/1706.01383","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:907fff4566c8f8d06dcdf1a93d063319f3226a6c8b41830765595f4232758b5a","observation_id":"e4c43df8-1e13-4e18-81f4-7806272d4427","resolution":{"observed_at":"2026-08-14T12:58:10.503854Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1507.07880","last_updated":"2016-02-24T17:10:17Z","snapshot_observed_at":"2026-08-16T18:56:30.643820Z","submitted_at":"2015-07-28T18:09:19Z","title":"Optimally Confident UCB: Improved Regret for Finite-Armed Bandits","version":3},"cited_work":{"arxiv_id":"1507.07880","doi":null,"metadata_source":"pith","pith_arxiv_id":"1507.07880","snapshot_observed_at":"2026-08-14T12:58:11.025102Z","title":"Optimally Confident UCB: Improved Regret for Finite-Armed Bandits","venue":"cs.LG","work_id":"0b5b1672-01df-4af9-a528-a4d669d5a9b3","year":2015},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.508386Z"},"links":{"cited_paper":"/paper/1507.07880","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:2a7d17ccae954d77556efc8837e2d6fdaa37993ab58b622c21dad88e7caca075","observation_id":"0b10026b-1927-491e-834e-99fc547aeb97","resolution":{"observed_at":"2026-08-14T12:58:11.029235Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1603.08661","last_updated":"2016-05-06T19:06:26Z","snapshot_observed_at":"2026-08-14T22:04:07.979635Z","submitted_at":"2016-03-29T07:12:14Z","title":"Regret Analysis of the Anytime Optimally Confident UCB Algorithm","version":2},"cited_work":{"arxiv_id":"1603.08661","doi":null,"metadata_source":"pith","pith_arxiv_id":"1603.08661","snapshot_observed_at":"2026-08-14T12:58:11.008574Z","title":"Regret Analysis of the Anytime Optimally Confident UCB Algorithm","venue":"cs.LG","work_id":"d48a310d-8e77-402d-8c5a-aebc5b1c8958","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.513353Z"},"links":{"cited_paper":"/paper/1603.08661","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:85aff346b5cf36f062496efc1931c5e0781fd63493c4c8225ff5c599f2b88930","observation_id":"461c5b44-5dd9-4a0a-bc2e-09c2cfa2a7fb","resolution":{"observed_at":"2026-08-14T12:58:11.013015Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"0704.3359","last_updated":"2007-04-25T12:36:55Z","snapshot_observed_at":"2026-08-15T10:42:06.912614Z","submitted_at":"2007-04-25T12:36:55Z","title":"Direct Optimization of Ranking Measures","version":1},"cited_work":{"arxiv_id":"0704.3359","doi":null,"metadata_source":"pith","pith_arxiv_id":"0704.3359","snapshot_observed_at":"2026-08-14T12:58:10.988349Z","title":"Direct Optimization of Ranking Measures","venue":"cs.IR","work_id":"41704fd3-3965-4049-a171-233031c12ab3","year":2007},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.518710Z"},"links":{"cited_paper":"/paper/0704.3359","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:8beaaff0981713758eeaa4cb795c79f05681b9d3a990199d23581f7a0e073700","observation_id":"4b04920d-a446-4908-b01c-9e8d997b0c59","resolution":{"observed_at":"2026-08-14T12:58:10.993900Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.442738Z","title":null,"venue":null,"work_id":"c322a6fb-4039-4cc6-861a-1ab28e056be7","year":2012},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.523192Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:fa29e57c65437a26f7a4fa0be79161da2a758f7d84e28216fb5bd1f04758e1b1","observation_id":"9e580260-ec0e-4405-bea6-73ddbaba1fa2","resolution":{"observed_at":"2026-08-14T12:58:11.446416Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.527135Z","title":null,"venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.527135Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:51c97d63941e6be641717fdf005f38d412e9eb2beb78b1260183bd617b2da706","observation_id":"75594cbd-4b4b-43b9-a502-987d1e311105","resolution":{"observed_at":"2026-08-14T12:58:10.527135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.425238Z","title":null,"venue":null,"work_id":"a6a7dfbf-248d-48a5-a43e-6b75f16cd636","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.530956Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:3a8f5657a85f763a5cf39e8e2c1dfb25daaac9f8a58ac39adbe23dbbbe8cb8ac","observation_id":"7a8edfc2-4c2b-47f9-8f30-cfe9384aded2","resolution":{"observed_at":"2026-08-14T12:58:11.428726Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.415341Z","title":null,"venue":null,"work_id":"f3f108a1-e77b-4cbe-8c9c-b08a07c81a48","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.535144Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:334bca3a2218f3f10910fa53b7e5551940c92e339339da8c50b632d22267bff7","observation_id":"44269619-4f85-426c-832f-2876df6c99e4","resolution":{"observed_at":"2026-08-14T12:58:11.418737Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1804.05929","last_updated":"2018-04-16T20:44:28Z","snapshot_observed_at":"2026-08-14T19:25:22.374742Z","submitted_at":"2018-04-16T20:44:28Z","title":"UCBoost: A Boosting Approach to Tame Complexity and Optimality for Stochastic Bandits","version":1},"cited_work":{"arxiv_id":"1804.05929","doi":null,"metadata_source":"pith","pith_arxiv_id":"1804.05929","snapshot_observed_at":"2026-08-14T12:58:10.967068Z","title":"UCBoost: A Boosting Approach to Tame Complexity and Optimality for Stochastic Bandits","venue":"cs.LG","work_id":"d5f28320-2b7f-4e7c-89b2-755ed47041a6","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.539292Z"},"links":{"cited_paper":"/paper/1804.05929","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:27a020ae21c6a7531a5a2236e839bcc9fd82d77f979d1314cc789ae269431cc9","observation_id":"b216a22f-065e-4b03-ae6c-5bcbc11bc3a6","resolution":{"observed_at":"2026-08-14T12:58:10.971843Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.405331Z","title":null,"venue":null,"work_id":"bea0fdf1-94d1-4f6e-b9a6-743124256ef8","year":2010},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.543637Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:75bcb120f3596cd1a66e0d5c6ff75635703f581036be49ab9c8f4466a5966de3","observation_id":"6a6a0ea8-3b71-4212-a9b5-c593bda1cf53","resolution":{"observed_at":"2026-08-14T12:58:11.408458Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.395189Z","title":null,"venue":null,"work_id":"40ee00cd-6ec1-4e87-bcee-03f75f31e3d5","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.547563Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:85491ad64db9536d4911da875d80a32351127a7d9662392224610e7822fb8402","observation_id":"4c9742a6-9c27-48ee-8eff-83f7292aabd7","resolution":{"observed_at":"2026-08-14T12:58:11.398465Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1808.08416","last_updated":"2021-04-04T22:32:27Z","snapshot_observed_at":"2026-08-16T12:13:53.481245Z","submitted_at":"2018-08-25T12:06:17Z","title":"Multiplayer bandits without observing collision information","version":2},"cited_work":{"arxiv_id":"1808.08416","doi":null,"metadata_source":"pith","pith_arxiv_id":"1808.08416","snapshot_observed_at":"2026-08-14T12:58:10.950651Z","title":"Multiplayer bandits without observing collision information","venue":"cs.LG","work_id":"d6d665d4-bbd3-46dd-a763-7cb067160134","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.551434Z"},"links":{"cited_paper":"/paper/1808.08416","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:8687ecafe9dc6bc0551f3a58e79bd59fca64332da28b7ac70c04883f80c372ca","observation_id":"4be96a48-523b-43d8-9306-563f668ececa","resolution":{"observed_at":"2026-08-14T12:58:10.955198Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.384110Z","title":null,"venue":null,"work_id":"aadf5422-7ec3-4291-9efb-221feaeca68f","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.556124Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:ee47e333896a6b9800c70ca575872596830ae4d5cd61239d379120ce0e7e10e8","observation_id":"bd9b4953-fd66-488b-b9ba-bc555e57ddb0","resolution":{"observed_at":"2026-08-14T12:58:11.387613Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1703.10254","last_updated":"2017-03-29T22:15:46Z","snapshot_observed_at":"2026-08-14T21:09:18.369041Z","submitted_at":"2017-03-29T22:15:46Z","title":"Bandit-Based Model Selection for Deformable Object Manipulation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1703.10254","snapshot_observed_at":"2026-08-14T12:58:10.561207Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.561207Z"},"links":{"cited_paper":"/paper/1703.10254","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:b0c4a4df64f0515cfec977d54e4fd4087331391b24ee0ea62fa86b5fe15014f4","observation_id":"755d5409-a726-46be-8b18-3c98da11bd36","resolution":{"observed_at":"2026-08-14T12:58:10.561207Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.372440Z","title":null,"venue":null,"work_id":"32c29cd5-ecd9-4cda-a702-9b3c69e23495","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.566333Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:b73fba7993d4fee9d0d7a90e78ad46d589aa76e1297239876dac5b6d919e1a40","observation_id":"1c3bd982-b3b5-487d-86ae-80b62c4be1f7","resolution":{"observed_at":"2026-08-14T12:58:11.376373Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.361477Z","title":null,"venue":null,"work_id":"fcfca416-3b67-4511-b705-3c1c3c9a10b2","year":null},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.571968Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:c5f08514b17c57c209950b61acb0f41e3ce51dd3ae6e2a7a3409c21cfd62f124","observation_id":"873c4a56-9320-436e-8f90-89fcad190b2a","resolution":{"observed_at":"2026-08-14T12:58:11.365279Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.580531Z","title":null,"venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.580531Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:1c223629385a4ac8c56ab65d30f212f60c46dbb433aa79c42ecc22d29d9abba9","observation_id":"95d846f9-0824-4e4b-9efe-0c1dbee75371","resolution":{"observed_at":"2026-08-14T12:58:10.580531Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.344068Z","title":null,"venue":null,"work_id":"456e9efc-1ac4-4777-a88a-5da7290f59be","year":1985},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.584817Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:78d406cbf45a23898af10b5d04e0a2918e2cd36297d388af69f8f9760d25a0be","observation_id":"ff809624-73ad-4df9-b0f1-a5f7fad2978d","resolution":{"observed_at":"2026-08-14T12:58:11.347417Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.332219Z","title":null,"venue":null,"work_id":"8d089322-97c2-42c1-8f82-f8cccc3118de","year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.589425Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:4a1d76376bab2f9fd7a979b3ba0d578146a947eb21a44ebdeb756ebe11fd1336","observation_id":"26799f95-8b40-48e5-8990-18ccac80b120","resolution":{"observed_at":"2026-08-14T12:58:11.336646Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.593394Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.593394Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:df7154db8a62d9999bf72dfe8ffd1b3733fc6465ebb30fb03bad156212a9ed40","observation_id":"632df4bb-7132-404e-88eb-95fcaa4b381c","resolution":{"observed_at":"2026-08-14T12:58:10.593394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1809.03186","last_updated":"2020-06-09T07:45:50Z","snapshot_observed_at":"2026-08-16T17:47:47.140098Z","submitted_at":"2018-09-10T08:52:55Z","title":"Off-line vs. On-line Evaluation of Recommender Systems in Small E-commerce","version":3},"cited_work":{"arxiv_id":"1809.03186","doi":null,"metadata_source":"pith","pith_arxiv_id":"1809.03186","snapshot_observed_at":"2026-08-14T12:58:10.905573Z","title":"Off-line vs. On-line Evaluation of Recommender Systems in Small E-commerce","venue":"cs.IR","work_id":"9da58f27-03b9-456d-95ba-c28ecd5bcf7b","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.597931Z"},"links":{"cited_paper":"/paper/1809.03186","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:5ddebc6776456b059481ee0ff9c731dc7c917ea0900509452db3c64a088981e1","observation_id":"da0db73e-4273-4a5d-9634-3adc46fe3e2c","resolution":{"observed_at":"2026-08-14T12:58:10.910406Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.09727","last_updated":"2017-07-31T05:46:39Z","snapshot_observed_at":"2026-08-14T20:44:19.342054Z","submitted_at":"2017-07-31T05:46:39Z","title":"Taming Non-stationary Bandits: A Bayesian Approach","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.09727","snapshot_observed_at":"2026-08-14T12:58:10.602490Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.602490Z"},"links":{"cited_paper":"/paper/1707.09727","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:953c423ffc2a340bdd93a06242a3e99baad83699f1a1085309c3ee24f2453c52","observation_id":"b8838379-4146-4767-992b-a3435347156a","resolution":{"observed_at":"2026-08-14T12:58:10.602490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.313175Z","title":null,"venue":null,"work_id":"e622a5c8-6c95-47dd-957d-c1f2a7778c4c","year":2010},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.607066Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:a9d223f784075015affbdef53eb29b110328ef579060a3e19835c556e2f4019d","observation_id":"b2593d22-4bf2-4bad-8512-477e44074595","resolution":{"observed_at":"2026-08-14T12:58:11.317034Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1905.01395","last_updated":"2019-05-04T00:27:22Z","snapshot_observed_at":"2026-08-14T16:37:28.644746Z","submitted_at":"2019-05-04T00:27:22Z","title":"On the Difficulty of Evaluating Baselines: A Study on Recommender Systems","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1905.01395","snapshot_observed_at":"2026-08-14T12:58:10.611039Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.611039Z"},"links":{"cited_paper":"/paper/1905.01395","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:104a0f5a7674aeca7527081e973670cf1a1967eaae28f5c1d82cd29d0ee6c321","observation_id":"2133227c-3fb2-46bf-b2db-3a4fdef4b991","resolution":{"observed_at":"2026-08-14T12:58:10.611039Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.301255Z","title":null,"venue":null,"work_id":"72a2edac-8903-474b-b63e-ab2da9947cf5","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.615363Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:508ca8a0cdc0f645c6c44e95e6ec6ee31ff875f1d2f3282ba7c2849cc53373df","observation_id":"f1f4b68c-3cae-4572-b91e-72ff3ba2e28e","resolution":{"observed_at":"2026-08-14T12:58:11.305181Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1305.6646","last_updated":"2013-05-28T22:12:59Z","snapshot_observed_at":"2026-08-15T00:16:18.002318Z","submitted_at":"2013-05-28T22:12:59Z","title":"Normalized Online Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1305.6646","snapshot_observed_at":"2026-08-14T12:58:10.619849Z","title":null,"venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.619849Z"},"links":{"cited_paper":"/paper/1305.6646","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:34dbb3e11506ef5256d914d837b713596ae085e197d75155724696ff85c3dc17","observation_id":"63bf1b85-3d7d-44d0-9d79-48e2a8c8e481","resolution":{"observed_at":"2026-08-14T12:58:10.619849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.288134Z","title":null,"venue":null,"work_id":"80aa080e-c779-4574-aa94-63c601bbcd3c","year":2015},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.626674Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:c3562703ca826798f8f9f86e0e321f6bf9c2c605a1603867f265e78e94389387","observation_id":"035af544-2c19-439d-8a20-af8d58cd84c6","resolution":{"observed_at":"2026-08-14T12:58:11.292175Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1702.06103","last_updated":"2017-05-09T14:29:38Z","snapshot_observed_at":"2026-08-14T21:15:39.271560Z","submitted_at":"2017-02-20T18:43:05Z","title":"An Improved Parametrization and Analysis of the EXP3++ Algorithm for Stochastic and Adversarial Bandits","version":2},"cited_work":{"arxiv_id":"1702.06103","doi":null,"metadata_source":"pith","pith_arxiv_id":"1702.06103","snapshot_observed_at":"2026-08-14T12:58:10.845192Z","title":"An Improved Parametrization and Analysis of the EXP3++ Algorithm for Stochastic and Adversarial Bandits","venue":"cs.LG","work_id":"696c073b-1adf-4c65-b650-ae1e2bb8e9eb","year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.632132Z"},"links":{"cited_paper":"/paper/1702.06103","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:8922ea92729951dca592440d67c771eee4c9ebf6cc74dc1d6f482713d6847281","observation_id":"4fa9a37c-a5ae-46eb-a595-294668ffe12d","resolution":{"observed_at":"2026-08-14T12:58:10.849690Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.274854Z","title":null,"venue":null,"work_id":"d0a7f5f6-3378-46ee-b5eb-e3120ed42f53","year":2012},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.636776Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:b1f4a091add1ecddaa82e9076ee7e75d3d31bc071490176033e27b59e33a4bf4","observation_id":"9361d072-87bc-42bc-b58d-8af3133c3a43","resolution":{"observed_at":"2026-08-14T12:58:11.278523Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.262908Z","title":null,"venue":null,"work_id":"bb7b9722-dd28-4ab1-b476-66bf7366d5e5","year":1972},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.642176Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:47738429578232dac505fa31640281af170decf6de9b00f2431d4cd6b71ad337","observation_id":"d0ba269a-745d-4453-bbe3-8719691446d1","resolution":{"observed_at":"2026-08-14T12:58:11.267073Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1702.04825","last_updated":"2017-02-17T21:39:29Z","snapshot_observed_at":"2026-08-14T21:16:16.674457Z","submitted_at":"2017-02-16T00:22:16Z","title":"Learning to Use Learners' Advice","version":2},"cited_work":{"arxiv_id":"1702.04825","doi":null,"metadata_source":"pith","pith_arxiv_id":"1702.04825","snapshot_observed_at":"2026-08-14T12:58:10.826225Z","title":"Learning to Use Learners' Advice","venue":"cs.LG","work_id":"3a5b03bc-9498-41f0-ae5e-fed38df6919e","year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.647170Z"},"links":{"cited_paper":"/paper/1702.04825","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:b3b3441d7dec470ec7734f1c0268c83b7e6e3a01da99fe707219abb8b4164213","observation_id":"a7ec3c55-7f96-4ed1-82c8-af44aa63f9f1","resolution":{"observed_at":"2026-08-14T12:58:10.831140Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1302.2318","last_updated":"2013-02-10T10:50:56Z","snapshot_observed_at":"2026-08-15T00:29:18.649938Z","submitted_at":"2013-02-10T10:50:56Z","title":"On Search Engine Evaluation Metrics","version":1},"cited_work":{"arxiv_id":"1302.2318","doi":null,"metadata_source":"pith","pith_arxiv_id":"1302.2318","snapshot_observed_at":"2026-08-14T12:58:10.807425Z","title":"On Search Engine Evaluation Metrics","venue":"cs.IR","work_id":"e666e6b2-8d81-4b18-a6a3-7edbfc650f61","year":2013},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.651913Z"},"links":{"cited_paper":"/paper/1302.2318","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:b44c989fd527829c563f1dedf7227ee02a2f1b34be4d8e9729a35a7ac01a8aac","observation_id":"084d940e-263a-4e03-b9d4-618adc8ac6a8","resolution":{"observed_at":"2026-08-14T12:58:10.813867Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.656919Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.656919Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:c947236dd89ecc918b851ea3711fcc173e865113ea1038ba9f148053efedfb0b","observation_id":"13776622-4b82-4323-9e6e-1fe5e191ddd7","resolution":{"observed_at":"2026-08-14T12:58:10.656919Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.244037Z","title":null,"venue":null,"work_id":"b7a853ae-f8ab-408a-9e60-d042935b8526","year":2014},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.660886Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:80acb94e0ff1fdb019090863b7500dd344fffa5c997af5ab8e3542c128b6b356","observation_id":"d5ca3b17-857e-4638-aed6-9678a85e9512","resolution":{"observed_at":"2026-08-14T12:58:11.247319Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.232847Z","title":null,"venue":null,"work_id":"c81d1c5b-54a8-49a7-ab46-c1c35442e566","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.665070Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:a321fc6ebba3ce195ebd26de7fdc05ab2ce917809f46ef889bc4fadd2b41fc65","observation_id":"fa599748-418f-4c35-99a2-7374d25050d8","resolution":{"observed_at":"2026-08-14T12:58:11.236271Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.219271Z","title":null,"venue":null,"work_id":"f1bbc4a7-c5f4-48f8-98b0-37b84dbf8080","year":2015},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.669325Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:a2a04267201b1afc3a32bdea8bafcfdffbc82889a0afff4982a7e97bb493d54f","observation_id":"c346406b-b24b-4048-bd60-c8991b875f47","resolution":{"observed_at":"2026-08-14T12:58:11.223109Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.207398Z","title":null,"venue":null,"work_id":"1afad2d1-3e5c-4f49-9b63-94de91c6f369","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.673865Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:a111c87ac8787303043f9744fbb3facb6f32652d0dc4ec7c90efa22435f9d7f0","observation_id":"5ae5a7e7-8360-4b46-9439-ac9d397dacb1","resolution":{"observed_at":"2026-08-14T12:58:11.211514Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1902.04864","last_updated":"2021-05-15T13:28:09Z","snapshot_observed_at":"2026-08-14T17:16:59.021623Z","submitted_at":"2019-02-13T11:12:47Z","title":"A Survey on Session-based Recommender Systems","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1902.04864","snapshot_observed_at":"2026-08-14T12:58:10.678430Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.678430Z"},"links":{"cited_paper":"/paper/1902.04864","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:261f8f3e0abb62b0c5ebebc3cfb41dbe1016346bd8d25c1d6bdc67cd15d81c25","observation_id":"16fc02b2-9cef-4934-8346-3392903146c4","resolution":{"observed_at":"2026-08-14T12:58:10.678430Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.191634Z","title":null,"venue":null,"work_id":"464530d5-5891-48d8-9fb0-39b8888aa2a3","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.683152Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:b7ed2a3821c0c80ab43fcd7e3a2b409d441a591a424a5f1e7ee03c09493c77de","observation_id":"4d7d267b-2db3-4bbb-b8e1-8de8ce4d83f3","resolution":{"observed_at":"2026-08-14T12:58:11.197173Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.177753Z","title":null,"venue":null,"work_id":"e5ca3d2b-d0ca-40e1-8cc8-7cda52f4bc38","year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.687697Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:d478fb4e2fe22366217e45207188136ac627e991960605ac42148ee147b2472c","observation_id":"294c78f2-be80-4642-a75d-000b40b10e18","resolution":{"observed_at":"2026-08-14T12:58:11.182182Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1811.00855","last_updated":"2019-01-24T08:12:19Z","snapshot_observed_at":"2026-08-16T18:58:18.474740Z","submitted_at":"2018-11-01T02:44:16Z","title":"Session-based Recommendation with Graph Neural Networks","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1811.00855","snapshot_observed_at":"2026-08-14T12:58:10.693416Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.693416Z"},"links":{"cited_paper":"/paper/1811.00855","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:71d6cb121b27035725a7ef9c76ecad51c013eaeb70f404610911a954b4647d04","observation_id":"0a67458f-6dcf-4933-8f42-e73ea8601d1e","resolution":{"observed_at":"2026-08-14T12:58:10.693416Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.165304Z","title":null,"venue":null,"work_id":"df040619-21fc-4cbd-93d9-b86e771fb0ef","year":2016},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.698388Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:79b3331a902d655f14ae364c87554b9d2767c5297388fe5ad382ea3a1764315f","observation_id":"71ba6c24-819e-45e5-bc0c-071b484f0c62","resolution":{"observed_at":"2026-08-14T12:58:11.169436Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1807.07623","last_updated":"2022-03-02T09:02:05Z","snapshot_observed_at":"2026-08-14T18:49:48.048465Z","submitted_at":"2018-07-19T19:42:19Z","title":"Tsallis-INF: An Optimal Algorithm for Stochastic and Adversarial Bandits","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1807.07623","snapshot_observed_at":"2026-08-14T12:58:10.702721Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.702721Z"},"links":{"cited_paper":"/paper/1807.07623","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:d11ca260255c10dffdb798bc3bd03e965a5933a4f53532ad429fd62c7dc00781","observation_id":"2f9ef6bb-9d95-4a9d-802b-43bfb219cd0a","resolution":{"observed_at":"2026-08-14T12:58:10.702721Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.337450Z","title":"InProceedings of the international workshop on reproducibility and replication in recommender systems evaluation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":2013,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.337450Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:1c8b8b84d77c27ab9d64dea9c06f6a81774a7ee03d3afe6b54d81e099d2cd116","observation_id":"703a22f1-62af-4d92-b651-479145bc1485","resolution":{"observed_at":"2026-08-14T12:58:10.337450Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:10.399503Z","title":"In Proceedings of the 1st workshop on deep learning for recommender systems","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":2016,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.399503Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:41d007c6e7f37431e5d1d5a950f431c104a9a89c579b53f96a461dee4b7bcd45","observation_id":"51afb35d-1e0c-4501-aff2-ae610a2023f3","resolution":{"observed_at":"2026-08-14T12:58:10.399503Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T12:58:11.505351Z","title":"In Proceedings of the 23rd ACM SIGKDD International Conference on Knowledge Discovery and Data Mining","venue":null,"work_id":"a129f6bf-2053-43ac-a800-f8abdbfcc504","year":null},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.468980Z"},"links":{"citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:e6a149a8272388c6303db7ea698c36dc937c88fd34ff6e40b3243e893198703c","observation_id":"7d482a32-f08e-497f-ac04-440470b96726","resolution":{"observed_at":"2026-08-14T12:58:11.509550Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1906.11336","last_updated":"2019-08-06T14:50:15Z","snapshot_observed_at":"2026-08-14T16:11:10.001888Z","submitted_at":"2019-06-26T20:27:38Z","title":"A Simple Deep Personalized Recommendation System","version":2},"cited_work":{"arxiv_id":"1906.11336","doi":null,"metadata_source":"pith","pith_arxiv_id":"1906.11336","snapshot_observed_at":"2026-08-14T12:58:10.922726Z","title":"A Simple Deep Personalized Recommendation System","venue":"cs.IR","work_id":"a598c975-f3a5-43ad-919c-9eca4c3fc8c3","year":2019},"citing_paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit","version":1},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-14T12:58:10.576164Z"},"links":{"cited_paper":"/paper/1906.11336","citing_paper":"/paper/1908.06158"},"observation_digest":"sha256:a94c55b56a355c82e4c22ff537bd8c0caae6399f2d86576927a88940a76f6dfc","observation_id":"e1348669-2278-4e15-a8a7-451e00acb882","resolution":{"observed_at":"2026-08-14T12:58:10.926947Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"1908.06158","last_updated":"2019-08-16T20:44:01Z","latest_version":1,"primary_category":"cs.IR","snapshot_observed_at":"2026-08-16T11:26:00.431548Z","submitted_at":"2019-08-16T20:44:01Z","title":"Accelerated learning from recommender systems using multi-armed bandit"},"reference_resolution":{"displayed":90,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":74,"verified_exact":13,"verified_fuzzy":2},"total_outbound_references":90},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 90 of 90 outbound references and 1 inbound Pith citation observation for arXiv:1908.06158."}