{"as_of":"2026-08-18T18:49:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f18ae07767b55f403ff462c23d85cafc5a66fc01664aacbeab52cacaa8180774","coverage":[{"denominator":25,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":25,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T20:34:16.067949Z","state":"measured"},{"denominator":25,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":25,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2501.08109/citation-record","integrity":"/paper/2501.08109/integrity","json":"/paper/2501.08109/citation-record.json","paper":"/paper/2501.08109"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.468206Z","title":"Optimal ordering, issuance and disposal policies for inventory management of perishable products,","venue":null,"work_id":"f9478224-a8e2-4a88-9723-e78d5fd48da1","year":2014},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.933992Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:5cb399a994565ba63129f1cef7e627ba469154b455cc220a55fe08906f4a00f3","observation_id":"7995875f-fbee-41bd-8a25-410644076b51","resolution":{"observed_at":"2026-08-10T20:34:16.473122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.449910Z","title":"Addressing the cold-start problem of recommendation systems for financial products by using few-shot deep learning,","venue":null,"work_id":"6136a2e6-fb98-4ac1-9083-8353aa6a0b4c","year":2022},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.938748Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:9cf720735074370ba6c841e6488ae32e56e625dac45552e6aae2decbfa693e81","observation_id":"1557d025-0a68-439a-8121-244757c70394","resolution":{"observed_at":"2026-08-10T20:34:16.456898Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.431480Z","title":"Increasing supply chain robustness through process flexibility and inventory,","venue":null,"work_id":"1460ed00-3c18-4793-b38d-db8ffefd4d18","year":2018},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.942863Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:55e2b1b5efd8459b28d5b1f4589342863b35dbd5009c3e1e55aa4a5a3fe40024","observation_id":"04c92d28-b3b2-4b6c-97a8-4ca1456467e8","resolution":{"observed_at":"2026-08-10T20:34:16.437735Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.415959Z","title":null,"venue":null,"work_id":"bf46fee5-905f-4161-9cd8-49c94f97480a","year":1996},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.947491Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:c2f55e3793141c57b3d091f6b1527e4aa721be7bd12a4a3eb57e951859ba7eef","observation_id":"f5168519-5d7a-4ef1-9072-305afb6130f4","resolution":{"observed_at":"2026-08-10T20:34:16.420303Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.400806Z","title":"Inventory management in supply chains: a reinforcement learning approach,","venue":null,"work_id":"438f1069-9ae3-4258-80dd-93551786e167","year":2002},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.955901Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:32fc72c37abf956d173898cfe01b56deb4c5094b1938a8cc300e54e09df085b8","observation_id":"bdb0e83c-f9b1-4f69-bfaa-6810ac09b9ff","resolution":{"observed_at":"2026-08-10T20:34:16.405301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.385022Z","title":"Global supply chain management: a reinforcement learning approach,","venue":null,"work_id":"06a3caa2-2378-46f7-9313-b3e62e30d7c7","year":2002},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.962124Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:de4b410e8ca5856cee9cfa43eb4fc40ef8bbe59492b88b09730bc2fec0d9f15a","observation_id":"383f1c1b-1bbb-416c-8115-f00e8abe0786","resolution":{"observed_at":"2026-08-10T20:34:16.390881Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.366682Z","title":"Inventory management of new products in retailers using model-based deep reinforcement learning,","venue":null,"work_id":"93573dc0-1760-41e9-91ea-7cbbf6f86bed","year":2023},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.967166Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:7a8c2920ab99c24e6f8a9a826dc6074f0a0430e6a98deacee03ea659b8068771","observation_id":"6ad26bd8-fd7b-4216-9e0b-442ad3e3fb1d","resolution":{"observed_at":"2026-08-10T20:34:16.370925Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.352686Z","title":"Deep reinforcement learning for inventory control: A roadmap,","venue":null,"work_id":"fa50f1a4-47e5-4ef9-88a0-f021689153de","year":2022},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.971347Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:d5082ae869a6c608d27bf57911f8c6e8dff8ff16fa69de71e8ffe7fb81576c61","observation_id":"4706c330-27b1-41b8-b8a5-7d15c9769ae9","resolution":{"observed_at":"2026-08-10T20:34:16.357700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2012.02476","last_updated":"2020-12-04T08:58:35Z","snapshot_observed_at":"2026-08-16T19:01:02.203383Z","submitted_at":"2020-12-04T08:58:35Z","title":"Offline Meta-level Model-based Reinforcement Learning Approach for Cold-Start Recommendation","version":1},"cited_work":{"arxiv_id":"2012.02476","doi":null,"metadata_source":"pith","pith_arxiv_id":"2012.02476","snapshot_observed_at":"2026-08-10T20:34:16.134637Z","title":"Offline Meta-level Model-based Reinforcement Learning Approach for Cold-Start Recommendation","venue":"cs.LG","work_id":"58e62b29-e636-4d5b-a2ee-6dcc161b2e31","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.975250Z"},"links":{"cited_paper":"/paper/2012.02476","citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:c39ad80b83d72af63a11b44498d0e14fdc6d6a9a70a3a79ac85301c343f389fc","observation_id":"4de4fad2-1e5c-4afb-9132-faa82b15bf24","resolution":{"observed_at":"2026-08-10T20:34:16.141926Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.335721Z","title":"Dyna, an integrated architecture for learning, planning, and reacting,","venue":null,"work_id":"014d0024-1250-4fdb-9f4e-8410573493c0","year":1991},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.980823Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:55351eda65adb2cf90fffe8fc6822d655497b6a35dc215586da6c511bf47beda","observation_id":"954854cc-bb9e-4b4f-be2b-8931bfcbe337","resolution":{"observed_at":"2026-08-10T20:34:16.342970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.322823Z","title":"An improved dyna-q algorithm for mobile robot path planning in unknown dynamic environment,","venue":null,"work_id":"47666b98-980e-4055-a15d-cc297e0cbdf0","year":2021},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.985038Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:ff9e8cb6881d241740287ac7af6b1813245582dce18001b1739b7b8fbf18e927","observation_id":"5e426abb-4369-4a04-9b28-eb0dea26cb16","resolution":{"observed_at":"2026-08-10T20:34:16.327658Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.310125Z","title":"Pseudo dyna-q: A reinforcement learning framework for interactive recommendation,","venue":null,"work_id":"10eadf9d-9928-478a-b3d2-5250de76b5c8","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.989023Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:56891cb00a9eacdac61bbfa39aee8a46d051560fc2c3dcc58bd59ec275a64d78","observation_id":"4672f25e-498e-431f-a7a3-8c61e11ba477","resolution":{"observed_at":"2026-08-10T20:34:16.314469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:15.994388Z","title":"A survey of transfer learning,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.994388Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:ec3f0bacb5e4a7945530894795e1b6d3bfac3a2a67a640feb5f60e1524e9eb6f","observation_id":"84b4c097-d190-4430-9681-c4183d8f8c56","resolution":{"observed_at":"2026-08-10T20:34:15.994388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.289085Z","title":"Imitation and transfer q-learning-based parameter identification for composite load modeling,","venue":null,"work_id":"bd4d55f5-4275-4bee-8642-fd0a8d7d3afc","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.998736Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:c71471d1046f37e72952fae06b2a447402e057b406e6e20f9e1fc5caadad7b2b","observation_id":"3615bd0d-7a93-4902-9c5f-31a6cec073e0","resolution":{"observed_at":"2026-08-10T20:34:16.293454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.275093Z","title":"Target transfer q-learning and its convergence analysis,","venue":null,"work_id":"ff8296c2-3b45-4276-aed3-ec2bd53945b1","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.003823Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:5949a5277ee6b17c9fe17c7d9ad8a187e859c049adc910e049d72afa1a0f07f8","observation_id":"3506d899-3ce6-41be-b58f-8d1b3c2703fb","resolution":{"observed_at":"2026-08-10T20:34:16.280093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.262064Z","title":"Transferring models in hybrid reinforcement learning agents,","venue":null,"work_id":"2cf72bf6-64f9-4662-8306-f74602a89fa9","year":2011},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.008242Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:6c15dd22ad8cdbcc92de7b887b7f4fd558c2da6293d21d60746101ce14a44ad6","observation_id":"278fbc54-bd8a-4a56-b796-1881057871b0","resolution":{"observed_at":"2026-08-10T20:34:16.266682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.244480Z","title":"q-transfer: A novel framework for efficient deep transfer learning in networking,","venue":null,"work_id":"2056762d-c4fd-48b6-93b1-30ff85921ae0","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.019855Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:062f024058603b0f6a266b6e4efad5349945f36bbd62a159352d7be27649e952","observation_id":"e597d587-1d9e-433f-a466-78bec833af87","resolution":{"observed_at":"2026-08-10T20:34:16.249961Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.230524Z","title":"Learning rate schedules for faster stochastic gradient search,","venue":null,"work_id":"b9c875b1-4f4f-40aa-915d-fc0edc4a25f1","year":1992},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.025478Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:53b3052674495695f16e4e366f72fb5049a0d68fca38ca52df3d5771629b4a33","observation_id":"c5b55b61-cd2a-4d4e-915a-52971a4947e7","resolution":{"observed_at":"2026-08-10T20:34:16.235003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.216491Z","title":"Survey of model-based reinforcement learning: Applications on robotics,","venue":null,"work_id":"eee55622-9ed3-4aaa-8a97-02c31b80ae4d","year":2017},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.030976Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:a5d24947bda0cf28d499f8b2f48636d2ff038e02bc4de15136d194c8b86b1ed3","observation_id":"89015f1a-1864-4eb2-8dee-7d0146b5d4a2","resolution":{"observed_at":"2026-08-10T20:34:16.221788Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1811.01848","last_updated":"2019-01-28T16:39:23Z","snapshot_observed_at":"2026-08-14T18:04:02.877945Z","submitted_at":"2018-11-05T17:09:18Z","title":"Plan Online, Learn Offline: Efficient Learning and Exploration via Model-Based Control","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1811.01848","snapshot_observed_at":"2026-08-10T20:34:16.035861Z","title":"Plan online, learn offline: Efficient learning and exploration via model-based control,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.035861Z"},"links":{"cited_paper":"/paper/1811.01848","citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:6f8bdff9a797f52c2bdadc8d9cc46dad93aa42c524482caf0a2a91f785e87098","observation_id":"c53846d5-91f1-48d5-8974-f10e77af222d","resolution":{"observed_at":"2026-08-10T20:34:16.035861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1807.03858","last_updated":"2021-02-15T17:29:47Z","snapshot_observed_at":"2026-08-14T18:53:32.384767Z","submitted_at":"2018-07-10T20:53:04Z","title":"Algorithmic Framework for Model-based Deep Reinforcement Learning with Theoretical Guarantees","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1807.03858","snapshot_observed_at":"2026-08-10T20:34:16.041618Z","title":"Algorithmic framework for model-based deep reinforcement learning with theoretical guarantees,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.041618Z"},"links":{"cited_paper":"/paper/1807.03858","citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:ec4dc140ec03f8740bba55941ecd527cea0d9b63042e7ccdfcc41d78498e10ef","observation_id":"c7381fb9-0998-4cf9-9e7b-b5bc5285c746","resolution":{"observed_at":"2026-08-10T20:34:16.041618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.198594Z","title":"Hands-on bayesian neural networks–a tutorial for deep learning users,","venue":null,"work_id":"acd4eba4-b214-4e29-b1b9-650b3e8d7724","year":2022},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.047509Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:09338c931bcae7c74af7ef2355f5797c80e0559bd9c3bc0dedb40ffd9558adb9","observation_id":"7a162c04-52f1-4a4a-9065-842dd722a8ad","resolution":{"observed_at":"2026-08-10T20:34:16.205643Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.181826Z","title":"French bakery daily sales,","venue":null,"work_id":"88c21e21-85e1-41d3-9e85-bff80993dfd3","year":null},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.052924Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:8149bfc09cce962dd7c66b3b4c2c8837d85c2d65ef613bf2f38d4ebbe4f433f5","observation_id":"9882ca4f-61ae-4b5f-b363-5145f34732c2","resolution":{"observed_at":"2026-08-10T20:34:16.187279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.067949Z","title":"Multilayer perceptron and neural networks,","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.067949Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:d554695e87fcb7e5ff3b3205175e4cb80723f62b4d1001a7e8714c2d8680a220","observation_id":"079c737a-0a18-4aba-9438-c3c283954055","resolution":{"observed_at":"2026-08-10T20:34:16.067949Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.163218Z","title":"Available: https://www.kaggle.com/datasets/ matthieugimbert/french-bakery-daily-sales","venue":null,"work_id":"64b01905-370e-4e46-b158-9f495ff97625","year":null},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.063298Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:93d507577f9f313b12aa10ddfc8b9bec87db39239cb1a12fee50be1bcc8f80d1","observation_id":"6813897c-cfd9-4e62-8831-88fed5640d60","resolution":{"observed_at":"2026-08-10T20:34:16.167858Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","latest_version":4,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning"},"reference_resolution":{"displayed":25,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":5,"verified_exact":1,"verified_fuzzy":19},"total_outbound_references":25},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 25 of 25 outbound references and 0 inbound Pith citation observations for arXiv:2501.08109."}