{"as_of":"2026-08-16T21:38:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:550719ef298a82d5a67a88aa4cf2706c0824cb630bf97a0e6b891ff17686c8cf","coverage":[{"denominator":25,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":25,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T20:34:16.067949Z","state":"measured"},{"denominator":25,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":25,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2501.08109/citation-record","integrity":"/paper/2501.08109/integrity","json":"/paper/2501.08109/citation-record.json","paper":"/paper/2501.08109"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.468206Z","title":"Optimal ordering, issuance and disposal policies for inventory management of perishable products,","venue":null,"work_id":"f9478224-a8e2-4a88-9723-e78d5fd48da1","year":2014},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.933992Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:a76956a6f43c5a7a4dedffd6cab16d3e610dd9279daf4d3311fdb0466a4367a7","observation_id":"7995875f-fbee-41bd-8a25-410644076b51","resolution":{"observed_at":"2026-08-10T20:34:16.473122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.449910Z","title":"Addressing the cold-start problem of recommendation systems for financial products by using few-shot deep learning,","venue":null,"work_id":"6136a2e6-fb98-4ac1-9083-8353aa6a0b4c","year":2022},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.938748Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:522b1dca5deef6040374b8fca33b7f0fbbb7796c86e33c79ebb8ea8ba10c965c","observation_id":"1557d025-0a68-439a-8121-244757c70394","resolution":{"observed_at":"2026-08-10T20:34:16.456898Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.431480Z","title":"Increasing supply chain robustness through process flexibility and inventory,","venue":null,"work_id":"1460ed00-3c18-4793-b38d-db8ffefd4d18","year":2018},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.942863Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:0115be93bc4c726a847543ac71eb25d842575341649f0d4ed41d1f62295d2c32","observation_id":"04c92d28-b3b2-4b6c-97a8-4ca1456467e8","resolution":{"observed_at":"2026-08-10T20:34:16.437735Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.415959Z","title":null,"venue":null,"work_id":"bf46fee5-905f-4161-9cd8-49c94f97480a","year":1996},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.947491Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:1fb0cc0ca3f1d2a0e4ef35a1625698bb336103206746095f4b8255a92e25d8ad","observation_id":"f5168519-5d7a-4ef1-9072-305afb6130f4","resolution":{"observed_at":"2026-08-10T20:34:16.420303Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.400806Z","title":"Inventory management in supply chains: a reinforcement learning approach,","venue":null,"work_id":"438f1069-9ae3-4258-80dd-93551786e167","year":2002},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.955901Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:ff7b54f4aa20082527385b1a0d99bfa2eb1e82a8e6707a2f3e42ad5cdf08efb4","observation_id":"bdb0e83c-f9b1-4f69-bfaa-6810ac09b9ff","resolution":{"observed_at":"2026-08-10T20:34:16.405301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.385022Z","title":"Global supply chain management: a reinforcement learning approach,","venue":null,"work_id":"06a3caa2-2378-46f7-9313-b3e62e30d7c7","year":2002},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.962124Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:5cc81b3b58cff45d4e5eb16d5c1fa7864a80766a49803db63dddb9bdbc304f77","observation_id":"383f1c1b-1bbb-416c-8115-f00e8abe0786","resolution":{"observed_at":"2026-08-10T20:34:16.390881Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.366682Z","title":"Inventory management of new products in retailers using model-based deep reinforcement learning,","venue":null,"work_id":"93573dc0-1760-41e9-91ea-7cbbf6f86bed","year":2023},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.967166Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:233a92c2aab93ec7e4663ec88625716170e8004269ef2ffa16e1fa3ce676234d","observation_id":"6ad26bd8-fd7b-4216-9e0b-442ad3e3fb1d","resolution":{"observed_at":"2026-08-10T20:34:16.370925Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.352686Z","title":"Deep reinforcement learning for inventory control: A roadmap,","venue":null,"work_id":"fa50f1a4-47e5-4ef9-88a0-f021689153de","year":2022},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.971347Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:454f40b05a6ab11a35c5bf9ff50f1b9731be312b94af3e350dfb0fa613a7f608","observation_id":"4706c330-27b1-41b8-b8a5-7d15c9769ae9","resolution":{"observed_at":"2026-08-10T20:34:16.357700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2012.02476","last_updated":"2020-12-04T08:58:35Z","snapshot_observed_at":"2026-08-16T19:01:02.203383Z","submitted_at":"2020-12-04T08:58:35Z","title":"Offline Meta-level Model-based Reinforcement Learning Approach for Cold-Start Recommendation","version":1},"cited_work":{"arxiv_id":"2012.02476","doi":null,"metadata_source":"pith","pith_arxiv_id":"2012.02476","snapshot_observed_at":"2026-08-10T20:34:16.134637Z","title":"Offline Meta-level Model-based Reinforcement Learning Approach for Cold-Start Recommendation","venue":"cs.LG","work_id":"58e62b29-e636-4d5b-a2ee-6dcc161b2e31","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.975250Z"},"links":{"cited_paper":"/paper/2012.02476","citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:db6073ad5de43f31e8302525a72349a42a59ba445f8f7ff8cd1b049923eb2daf","observation_id":"4de4fad2-1e5c-4afb-9132-faa82b15bf24","resolution":{"observed_at":"2026-08-10T20:34:16.141926Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.335721Z","title":"Dyna, an integrated architecture for learning, planning, and reacting,","venue":null,"work_id":"014d0024-1250-4fdb-9f4e-8410573493c0","year":1991},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.980823Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:56e79b4fab7292939c60b5ca8b1fadd61e5550e73a7a98cbee20c2318de501b2","observation_id":"954854cc-bb9e-4b4f-be2b-8931bfcbe337","resolution":{"observed_at":"2026-08-10T20:34:16.342970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.322823Z","title":"An improved dyna-q algorithm for mobile robot path planning in unknown dynamic environment,","venue":null,"work_id":"47666b98-980e-4055-a15d-cc297e0cbdf0","year":2021},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.985038Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:7eb6b63013affb2cfdd9cea12ef4e6eea0a9113b76d7216df49abba4164fd45c","observation_id":"5e426abb-4369-4a04-9b28-eb0dea26cb16","resolution":{"observed_at":"2026-08-10T20:34:16.327658Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.310125Z","title":"Pseudo dyna-q: A reinforcement learning framework for interactive recommendation,","venue":null,"work_id":"10eadf9d-9928-478a-b3d2-5250de76b5c8","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.989023Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:d13fb7b468497e56ebcb8349b227ae625a19124afb2dd748363293fa8fe6a4e8","observation_id":"4672f25e-498e-431f-a7a3-8c61e11ba477","resolution":{"observed_at":"2026-08-10T20:34:16.314469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:15.994388Z","title":"A survey of transfer learning,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.994388Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:ec3f0bacb5e4a7945530894795e1b6d3bfac3a2a67a640feb5f60e1524e9eb6f","observation_id":"84b4c097-d190-4430-9681-c4183d8f8c56","resolution":{"observed_at":"2026-08-10T20:34:15.994388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.289085Z","title":"Imitation and transfer q-learning-based parameter identification for composite load modeling,","venue":null,"work_id":"bd4d55f5-4275-4bee-8642-fd0a8d7d3afc","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:15.998736Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:4824e710bb84e07686ed90003d5691f22441b36418a3bf03a131bbb302b5b1c3","observation_id":"3615bd0d-7a93-4902-9c5f-31a6cec073e0","resolution":{"observed_at":"2026-08-10T20:34:16.293454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.275093Z","title":"Target transfer q-learning and its convergence analysis,","venue":null,"work_id":"ff8296c2-3b45-4276-aed3-ec2bd53945b1","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.003823Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:79ddba56f0d4bc20706675f8ed834dab27fa0b6364a81f8e7a2594af0abf498c","observation_id":"3506d899-3ce6-41be-b58f-8d1b3c2703fb","resolution":{"observed_at":"2026-08-10T20:34:16.280093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.262064Z","title":"Transferring models in hybrid reinforcement learning agents,","venue":null,"work_id":"2cf72bf6-64f9-4662-8306-f74602a89fa9","year":2011},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.008242Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:89e865d9a35131248bd93f2ba15228d4fc063697ed51681cbdac9a43738f9001","observation_id":"278fbc54-bd8a-4a56-b796-1881057871b0","resolution":{"observed_at":"2026-08-10T20:34:16.266682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.244480Z","title":"q-transfer: A novel framework for efficient deep transfer learning in networking,","venue":null,"work_id":"2056762d-c4fd-48b6-93b1-30ff85921ae0","year":2020},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.019855Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:dc8de307282c259526f98e315d490a2ae64a37e0e243b6308874eba200e62ec5","observation_id":"e597d587-1d9e-433f-a466-78bec833af87","resolution":{"observed_at":"2026-08-10T20:34:16.249961Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.230524Z","title":"Learning rate schedules for faster stochastic gradient search,","venue":null,"work_id":"b9c875b1-4f4f-40aa-915d-fc0edc4a25f1","year":1992},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.025478Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:62924b4b2f58ebce5741bbfaf3b8bb44bdedd78ccf3dc0e20a6917aae97e382b","observation_id":"c5b55b61-cd2a-4d4e-915a-52971a4947e7","resolution":{"observed_at":"2026-08-10T20:34:16.235003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.216491Z","title":"Survey of model-based reinforcement learning: Applications on robotics,","venue":null,"work_id":"eee55622-9ed3-4aaa-8a97-02c31b80ae4d","year":2017},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.030976Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:d67fc02b776a0b36d374b4a70ecb1b614e8d80df0c24f9c5c24195fc48e468d6","observation_id":"89015f1a-1864-4eb2-8dee-7d0146b5d4a2","resolution":{"observed_at":"2026-08-10T20:34:16.221788Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1811.01848","last_updated":"2019-01-28T16:39:23Z","snapshot_observed_at":"2026-08-14T18:04:02.877945Z","submitted_at":"2018-11-05T17:09:18Z","title":"Plan Online, Learn Offline: Efficient Learning and Exploration via Model-Based Control","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1811.01848","snapshot_observed_at":"2026-08-10T20:34:16.035861Z","title":"Plan online, learn offline: Efficient learning and exploration via model-based control,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.035861Z"},"links":{"cited_paper":"/paper/1811.01848","citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:6f8bdff9a797f52c2bdadc8d9cc46dad93aa42c524482caf0a2a91f785e87098","observation_id":"c53846d5-91f1-48d5-8974-f10e77af222d","resolution":{"observed_at":"2026-08-10T20:34:16.035861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1807.03858","last_updated":"2021-02-15T17:29:47Z","snapshot_observed_at":"2026-08-14T18:53:32.384767Z","submitted_at":"2018-07-10T20:53:04Z","title":"Algorithmic Framework for Model-based Deep Reinforcement Learning with Theoretical Guarantees","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1807.03858","snapshot_observed_at":"2026-08-10T20:34:16.041618Z","title":"Algorithmic framework for model-based deep reinforcement learning with theoretical guarantees,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.041618Z"},"links":{"cited_paper":"/paper/1807.03858","citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:ec4dc140ec03f8740bba55941ecd527cea0d9b63042e7ccdfcc41d78498e10ef","observation_id":"c7381fb9-0998-4cf9-9e7b-b5bc5285c746","resolution":{"observed_at":"2026-08-10T20:34:16.041618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.198594Z","title":"Hands-on bayesian neural networks–a tutorial for deep learning users,","venue":null,"work_id":"acd4eba4-b214-4e29-b1b9-650b3e8d7724","year":2022},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.047509Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:537265376f42039cc65f03613bb846b1e0476fdd8e6a39156c1fc42c737e2de3","observation_id":"7a162c04-52f1-4a4a-9065-842dd722a8ad","resolution":{"observed_at":"2026-08-10T20:34:16.205643Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.181826Z","title":"French bakery daily sales,","venue":null,"work_id":"88c21e21-85e1-41d3-9e85-bff80993dfd3","year":null},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.052924Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:93c83e825d1779177cf36551a6895b31e31565e1e9032d191f54a52f0e4de83a","observation_id":"9882ca4f-61ae-4b5f-b363-5145f34732c2","resolution":{"observed_at":"2026-08-10T20:34:16.187279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.067949Z","title":"Multilayer perceptron and neural networks,","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.067949Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:d554695e87fcb7e5ff3b3205175e4cb80723f62b4d1001a7e8714c2d8680a220","observation_id":"079c737a-0a18-4aba-9438-c3c283954055","resolution":{"observed_at":"2026-08-10T20:34:16.067949Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T20:34:16.163218Z","title":"Available: https://www.kaggle.com/datasets/ matthieugimbert/french-bakery-daily-sales","venue":null,"work_id":"64b01905-370e-4e46-b158-9f495ff97625","year":null},"citing_paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning","version":4},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:16.063298Z"},"links":{"citing_paper":"/paper/2501.08109"},"observation_digest":"sha256:aacf28be521e60be16c962ae6a208fa9aa97c04ec5047c67aa86ac60198a0551","observation_id":"6813897c-cfd9-4e62-8831-88fed5640d60","resolution":{"observed_at":"2026-08-10T20:34:16.167858Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.08109","last_updated":"2025-06-09T11:45:53Z","latest_version":4,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-16T05:52:00.405109Z","submitted_at":"2025-01-14T13:40:08Z","title":"Data-driven inventory management for new products: An adjusted Dyna-$Q$ approach with transfer learning"},"reference_resolution":{"displayed":25,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":5,"verified_exact":1,"verified_fuzzy":19},"total_outbound_references":25},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 25 of 25 outbound references and 0 inbound Pith citation observations for arXiv:2501.08109."}