{"as_of":"2026-08-08T01:11:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:afbb0d78ae59dd35c646a869ccfce5662701de72ceb5045947af9cd78902ee5a","coverage":[{"denominator":32,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":32,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T06:00:33.591579Z","state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.06484/citation-record","integrity":"/paper/2506.06484/integrity","json":"/paper/2506.06484/citation-record.json","paper":"/paper/2506.06484"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:32.268126Z","title":"Enhancing Battery Storage Energy Arbitrage With Deep Reinforcement Learning and Time-Series Forecasting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.268126Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:0071465d2ca09021c5b022f486e0147b54bca4db211567842f8ba130cc2d05c2","observation_id":"0e624b34-e94d-4e70-ae75-77166058e838","resolution":{"observed_at":"2026-08-07T06:00:32.268126Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2025.11542","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:37.104487Z","title":"Deep reinforce- ment learning for economic battery dispatch: A compre- hensive comparison of algorithms and experiment design choices","venue":null,"work_id":"3d648496-2566-4417-aae2-593863568787","year":2025},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.277089Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:99562873392a4a8c14465083c2711af289e098b5c693fe0a4e10749173cfc8e7","observation_id":"4c64a3ef-0963-4d7f-9a4e-754d752f4324","resolution":{"observed_at":"2026-08-07T06:00:37.192796Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2020.30478","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:36.882205Z","title":"Deep-Reinforcement- Learning-Based Capacity Scheduling for PV-Battery Stor- age System","venue":null,"work_id":"0e25dc99-0afa-44ee-803d-189128835fd2","year":2021},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.286643Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:5244188a788ea19b04cf30c189648875439f32f0c49b7899051f4ee45c794c36","observation_id":"9dc97af9-2c9e-41ef-9307-20cf0b99fb10","resolution":{"observed_at":"2026-08-07T06:00:36.957875Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:32.297626Z","title":"Deep Reinforcement Learning-Based Energy Storage Arbitrage With Accurate Lithium-Ion Battery Degradation Model","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.297626Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:d6a21f853a869d52362940508d4db149de08fc9f9d6422cb61a971a699562ce1","observation_id":"547d1e26-fcb3-4f37-b4b3-6719e9176733","resolution":{"observed_at":"2026-08-07T06:00:32.297626Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2021.12195","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:36.551832Z","title":"Data- driven battery operation for energy arbitrage using rainbow deep reinforcement learning","venue":null,"work_id":"647c9549-9fa0-4abc-88d3-24269431a261","year":2022},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.308976Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:fae6b5ae6118b7960fdc9c006c207a5022d4121a16feb117b79da6ea88c55bc4","observation_id":"d437ff1b-1ade-4a64-b997-b64595917182","resolution":{"observed_at":"2026-08-07T06:00:36.665239Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:32.318404Z","title":"Dynamicen- ergy conversion and management strategy for an integrated electricity and natural gas system with renewable energy: Deep reinforcement learning approach","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.318404Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:9706135414bf26e0999d3a7196b27e9f1fdc377032eaa2be4702cd7924f125bf","observation_id":"501e2b0d-9eae-4446-88a8-d70d07fd04ba","resolution":{"observed_at":"2026-08-07T06:00:32.318404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:38.870502Z","title":"Inte- grated Electricity-Gas System Optimal Dispatch Based on Deep Reinforcement Learning","venue":null,"work_id":"2bbad681-0aa0-4599-a328-98059a1cda41","year":2021},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.327706Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:6def56974b4a846c3f8f2edc73759e17f93f6b127e7e5eb19d9ed37bf56fbacb","observation_id":"be40b532-3eb5-4de4-97cd-97ee57ba881e","resolution":{"observed_at":"2026-08-07T06:00:38.949458Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:32.348341Z","title":"Dynamicoptimizationofanintegrateden- ergysystemwithcarboncaptureandpower-to-gasintercon- nection: A deep reinforcement learning-based scheduling strategy","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.348341Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:ba6b873808cd6b33d2ad38f0edc6a0742507f2b5b93c4f5601a04d55652dfde6","observation_id":"da58fa93-0b88-4442-a3dd-1e19713243e6","resolution":{"observed_at":"2026-08-07T06:00:32.348341Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1007/s40565-016-0238-z","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:34.770220Z","title":"Stochastic coordinated operation of wind and battery energy storage system considering bat- tery degradation","venue":null,"work_id":"cadeaaa7-61f8-4361-8d2c-b6e1928f98db","year":2016},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.359983Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:44a61f66b1e1aea20bb64927e278ec8d9277f2e8fbf06fb8155a425135757bda","observation_id":"5c67bdb3-df0e-4848-85e1-4f34fadae142","resolution":{"observed_at":"2026-08-07T06:00:34.872136Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2015.24243","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:35.823437Z","title":"OptimalBiddingStrategyofBattery Storage in Power Markets Considering Performance-Based Regulation and Battery Cycle Life","venue":null,"work_id":"7f5620c7-ef41-4973-86d7-6a1d29f04f0f","year":2016},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.370354Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:4f4cbba4b3f72952e92126107a5d04822d377981a2ba7b8890dfb88383766eea","observation_id":"10cd732e-d107-4ecc-ae41-6535c5b2dd17","resolution":{"observed_at":"2026-08-07T06:00:35.922265Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1016/j.ifacol.2023.10.871","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:34.574222Z","title":"Optimal Economic Gas Turbine Dis- patch with Deep Reinforcement Learning","venue":null,"work_id":"e668dad0-3394-4050-9255-7b6ef3d5ac8e","year":2023},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.382775Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:f6c11a4b3aa87426744b4f53b571d4c99f2f6c6e3cd50173636b3968cd4d8204","observation_id":"816dfa36-ab19-47da-a08b-20112a3c9971","resolution":{"observed_at":"2026-08-07T06:00:34.655094Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:38.659315Z","title":"DeepReinforcementLearningforJointDispatchofBattery Energy Storage Systems and Gas Turbines in Microgrids withRenewableEnergy","venue":null,"work_id":"86b5da16-b9fa-4ed9-80cc-71abebb62a2d","year":2025},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.394320Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:1d8d7cc4ca7e38c31cea4fe8a071eb66c7e992274d767e3deeab310f8a83dfb5","observation_id":"62bb98b7-4f73-4ac1-a955-4fb485f81cfe","resolution":{"observed_at":"2026-08-07T06:00:38.779862Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:38.492618Z","title":"MIT press (2018)","venue":null,"work_id":"1c1c8a6c-2879-408a-b19e-11770695a323","year":2018},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.406465Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:3a1a0148a468beb273dd71da793b808461632f8372f2e301813d56045ead60a1","observation_id":"125f1032-b79d-41be-b389-668a708c2274","resolution":{"observed_at":"2026-08-07T06:00:38.563132Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:32.417700Z","title":"Human-level control through deep reinforce- ment learning","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.417700Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:f065f342db6c4120cb624cbb36a72ae7ae93a62dec2d80932dd4df58e098fe06","observation_id":"b1bd297c-7835-46d2-93ee-ffe351095506","resolution":{"observed_at":"2026-08-07T06:00:32.417700Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-07T06:00:32.427757Z","title":"Proximal policy optimization al- gorithms","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.427757Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:9706551ec4e9fec7e0179c78ce7890d20075ffbfa6d662e877b962254cd0f14c","observation_id":"a86d06c0-362a-46dc-b655-3fa83ff987bd","resolution":{"observed_at":"2026-08-07T06:00:32.427757Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:38.333747Z","title":"Market and sys- tem reporting","venue":null,"work_id":"cb161d0b-047f-40bb-976d-dec980ded092","year":2023},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.436894Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:14959fac9e5202ee96dffc638043aecb8f4fbd7f411614c72247b8e89ce82f5a","observation_id":"b52ca242-5bd7-48d3-b8b5-73ed44e60d65","resolution":{"observed_at":"2026-08-07T06:00:38.417415Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:32.457935Z","title":"Using bias-corrected reanalysis to simulate current and future wind power out- put","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.457935Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:ea2ba4c4d4e5b4c36489f578bf724727b54c87492f03876dacc04e16f9e5f94a","observation_id":"316a02eb-0b35-4643-974a-e6dccab31351","resolution":{"observed_at":"2026-08-07T06:00:32.457935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:38.190584Z","title":"Hydrogen and Power-to-X solutions - Elyzer P-300 - Technical Data","venue":null,"work_id":"789f377a-3026-42cd-becf-34de75f4f996","year":2024},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.513714Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:44ac17d2e8895f80743749257f780cff2e0bd25dbe55e0e580684a6c92b71ce9","observation_id":"60b9f76d-f5bb-4355-8dc6-b8beae868102","resolution":{"observed_at":"2026-08-07T06:00:38.256561Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1016/j.rser.2017.08.004","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:34.296385Z","title":"Power-to- Methane: A state-of-the-art review","venue":null,"work_id":"3f3c07aa-4cf7-46bc-a393-79247ed3b527","year":2018},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.570567Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:fa463fabf745bb489c63405a73b636f8eaa6eda90921da1286b4e214b3162787","observation_id":"1a33ad73-3c4f-4c79-ba27-675388c3c7b5","resolution":{"observed_at":"2026-08-07T06:00:34.439500Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2023.14472","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:35.506569Z","title":"Techno- economic evaluation of a power-to-methane plant : Lev- elized cost of methane, financial performance met- rics, and sensitivity analysis","venue":null,"work_id":"a30d927e-3288-4faa-a358-1218ed1026c2","year":2023},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.660252Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:16d3aba549e8436fb34d63c5a22cb970ca41f8c052c7ffa8b67fff47f6c6dbdb","observation_id":"a1c61de7-2868-4566-9faa-4ca57597da1a","resolution":{"observed_at":"2026-08-07T06:00:35.624995Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:38.046065Z","title":"Putting CO2 to Use","venue":null,"work_id":"8718494a-e5a1-4cf2-a789-72f66267f076","year":2019},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.783443Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:bbc32052f9cb8d10d0ae1168f786973a3ba725956a9e04ed2f120ffa6c78ec9e","observation_id":"3f31f088-2f4a-4195-81fa-7302b3f61779","resolution":{"observed_at":"2026-08-07T06:00:38.112434Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17032","last_updated":"2025-11-02T13:42:19Z","snapshot_observed_at":"2026-07-06T18:51:06.750136Z","submitted_at":"2024-07-24T06:35:05Z","title":"Gymnasium: A Standard Interface for Reinforcement Learning Environments","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.17032","snapshot_observed_at":"2026-08-07T06:00:32.846274Z","title":"Gymnasium: A Standard Interface for Reinforcement Learning Environments","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.846274Z"},"links":{"cited_paper":"/paper/2407.17032","citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:21b082c07e8eb837cccc25dfe300bf236e203f9a3bfe3e812a2b9bc090119f11","observation_id":"f03176e2-ed88-4b1c-8b2d-272d109d4eed","resolution":{"observed_at":"2026-08-07T06:00:32.846274Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:37.828748Z","title":"Stable- Baselines3: Reliable Reinforcement Learning Implemen- tations","venue":null,"work_id":"767f78de-b044-413d-8e9a-1af56ff7cbcc","year":2021},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.927146Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:bdda2c205d8b62bcfd551cabf7679424efb789d21f3b9e36e3ef245ee9e707aa","observation_id":"4d71ac74-88cc-4a5b-b21b-bcfde37b3a81","resolution":{"observed_at":"2026-08-07T06:00:37.948429Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:33.042659Z","title":"Optuna: Anext-generation hyperparameter optimization framework","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.042659Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:563370916b23804aea2d48fd32eea7058455df8c9831f4d4035c9bbfb03f5b6a","observation_id":"8a38627c-dc69-459b-9a23-0db1c05d1355","resolution":{"observed_at":"2026-08-07T06:00:33.042659Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1162/neco.2006.18.12.2936","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:34.116957Z","title":"Learning Tetris Us- ing the Noisy Cross-Entropy Method","venue":null,"work_id":"0c16db11-4f5f-44c5-94ff-29913ec2fb8c","year":2006},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.140416Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:fbc4fabfc8bb558357b9b296a9e51b9ba2cd85dcd90e3e5f666f4683a2aaec02","observation_id":"854e85a6-8ab5-40e1-9540-20259b0c54a7","resolution":{"observed_at":"2026-08-07T06:00:34.182270Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:37.641566Z","title":"Gurobi Optimizer Reference Manual","venue":null,"work_id":"9839ccac-20f9-46f2-9c62-54aab35d7142","year":2024},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.227795Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:4f2bd95dbec6ad5b3e3819d7405483dad890386ad9dc4f99f1d13ab65aa38fa9","observation_id":"b1ee8ff6-02cb-49d6-813e-706cc57ce994","resolution":{"observed_at":"2026-08-07T06:00:37.728083Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2013.22728","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:35.084506Z","title":"EnergyMan- agement for Lifetime Extension of Energy Storage Sys- tem in Micro-Grid Applications","venue":null,"work_id":"490b4d45-9300-4fe7-a164-04673f90fcc6","year":2013},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.314056Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:36194d17d7ac7988e0b6f10fa34863b5bd53ebc5ace1423815c24e64e7fa7ae3","observation_id":"13362af9-8821-42ab-9ce1-3cfabbdf3d97","resolution":{"observed_at":"2026-08-07T06:00:35.240447Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.3390/en11020469","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:33.925090Z","title":"A PSO- Optimized Fuzzy Logic Control-Based Charging Method for Individual Household Battery Storage Systems within a Community","venue":null,"work_id":"61968411-26e3-457a-9851-bd91f7d90514","year":2018},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.347120Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:ab8a0d2a28b561b6b4599f9ba5215860db290958731d159b15694c30b07d29ff","observation_id":"6ab02f95-e51e-474d-a316-902df3bc901b","resolution":{"observed_at":"2026-08-07T06:00:34.008838Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.2172/1984976","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:33.713131Z","title":"Cost Projections for Utility-Scale Battery Storage: 2023 Update","venue":null,"work_id":"df13c8a3-6946-400c-9949-8e9eb447adee","year":2023},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.427314Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:3580dc58c29dfadc6d06d3d6e2d81f6d552b99013e675eaecffe56e0d591c751","observation_id":"c8722501-de9e-4168-aed7-18320401e09b","resolution":{"observed_at":"2026-08-07T06:00:33.797856Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:37.471315Z","title":"Aeroderivative gas turbines","venue":null,"work_id":"49be2422-8d31-46e3-bcae-d269e808a7de","year":2017},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.522111Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:fbb721b31be9f408f1b2307a1901de5eede83037c2691b85049f119b3331fa86","observation_id":"095c15f6-4279-4901-ab05-dcb05de9a088","resolution":{"observed_at":"2026-08-07T06:00:37.545099Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:37.295685Z","title":"Capital Cost and Performance Characteristic Estimates for Utility Scale Electric Power Generating Technolo- gies","venue":null,"work_id":"583561e9-3714-480c-b19b-db6a1800aeaa","year":2020},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:33.591579Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:39e80d90721e4e12ba2bb5cabf3b13ac7d23e9156d73398a64e4acfa44510322","observation_id":"86960afa-0c98-4a84-9a73-d8616d2e6026","resolution":{"observed_at":"2026-08-07T06:00:37.372890Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"3008.2021","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T06:00:36.215926Z","title":null,"venue":null,"work_id":"ee30ec6c-effb-4839-9a36-ff2cc343f224","year":2021},"citing_paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:32.337031Z"},"links":{"citing_paper":"/paper/2506.06484"},"observation_digest":"sha256:0eec2a09ff4c8e6facd0c13782e6c9337ae5b608eca18c69dd344508f1e135a1","observation_id":"0491b14c-566b-4e81-8ba5-b18c2d32725a","resolution":{"observed_at":"2026-08-07T06:00:36.320265Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.06484","last_updated":"2025-06-06T19:20:08Z","latest_version":1,"primary_category":"eess.SY","snapshot_observed_at":"2026-08-07T05:53:19.223987Z","submitted_at":"2025-06-06T19:20:08Z","title":"The Economic Dispatch of Power-to-Gas Systems with Deep Reinforcement Learning:Tackling the Challenge of Delayed Rewards with Long-Term Energy Storage"},"reference_resolution":{"displayed":32,"state_counts":{"malformed_identifier":0,"metadata_mismatch":6,"parse_uncertain":0,"unresolved":9,"verified_exact":7,"verified_fuzzy":10},"total_outbound_references":32},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 32 of 32 outbound references and 0 inbound Pith citation observations for arXiv:2506.06484."}