{"as_of":"2026-08-07T20:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:12b4a569e081db612195692e259139e386b983c193ffdc313fd44e4e1751097d","coverage":[{"denominator":27,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":27,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-24T11:16:13.986064Z","state":"measured"},{"denominator":27,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":27,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2207.09845/citation-record","integrity":"/paper/2207.09845/integrity","json":"/paper/2207.09845/citation-record.json","paper":"/paper/2207.09845"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mastering the Game of Go Without Human Knowledge","venue":null,"work_id":"a5286b2e-0218-4aa7-bbac-5acdeada8c45","year":2017},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:808d3557c72e2d9b9cd648969aef583556cc323cf31a86108651c1e1d61581fa","observation_id":"06256d11-4f35-48cb-9949-bbb3b2576639","resolution":{"observed_at":"2026-05-24T11:19:24.007176Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A Survey on Inter- active Reinforcement Learning: Design Principles and Open Challenges","venue":null,"work_id":"7990dc75-e6b5-4e89-8999-a9691246869d","year":2020},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:45e254a62dcaf24d5d99b427243cc0b279b62577dcb23c9888a07b0c2207cc13","observation_id":"4bda5c47-1a60-4ebb-9792-c2dca6e69334","resolution":{"observed_at":"2026-05-24T11:19:24.009987Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Multi-Agent Reinforcement Learning: Independent Vs","venue":null,"work_id":"73184aef-920d-47a5-b637-3d2b7335b638","year":null},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:9dc1a1a49388a4dcfda34b75e634f493b9eb3825f9a73c04350babe0d6b1cce4","observation_id":"7c2c2d2b-9e8d-43f8-af82-5146e77e8919","resolution":{"observed_at":"2026-05-24T11:19:24.012417Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Agents Teaching Agents: A Survey on Inter-Agent Transfer Learning","venue":null,"work_id":"8edf3093-92fa-4bee-b8c9-13cb34efa420","year":2019},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:80ff735e5e44a036ac5abd8ab79b8e894fdda7ac54a5f5b440a57cfbe77e8a1d","observation_id":"ace26c2d-3362-4156-9998-b95541231519","resolution":{"observed_at":"2026-05-24T11:19:24.062618Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Policy Invari- ance Under Reward Transformations: Theory and Application to Reward Shaping","venue":null,"work_id":"3cac094f-ace6-4b17-a106-43dfe3698aa3","year":1999},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:e22234e429336dffe921a66d446630d563cf1d68416f2d89db20b451982eff12","observation_id":"5ed2930e-8501-42e3-99f2-cc68b0059db8","resolution":{"observed_at":"2026-05-24T11:19:24.059758Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Policy Shaping: Integrating Human Feedback with Reinforcement Learning","venue":null,"work_id":"3f1080ac-3ab1-49f2-afd3-53c6e23c15e2","year":2013},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:e4575b6bff8eb618f6a5c9aa3fb37a172c6abb3e0a8abec7f4c53ffb8140eab9","observation_id":"d5ff91b4-e18b-4948-8811-9b5c51fb8c3e","resolution":{"observed_at":"2026-05-24T11:19:24.057236Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.17185/duepublico/","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Interaction in Reinforcement Learning Reduces the Need for Finely Tuned Hyperpa- rameters in Complex Tasks","venue":null,"work_id":"898d7f72-afec-4668-8f82-0365a8667e6d","year":2015},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:0861cc00c56d0b770cc27fcb4c9c6abb6436df866cc6dba39302faeca7b48e81","observation_id":"cbc1b3fc-12b0-45a9-a055-40112062358b","resolution":{"observed_at":"2026-05-24T11:19:23.112960Z","resolver_source":"doi_truncated","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"People Teach with Rewards and Punish- ments as Communication, Not Reinforcements","venue":null,"work_id":"605538bd-ddc2-4db0-93f6-6ab9f559e65d","year":2019},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:43394230930a7cadc1e3a1744212d15aa50c3fde921e110885adc19d7289aa9e","observation_id":"481ca3f6-f413-43e7-ad31-1c2fd2c5e26c","resolution":{"observed_at":"2026-05-24T11:19:24.054765Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1016/j.artint.2007.09.009","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Teachable Robots: Understanding Human Teaching Behavior to Build More Eﬀective Robot Learners","venue":"Artificial Intelligence","work_id":"c01aacc9-4631-430a-836d-c6b32d2facb5","year":2008},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:b7a532aba9710688cf3928ec029be6365ef059cd296d1af4e2d8878aae97abf2","observation_id":"ec027e70-bb4c-4542-bd74-8380756291d8","resolution":{"observed_at":"2026-05-24T11:19:23.098847Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A Strategy-Aware Technique for Learning Behaviors from Discrete Human Feedback","venue":null,"work_id":"81cefc89-574f-45af-a3bf-13c685d9a795","year":2014},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:c59c502032af280881c200a7914133f1121fa630d44ce5d8c8b3ef845373e69b","observation_id":"e7c2d1ac-5956-478b-b24d-007113b0c3d5","resolution":{"observed_at":"2026-05-24T11:19:24.052064Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Reinforcement Learning from Human Reward: Discounting in Episodic Tasks","venue":null,"work_id":"1e4e1e43-f65d-4acc-b16a-a643f9a16b41","year":2012},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:d5ab22cd0fd9ad19d3d907c12c5a68ae7e514bc9ff1414360f6a7ae852ddca76","observation_id":"a598f88e-624e-4058-846d-f0cdc8e58ed1","resolution":{"observed_at":"2026-05-24T11:19:24.049225Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"An Evaluation Methodology for Interactive Reinforcement Learning with Simulated Users","venue":null,"work_id":"77848305-49f8-4519-b3d4-5f229e920f2e","year":2021},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:d37e1c738969642328ff00b170e5ff182978f8ef5ce8e98e12dab576be6c4ee1","observation_id":"a5a68c3a-6c56-45ba-874d-2254d5bda5d1","resolution":{"observed_at":"2026-05-24T11:19:24.046849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"0091.2014","doi":"10.1080/09540091.2014.885279","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Reinforcement Learning Agents Providing Advice in Complex Video Games","venue":"Connection Science","work_id":"a66b1f0a-4c53-4182-8b47-3555fc708c23","year":2014},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:550a56af6de6ae95de02da74b94e5108d33544a3ed4441e9baee742d4d9c8797","observation_id":"a848e805-6b07-4314-9263-72191a268297","resolution":{"observed_at":"2026-05-24T11:19:23.108916Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Agent-Advising Approaches in an Interactive Rein- forcement Learning Scenario","venue":null,"work_id":"da6ac65f-8251-42d4-b59a-a269af4e2061","year":2017},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:396b9f12f9f4bc618328f6d53f1999172be3d07a618527cb709ba4a79b68f7ba","observation_id":"49781342-f7f3-4e2c-9f5c-829c206ceb6c","resolution":{"observed_at":"2026-05-24T11:19:24.044243Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Eﬀect of Human Guidance and State Space Size on Interactive Reinforcement Learning","venue":null,"work_id":"fada49fc-cb48-42bc-bdb8-ae881a5c81b0","year":2011},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:e0faf6c90ff92b658229d74512bbe0e5cecb5ada7c0e1db8788adb1aca7c37e1","observation_id":"a5a6beee-a984-480a-909f-1776e1f1af96","resolution":{"observed_at":"2026-05-24T11:19:24.041302Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Interaction Is More Beneﬁcial in Com- plex Reinforcement Learning Problems Than in Simple Ones","venue":null,"work_id":"35a840fa-be43-48f3-8f71-0f738f4cc868","year":2015},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:3de52beeab4af1de34874bfc23fffbc63b0c840c471d44b2752ed870a95172ca","observation_id":"fbf75d0f-362f-4013-ac13-6e8f3c847158","resolution":{"observed_at":"2026-05-24T11:19:24.038628Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A Robust Approach for Continuous Interactive Reinforcement Learning","venue":null,"work_id":"a79304b9-a730-4c31-86da-d57bd7973907","year":2020},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:36ea8ed10e65e5b50824078b29b7b11ac837a5a2b97257764ecafa145dbf2dcf","observation_id":"3063bfaa-8f43-42b2-a052-ffac519e0f6d","resolution":{"observed_at":"2026-05-24T11:19:24.035757Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2016.25438","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Training Agents with Interactive Reinforcement Learn- ing and Contextual Aﬀordances","venue":null,"work_id":"86e3e103-c4cc-45f3-a915-d68a122e3dbe","year":2016},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:1ebe5deda8d402bcfefc1181d81931b8d0b908d2b60add81e7da10178cd44baf","observation_id":"a49003e5-7ef3-4f4d-a72a-9244ab3cabce","resolution":{"observed_at":"2026-05-24T11:19:23.239320Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1007/s10846-013-0015-4","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Com- plete Analytical Forward and Inverse Kinematics for the NAO Humanoid Robot","venue":"Journal of Intelligent & Robotic Systems","work_id":"16dc6f03-fdb7-4b93-9434-3754e67b27f8","year":2015},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:a613e1dd04902a8f4cf0aeb78cf2a290cfec0a9841e300d588550c94c2cc4cb6","observation_id":"b02063a5-e926-4405-a962-39aae3274048","resolution":{"observed_at":"2026-05-24T11:19:23.103001Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Task- Oriented Rigidity Optimization for 7 DoF Redundant Manipulators","venue":null,"work_id":"f6ff5f6d-7e85-4dc4-823b-ff2e68abbb2a","year":2017},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:1ac4c7e69b54ab4a7a6351efb4a89b56b085819a8e3fde3b1ff396628fd1ef37","observation_id":"607bb113-7f51-4818-b40a-6f256ada2e7e","resolution":{"observed_at":"2026-05-24T11:19:24.032637Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2017.00010","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Improv- ing Robot Motor Learning with Negatively Valenced Reinforcement Signals","venue":null,"work_id":"4df9b42e-4126-4b2b-a20a-72cb3e04cc46","year":2017},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:c33094bc226a9608a8525d83dc866e95d5aa90313d904a4e6634fde8280d7a02","observation_id":"c74d950e-a9fb-4cb4-a063-3edb980551a6","resolution":{"observed_at":"2026-05-24T11:19:23.244566Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The Eﬀects on Adaptive Behaviour of Negatively Valenced Signals in Reinforcement Learning","venue":null,"work_id":"1d7e37b3-b83e-4a72-a9c1-bc71a2d4068e","year":2017},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:372e262ccdea34c05e5ff1c34a1f789b95addb4144cd60b2a94f9691aea457c4","observation_id":"fed71bae-0c28-4312-b00b-9a3543bb83c2","resolution":{"observed_at":"2026-05-24T11:19:24.027495Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Reinforcement Learning in Continuous Action Spaces","venue":null,"work_id":"7e7d5a55-b5b3-47f9-b374-7b43b3bbb883","year":2007},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:be954986be4dc3213d1542c7c090e5d02f930456b240781b562902a643ead59b","observation_id":"df1dd645-6362-42fe-aa15-8729a59d9f0c","resolution":{"observed_at":"2026-05-24T11:19:24.030001Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Adam: A Method for Stochastic Optimization","venue":null,"work_id":"d8ec1a96-3b33-42d8-99f1-0ad327bec089","year":2015},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:c1f71b05c8790ac50e8a6bd3f871a8d582dc32c7b37075d2c4f237e04c4431bf","observation_id":"663b4a81-cc9b-4466-9631-5ecf9c1ec8fe","resolution":{"observed_at":"2026-05-24T11:19:24.024737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Transfer Learning for Rein- forcement Learning Domains: A Survey","venue":null,"work_id":"5f50655c-c7be-4f16-b8c0-181f8d89568e","year":2009},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:fd2ab698d4c6add12e68f26beb7b15c693a0242cbd7b24ed0c375ee25da58048","observation_id":"8ecdd784-8736-42b9-9e23-e24311638866","resolution":{"observed_at":"2026-05-24T11:19:24.021661Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Hyperopt: A Python Library for Optimizing the Hyperparame- ters of Machine Learning Algorithms","venue":null,"work_id":"ee1e300d-9377-4525-84f2-590375623296","year":null},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:f49b1dee1f18ff0f3866d1a2e75b29a419aec1a6ae39b72493f753c2b977e1a0","observation_id":"b954f5c1-0043-49c8-ada5-fb0f28cc1602","resolution":{"observed_at":"2026-05-24T11:19:24.016299Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Understanding and Preventing Capacity Loss in Reinforcement Learning","venue":null,"work_id":"f8a1ee20-42b4-450d-9dfc-61c62a50240e","year":2022},"citing_paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-24T11:16:13.986064Z"},"links":{"citing_paper":"/paper/2207.09845"},"observation_digest":"sha256:c0c921eb9d76b7dbaa57eb86448b26e73e1fff58d7651cf1363dbf41ee64d6b6","observation_id":"b9e0d797-5951-4910-b360-803679850910","resolution":{"observed_at":"2026-05-24T11:19:24.018861Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2207.09845","last_updated":"2023-03-15T16:06:29Z","latest_version":2,"primary_category":"cs.RO","snapshot_observed_at":"2026-07-06T13:33:20.950270Z","submitted_at":"2022-07-20T12:17:02Z","title":"Quantifying the Effect of Feedback Frequency in Interactive Reinforcement Learning for Robotic Tasks"},"reference_resolution":{"displayed":27,"state_counts":{"malformed_identifier":6,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":3,"verified_fuzzy":18},"total_outbound_references":27},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 27 of 27 outbound references and 0 inbound Pith citation observations for arXiv:2207.09845."}