{"as_of":"2026-08-12T13:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:85ec0f191cea720ef78bd52b92e5c2554ddb60e373b7eaaf32264253a0173d2f","coverage":[{"denominator":34,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":34,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T20:13:24.303202Z","state":"measured"},{"denominator":34,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":34,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2502.05118/citation-record","integrity":"/paper/2502.05118/integrity","json":"/paper/2502.05118/citation-record.json","paper":"/paper/2502.05118"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1111/j.1439-","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:24.833065Z","title":"Die angeborenen Formen m ¨oglicher Erfahrung,","venue":null,"work_id":"a3bd2e28-4631-4939-8f41-5bff087afd24","year":1943},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.561557Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:e622913d40f8bd2e45bead05c846c4e1dae601016843061ab8aa15d1bd3374b2","observation_id":"587be668-2235-481a-9abf-fc80fa68409e","resolution":{"observed_at":"2026-08-08T20:13:24.892850Z","resolver_source":"doi_truncated","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:26.019868Z","title":"The (Ir)relevance of Robot Cuteness: An Exploratory Study of Emotionally Durable Robot Design,","venue":null,"work_id":"add9ed77-23bf-45ca-a649-210c90ff541b","year":2020},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.594901Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:c4325e252f2b99845694729764477de21d11c57c7dba1032f2b39bb73e19d0c6","observation_id":"8f68b775-60c0-4665-b225-12b892c39444","resolution":{"observed_at":"2026-08-08T20:13:26.022818Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:26.011263Z","title":"Research on the influence of the baby schema ef- fect on the cuteness and trustworthiness of social robot faces,","venue":null,"work_id":"506c4939-a034-4853-a3d4-000941245ce8","year":2023},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.678344Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:de580b9766018c480151318f82b4fb7cbc6762f6370f52bca45eeb9ffbf6dd06","observation_id":"151688b2-6a90-4709-a7df-8127aaefe889","resolution":{"observed_at":"2026-08-08T20:13:26.014601Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:23.704094Z","title":"Human-robot interaction: the impact of robotic aesthetics on anticipated human trust,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.704094Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:226796853e3c52a2bea7ee2c0915365ba6efc32696f87aba55e376668e1b8ce8","observation_id":"0f716e75-a2e2-41c5-820b-e4c67a58b96d","resolution":{"observed_at":"2026-08-08T20:13:23.704094Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:26.002744Z","title":"Reinforcement learning from human reward: Discounting in episodic tasks,","venue":null,"work_id":"a3774f8b-cd78-4e65-b888-2773f7f6bcb9","year":2012},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.744408Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:222f72031ed0a2d5549077eb2273d1519f0f73b6be4ea8a848aae6bede91ae24","observation_id":"f32c1291-5260-4d68-aa0b-b0f21cfcbdd1","resolution":{"observed_at":"2026-08-08T20:13:26.006315Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.993825Z","title":"How humans teach agents: A new experimental perspective,","venue":null,"work_id":"8cf1ac54-194c-42a5-81e9-61a152643058","year":2012},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.769538Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:eb0613609a3b2aac943ab676f53bd24817158e0053debe734db47fca8a83e991","observation_id":"302b29b9-3324-46c8-87f7-5f07741ac7d9","resolution":{"observed_at":"2026-08-08T20:13:25.996972Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.943489Z","title":"Teachable robots: Understanding human teaching behavior to build more effective robot learners,","venue":null,"work_id":"9a698268-3111-4fdf-bdb2-ab146024a6cd","year":2008},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.785065Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:d8a30314fbd6f38b0057ef9e887e0948698bb6c76ea38aee4e39647714ffc81c","observation_id":"c9b4701c-5580-49ce-b058-ae04f220d8d5","resolution":{"observed_at":"2026-08-08T20:13:25.971270Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.824500Z","title":"Learning and Comfort in Human–Robot Interaction: A Review,","venue":null,"work_id":"654552c1-1fbc-4050-887b-e796db21cd1e","year":2019},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.790486Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:ad4afb9c0ea8f8c02f4802453e6b76dcf8cfdd045688dad41c3c737a0f353b46","observation_id":"097515ea-f937-4499-928e-18613759f778","resolution":{"observed_at":"2026-08-08T20:13:25.880237Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.757646Z","title":"Recent advances in robot learning from demonstration,","venue":null,"work_id":"76d77de3-25e7-4571-a78f-d4233bd9dec9","year":2020},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.795946Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:21fc9af3782d0dedf23c73f626af0d6fade5bb18814e20bb6c0af4beaf4a7c3d","observation_id":"148b510e-def6-4550-a9ec-68d8c4784491","resolution":{"observed_at":"2026-08-08T20:13:25.783186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.748893Z","title":"A survey of robot learning from demonstration,","venue":null,"work_id":"07ca2894-dc29-47a5-a431-93916f53088b","year":2009},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.799930Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:2bd71194b46fb7ed6140d185ea5a820fed6ac2a47b03f3947a78aaee0f231963","observation_id":"866eaec8-3c00-4c0c-b702-a0384d94c9d1","resolution":{"observed_at":"2026-08-08T20:13:25.751891Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.739768Z","title":"Robot learning from demonstration,","venue":null,"work_id":"7ef639b3-3b61-49d4-a7c4-b0117b76dcc7","year":1997},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.803003Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:3ccdaf4f506779f521031a0ec5cf3d37c99d5cd7c4bb53118ea4c8c6dac92e18","observation_id":"45a5e61c-ca47-48f3-8b8d-c3e3b7d2473c","resolution":{"observed_at":"2026-08-08T20:13:25.743312Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.731828Z","title":"Robot learning by demonstration,","venue":null,"work_id":"4e74fb40-14b2-4fc2-b802-4f6c5bdc3d1d","year":2013},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.813881Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:55e4d7fa55fe3af14501a770dcd502c6e38f985f608811d4be51948b8355b1cf","observation_id":"f5a7673d-86da-4648-a8e6-73be5f9636d4","resolution":{"observed_at":"2026-08-08T20:13:25.734796Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.704677Z","title":"Learning from demonstration,","venue":null,"work_id":"2c121998-a1db-4cde-8042-b0468c839d6f","year":1996},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.829025Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:8a5f3a474740438bba45b6f71238b8c8500b63484b891b2d557fa7d038458980","observation_id":"ee314380-b9fd-47f4-b930-a88970dd4e2a","resolution":{"observed_at":"2026-08-08T20:13:25.722157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.657216Z","title":"A robot learning from demonstration framework to perform force-based manipulation tasks,","venue":null,"work_id":"27220944-40c5-4223-a499-abcc19a503f7","year":2013},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.834091Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:1288aecf7265386b9dadbf92b16b59751c176240c575ce541708835c282be134","observation_id":"a220d1d3-dac4-435e-b321-cdb8e1b8933a","resolution":{"observed_at":"2026-08-08T20:13:25.685037Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.562879Z","title":"Interactive and Explainable Robot Learning: A Comprehensive Review,","venue":null,"work_id":"eed8867f-e941-4c5d-b4e7-ee94330ff707","year":2024},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.836624Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:891e242eb812edc88b112ab5b71fecc2b1c258ecbf23fc5b39bfd7fc6a5ecf18","observation_id":"49bd8cc8-3cff-4514-a25e-afc7fde6e076","resolution":{"observed_at":"2026-08-08T20:13:25.623187Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.521481Z","title":"Human-Robot Alignment through Interactivity and Interpretability: Don’t Assume a ’Spherical Human’,","venue":null,"work_id":"593f6cee-43ac-473f-b905-6398fdb8e740","year":2024},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.839899Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:6d3d4813176ac80cd7d069844ef7cd5033ec0c2a48d7352a2c87a8047ea79b28","observation_id":"afe77120-6201-422f-a1aa-efa5765bf2cd","resolution":{"observed_at":"2026-08-08T20:13:25.524927Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.490628Z","title":"Extrapolating beyond suboptimal demonstrations via inverse reinforcement learning from observations,","venue":null,"work_id":"600b34a7-4337-46cc-991d-b37c0c2b8bb7","year":2019},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.843163Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:95b4b2804b5c22bc7c9b89439b434ac6724a384186549148cc117db9c52abc23","observation_id":"6517d7fa-8a00-413a-b9ab-3faf7fc9ff04","resolution":{"observed_at":"2026-08-08T20:13:25.502361Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.480068Z","title":"Learning from suboptimal demonstration via self-supervised reward regression,","venue":null,"work_id":"495ea184-7863-4f56-8379-275115a1a87f","year":2021},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.845953Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:74e1a299448ffa0d8581c63c4030a2cf70b9902b41f8a053ae60e6167c28c5e3","observation_id":"5a262fe9-2edd-4b21-af29-2ff6c44a857b","resolution":{"observed_at":"2026-08-08T20:13:25.483215Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.437140Z","title":"Droid: Learning from offline heterogeneous demonstrations via reward-policy distillation,","venue":null,"work_id":"acd841dc-b0a2-47b4-9842-63f5828a9606","year":2023},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:23.942748Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:294224df1a8ac009ca5f48e793d7d0ca3d0539f436bb382c530a59c4ac0ce25d","observation_id":"791a1de6-8646-43b9-91d8-5c1c8cc41d64","resolution":{"observed_at":"2026-08-08T20:13:25.448597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.385982Z","title":"Beyond Success: Quantifying Demonstration Quality in Learning from Demonstration,","venue":null,"work_id":"678d220d-2ceb-491d-919e-ec6fc622ddb4","year":2024},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.008318Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:ab952807fc8d40474022a30c19201c4a4b1a8b9044c00fddef460e7067f367e9","observation_id":"7cc1ef4e-99ab-42e3-8d1b-72dcde49b15c","resolution":{"observed_at":"2026-08-08T20:13:25.430215Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2008.46408","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:24.755586Z","title":"TAMER: Training an Agent Manually via Evaluative Reinforcement,","venue":null,"work_id":"2f5e7793-8d35-4393-80c8-c3596f746598","year":2008},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.088897Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:703e8c51d0fd1a339311cdc94f7025345c9d8ad3bbdd0efdcd3c715ffff148f0","observation_id":"4bf574e5-86eb-4dda-8443-0f4976c061f9","resolution":{"observed_at":"2026-08-08T20:13:24.761288Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.254912Z","title":"Deep TAMER: Interactive agent shaping in high-dimensional state spaces,","venue":null,"work_id":"4c0f55a4-ff51-41ce-9ceb-34af626841a5","year":2018},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.146614Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:04f95a438d17075bb01eae704e573a4b608642883ab8092e87c167bbfe69bf50","observation_id":"66105df6-c82b-4158-97e4-f1f0027b0bb7","resolution":{"observed_at":"2026-08-08T20:13:25.302181Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.133240Z","title":"COACH: Learning continuous actions from corrective advice communicated by humans,","venue":null,"work_id":"d56692ad-4e2a-4c8f-bc85-ed7e09e83066","year":2015},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.193316Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:fae85d8388880753801c583a20a0ec5c82a1d9dc6ef8fea7da528316916d015a","observation_id":"45d10715-ec1f-4e80-806b-036fcd6711c3","resolution":{"observed_at":"2026-08-08T20:13:25.174236Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1902.04257","last_updated":"2019-02-12T06:45:21Z","snapshot_observed_at":"2026-08-08T14:49:33.518245Z","submitted_at":"2019-02-12T06:45:21Z","title":"Deep Reinforcement Learning from Policy-Dependent Human Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1902.04257","snapshot_observed_at":"2026-08-08T20:13:24.221474Z","title":"Deep reinforcement learning from policy-dependent human feedback,","venue":null,"work_id":null,"year":1902},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.221474Z"},"links":{"cited_paper":"/paper/1902.04257","citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:535bd3a25af0f81046978b1a8ce50a11e182f5545a0bc4144605b98a8e896f83","observation_id":"318f390a-922d-4db5-a29c-305866b18b4a","resolution":{"observed_at":"2026-08-08T20:13:24.221474Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15217","last_updated":"2023-09-11T17:25:24Z","snapshot_observed_at":"2026-08-03T19:11:09.671782Z","submitted_at":"2023-07-27T22:29:25Z","title":"Open Problems and Fundamental Limitations of Reinforcement Learning from Human Feedback","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15217","snapshot_observed_at":"2026-08-08T20:13:24.242292Z","title":"Open problems and fundamental limitations of reinforcement learning from human feedback,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.242292Z"},"links":{"cited_paper":"/paper/2307.15217","citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:f2eb0bd579ddf16eeefb702660a5dba74fb0784d825d5384a2067a5f06bfd714","observation_id":"0a55315f-54a4-494a-955d-74c839955e2f","resolution":{"observed_at":"2026-08-08T20:13:24.242292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.124701Z","title":"Learning optimal advantage from preferences and mistaking it for reward,","venue":null,"work_id":"61645136-1e71-4026-bc9d-1d67e5b14b2e","year":2024},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.245786Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:d57540ec7886cff6d7d6f994e7c5a0e02b664b85b2c3c5ddb910fded6cc6393e","observation_id":"c0e0c9b0-8428-43f7-b11c-4802836ef26d","resolution":{"observed_at":"2026-08-08T20:13:25.127723Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.12773","last_updated":"2023-10-19T14:22:03Z","snapshot_observed_at":"2026-08-02T16:56:38.535065Z","submitted_at":"2023-10-19T14:22:03Z","title":"Safe RLHF: Safe Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.12773","snapshot_observed_at":"2026-08-08T20:13:24.249006Z","title":"Safe RLHF: Safe reinforcement learning from human feedback,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.249006Z"},"links":{"cited_paper":"/paper/2310.12773","citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:ea5758c1f968e2f52000c64a342ec5613e2e2a2335ab1efbe75f2d64685cb5a4","observation_id":"1e5280a4-2037-4345-8fd6-82e68c16fe61","resolution":{"observed_at":"2026-08-08T20:13:24.249006Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.113536Z","title":"RLHF Deciphered: A Critical Analysis of Reinforcement Learning from Human Feedback for LLMs,","venue":null,"work_id":"67dcf437-37cf-47c2-96d4-78864027ad9e","year":2024},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.252415Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:7e9358858d0d7e81e5269b873d0f8489cad071c46130f560cede5c6f3d6640a9","observation_id":"bf407659-ed73-481d-9d20-ad32704da0d7","resolution":{"observed_at":"2026-08-08T20:13:25.117641Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.102868Z","title":null,"venue":null,"work_id":"228ca849-2334-4100-ab0c-75250dfd37c1","year":2004},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.259869Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:f751c69db46f52a4417c23c9bde3ca484d1dad50ba0c00fa0789293705dc1312","observation_id":"f8aa5a43-1187-44e1-9070-b707507ab0a2","resolution":{"observed_at":"2026-08-08T20:13:25.106915Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.082833Z","title":"Fostering Cross-Cultural Research by Cross-Cultural Student Teams: A Case Study Related to (Cute) Robot Design,","venue":null,"work_id":"82cabc9a-8521-4b99-b759-4a469804cc7f","year":2020},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.269673Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:21487eeb69b3d814af1ce92f57d33864549ed0cc767641a74e7a1fa1201db214","observation_id":"949d124d-68dc-48e4-89ee-a17800aaf8ff","resolution":{"observed_at":"2026-08-08T20:13:25.096144Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2010.55986","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:24.582098Z","title":"Austermann, S","venue":null,"work_id":"6d616f01-eedb-4810-af97-d58ff292f458","year":2010},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.276956Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:3c66646be516fcbc1b1de09374aecb48bf63c79bc78717dde55577d5785aa946","observation_id":"4d8e95cb-4dd6-4ccf-8aac-780056ff80ce","resolution":{"observed_at":"2026-08-08T20:13:24.658480Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"0310.2008","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:24.411036Z","title":"Baby Schema in Infant Faces Induces Cuteness Perception and Motivation for Caretaking in Adults,","venue":null,"work_id":"c431ce05-84ff-4515-b2c1-b95dc65de0a5","year":2009},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.284407Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:e55d6fe24e5c1c2d7c3875f98bc7c04865128c1434f190c5523cbc2e97ca4baf","observation_id":"210f5f41-6e0c-4ddf-baa3-c81471197d1a","resolution":{"observed_at":"2026-08-08T20:13:24.455641Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:25.031654Z","title":"Isn’t it cute: An evolutionary perspective of baby-schema effects in visual product designs,","venue":null,"work_id":"c01929c5-1289-481a-a4e7-4d06c8278dd8","year":2011},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.292797Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:c9aa940933bf6a064b727ddca6d1f81a4fa99ea3b75c9d4a5d5dd13a3d577c46","observation_id":"d7d78565-3899-480b-9fb2-5da82cc6b1b9","resolution":{"observed_at":"2026-08-08T20:13:25.055152Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T20:13:24.953142Z","title":"Interactively shaping agents via human reinforcement: the TAMER framework,","venue":null,"work_id":"36cf2c3f-ecf1-4824-9d5c-646c1476e3c3","year":2009},"citing_paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-08T20:13:24.303202Z"},"links":{"citing_paper":"/paper/2502.05118"},"observation_digest":"sha256:72a314fd8096bf763d6cd940c6853000ee57610cfdb910bbe740d2db636d3e3b","observation_id":"c1460918-f45c-4fda-8dda-ea3585949a1e","resolution":{"observed_at":"2026-08-08T20:13:25.009509Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.05118","last_updated":"2025-02-07T17:41:29Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-09T16:18:13.396327Z","submitted_at":"2025-02-07T17:41:29Z","title":"Use of Winsome Robots for Understanding Human Feedback (UWU)"},"reference_resolution":{"displayed":34,"state_counts":{"malformed_identifier":1,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":5,"verified_exact":1,"verified_fuzzy":25},"total_outbound_references":34},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 34 of 34 outbound references and 0 inbound Pith citation observations for arXiv:2502.05118."}