{"as_of":"2026-08-09T23:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:4b6d62607d83f406047180d29d2da5b97442d690e57f3403f76a0311cbc11849","coverage":[{"denominator":23,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":23,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:20:28.654352Z","state":"measured"},{"denominator":23,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":23,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.22011/citation-record","integrity":"/paper/2505.22011/integrity","json":"/paper/2505.22011/citation-record.json","paper":"/paper/2505.22011"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:32.266246Z","title":"Learning human activities and object affordances from RGB-D videos,","venue":null,"work_id":"06274729-2002-4328-97b6-f8205c2c739b","year":2013},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.180321Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:14856d0d666552192145116278b58668af46e98c04a6bfaba21178afb825ba28","observation_id":"607879c8-3cb0-42ce-9530-e674ebcfb594","resolution":{"observed_at":"2026-08-07T13:20:32.305318Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:32.104494Z","title":"Learning human-object interactions by graph parsing neural networks,","venue":null,"work_id":"bd5a1396-0842-4466-8393-2a1bd1d96f90","year":2018},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.223685Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:ace1f323bcbf1f8deaa0feff740223bd7b7cbc61619e94e146e1bcf57c253be0","observation_id":"cc5f0cb2-e7d1-4c0f-bac7-876f8b6e690b","resolution":{"observed_at":"2026-08-07T13:20:32.189617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:31.990750Z","title":"Ramakrishnan G. Lighten: Learning interactions with graph and hierarchical temporal networks for HOI in videos ,","venue":null,"work_id":"015ee161-a653-41ce-a00f-3154e8ab0582","year":2020},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.356568Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:621f1d31bd418a6ced6a8cb739b81660ff9fbcaf449a8528680fbc6566f4d194","observation_id":"d39a3f09-fa35-4638-9ec1-9e02a00c3660","resolution":{"observed_at":"2026-08-07T13:20:32.035535Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:31.801436Z","title":"QPIC: Query -based pairwise human -object interaction detection with image -wide contextual information,","venue":null,"work_id":"63b7de09-ec81-440d-af1b-4ede51b8c197","year":2021},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.482485Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:478121c0151c173394383cd414887f6bc9d98009c2dd78bb860819cbb90c465d","observation_id":"1a8335fa-dd19-4a23-bfb0-d4f4eedf4265","resolution":{"observed_at":"2026-08-07T13:20:31.880537Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:31.703000Z","title":"Spatial-temporal transformer for dynamic scene graph generation ,","venue":null,"work_id":"cd6956ff-7a25-44c6-a09d-192930175329","year":2022},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.601723Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:e626d08dfbdf77c365d349d90b0f7cb1947bdb77abf9bdc4b66e5e04682e88ac","observation_id":"818852d9-7074-4824-90e3-fe59dfbf709f","resolution":{"observed_at":"2026-08-07T13:20:31.730239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:31.574037Z","title":"Video-based human- object interaction detection from tubelet tokens ,","venue":null,"work_id":"1d556ebf-ec85-4a7c-9e1f-51f3be759e3d","year":2022},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.730633Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:fba4f11690119799b3c337dbdba437f16b067be11d9d39f16e04f38ecff56a81","observation_id":"76c68c9d-0fb0-44e9-bfec-85f52d87fed1","resolution":{"observed_at":"2026-08-07T13:20:31.660405Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:31.402387Z","title":"End-to-end video scene graph generation with temporal propagation Transformer,","venue":null,"work_id":"12f59197-c138-489d-ae36-3b9979a8c99f","year":2024},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.876363Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:97db881ef75728b0ba95e213288ddec2f2b1d110e9f73fc9f33f693275009fb0","observation_id":"65e94e5f-2859-4072-9b1e-cb316f9ff586","resolution":{"observed_at":"2026-08-07T13:20:31.483271Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:31.150972Z","title":"ST- HOI: A spatial -temporal baseline for human -object interaction detection in videos ,","venue":null,"work_id":"10859ace-8bb3-439d-91df-1065c792917a","year":2021},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:26.989045Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:5b925b374ad1b11288aae05418b4d3cae96707235a5b2a70ac68d6f73518a132","observation_id":"9fbfde67-9e56-485b-821c-6acb0d64e9de","resolution":{"observed_at":"2026-08-07T13:20:31.282528Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:30.969141Z","title":"From detection to understanding: A survey on representation learning for human -object interaction,","venue":null,"work_id":"eece977d-c11d-41ee-8119-cb865ed2f496","year":2023},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.103137Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:f7868a073f4965230327485208b4534049f86599ebc022ee78288945c2c8b7c8","observation_id":"31ef7819-f7c2-440f-a63b-444239dfcd9a","resolution":{"observed_at":"2026-08-07T13:20:31.059386Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:30.756660Z","title":"Human-object interaction prediction in videos through gaze following ,","venue":null,"work_id":"1d20fc3f-7609-43d0-a2cc-8cc7c2e8b5c1","year":2023},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.247874Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:59d11b95bf15a407506eb4bcb1ad15c2c7e470b5d535e49abebf9ebfbf4e927e","observation_id":"14b80c33-7512-460f-8752-6291cfaf5df3","resolution":{"observed_at":"2026-08-07T13:20:30.827996Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:30.610608Z","title":"Highlighting object category immunity for the generalization of human -object interaction detection ,","venue":null,"work_id":"f8739cb2-899c-404d-8bdf-50f2b51f7c4f","year":2022},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.355426Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:ffd4e6f9ef12598b939484cf3a527a53e7988406a70f7468329a7887cb81c8fa","observation_id":"870bfc00-c1d9-4458-b7b7-39d696028136","resolution":{"observed_at":"2026-08-07T13:20:30.666037Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:30.440489Z","title":"Chairs can be stood on: Overcoming object bias in human -object interaction detection ,","venue":null,"work_id":"823f94f7-6942-4d64-8373-f3c1d2740556","year":2022},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.450488Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:73f8e1e7dd33bb96904d1d057553c4a3be1624b6f34ad0552f210279606c2729","observation_id":"5f36cd03-dfd0-4863-af7e-e565dbb7fb57","resolution":{"observed_at":"2026-08-07T13:20:30.491259Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:30.314778Z","title":"Prototype rectification for few -shot learning,","venue":null,"work_id":"a7f8167d-c4b4-4e5d-8785-fa50adcb57df","year":2020},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.547487Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:b9ea658dc690c0c39fe676ba121e823abaca9857e18e9ea4782584078a4a90ea","observation_id":"b4a5155c-da9c-4c91-9003-782ec6fcb58e","resolution":{"observed_at":"2026-08-07T13:20:30.358995Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:30.172225Z","title":"Prototype contrastive learning for point - supervised temporal action detection ,","venue":null,"work_id":"93730e3d-7a65-45f1-945b-bea490740774","year":2023},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.627641Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:e79a9eb2bb5a9ee51b38d9002ebb26690dff66a346ff79b0b9e642271c58a6b5","observation_id":"be686092-f57e-4834-bd7b-7e7fb8103b5e","resolution":{"observed_at":"2026-08-07T13:20:30.255093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:29.990109Z","title":"Ultralytics YOLOv5,","venue":null,"work_id":"a58e7e67-f7ec-4033-9de7-b90549347210","year":null},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.719614Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:5003ddf0e6f1c50614110833bde74ff60486400eae8b88d8c386601ccda10bb9","observation_id":"ea92ff6b-84c9-4ec8-9365-f2aea645977a","resolution":{"observed_at":"2026-08-07T13:20:30.103007Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:29.798720Z","title":"Simple online and realtime tracking with a deep association metric,","venue":null,"work_id":"479ce3ce-fb9c-4810-a6d4-4968607ffb01","year":2017},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:28.023288Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:ff496ad7960552422b6c649ea93a62f5ccc3436f97a0d965bd8f8ae20ac29f52","observation_id":"656613bc-f6e5-4809-8086-0619eab6d668","resolution":{"observed_at":"2026-08-07T13:20:29.900228Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:29.578320Z","title":"GloVe: Global vectors for word representation ,","venue":null,"work_id":"90d95af8-72d4-4c70-a063-45a67a505023","year":2014},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:28.144092Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:356be2a46a3f7a44fcdc98e5e5869699feb6ddd80b30080db03415b54620615c","observation_id":"1c8e8e0d-03bc-40ff-bfd8-7e071e877089","resolution":{"observed_at":"2026-08-07T13:20:29.652237Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:29.317693Z","title":"Extreme multi-label loss functions for recommendation, tagging, ranking & other missing label applications,","venue":null,"work_id":"52228aa0-2de7-4a2a-8bf3-41e305425855","year":2016},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:28.293609Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:316797a7685d7450c692d8eebb36b382c82343dc1ca5c6d693085002e319f529","observation_id":"046e3ce4-caba-49e9-880b-4263bc70a78b","resolution":{"observed_at":"2026-08-07T13:20:29.459397Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:29.180218Z","title":"Ramakrishnan G. Skew -robust human-object interactions in videos ,","venue":null,"work_id":"3ce9b99a-1a9b-4501-a929-8e045edba7ae","year":2023},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:28.377649Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:4b15947903a32a3778f97551af20a9c299783ffc6175d4d310d2fd28fd46ec63","observation_id":"10018902-baf6-4efa-ad40-0fcc9de412d5","resolution":{"observed_at":"2026-08-07T13:20:29.223791Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.00123","last_updated":"2018-04-30T22:49:54Z","snapshot_observed_at":"2026-07-06T06:36:33.930118Z","submitted_at":"2018-04-30T22:49:54Z","title":"CrowdHuman: A Benchmark for Detecting Human in a Crowd","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.00123","snapshot_observed_at":"2026-08-07T13:20:28.464764Z","title":"CrowdHuman: A benchmark for detecting human in a crowd,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:28.464764Z"},"links":{"cited_paper":"/paper/1805.00123","citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:bd7a28f13b80fa85b857f59f08982d7259efcd29b95256e6bdc8601b0660579f","observation_id":"87500ef7-2479-460a-babb-94fd67aa2708","resolution":{"observed_at":"2026-08-07T13:20:28.464764Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:28.975957Z","title":"Detecting attended visual targets in video ,","venue":null,"work_id":"e7b4fc2a-afb0-4838-8a7b-fd0533ac3c95","year":2020},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:28.521551Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:a2fa9a5917dfe9e7dfed12df9d157d46614a6ca59da7d530493bb97dad55246d","observation_id":"be063bc4-b1eb-4563-90b8-ccb572233579","resolution":{"observed_at":"2026-08-07T13:20:29.129815Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:28.798569Z","title":"Open Set Video HOI detection from Action-centric Chain -of-Look Prompting ,","venue":null,"work_id":"e72b7d13-fb3d-432e-9ca8-419b78f934a8","year":2024},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:28.654352Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:ed6f9d3912720ac7422d27be8391b9b6f03c0febe1f218400e77dbf6a1e48342","observation_id":"b814cb55-9b10-40fe-b41f-1608c3ae7f9a","resolution":{"observed_at":"2026-08-07T13:20:28.861859Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:20:27.891533Z","title":"Available: https://doi.org/10.5281/zenodo.3908559","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming","version":2},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:27.891533Z"},"links":{"citing_paper":"/paper/2505.22011"},"observation_digest":"sha256:c97be5c50ffe59b2e07727f10e0ec836504f1cd322a53e3be07f9d38826bf237","observation_id":"1360053a-d41a-495a-923b-f0de53783769","resolution":{"observed_at":"2026-08-07T13:20:27.891533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.22011","last_updated":"2025-08-04T14:48:23Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-08T18:37:42.815734Z","submitted_at":"2025-05-28T06:19:37Z","title":"Prototype Embedding Optimization for Human-Object Interaction Detection in Livestreaming"},"reference_resolution":{"displayed":23,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":2,"verified_exact":0,"verified_fuzzy":21},"total_outbound_references":23},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 23 of 23 outbound references and 0 inbound Pith citation observations for arXiv:2505.22011."}