{"as_of":"2026-08-08T01:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9dc87d22bdaacdf042902535ae212352e2ec1941a01a287e640f3bcfacdca01f","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":17,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":17,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":17,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":17,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:48:10.244869Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T03:39:29.096987Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-07T13:48:10.244869Z","title":"Motionllm: Understanding human behaviors from human motions and videos.arXiv preprint arXiv:2405.20340, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20920","last_updated":"2025-05-27T09:10:59Z","snapshot_observed_at":"2026-08-07T13:41:10.441862Z","submitted_at":"2025-05-27T09:10:59Z","title":"HuMoCon: Concept Discovery for Human Motion Understanding","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T13:48:10.244869Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2505.20920"},"observation_digest":"sha256:dad5d45aea367ad28c6952a8854b907ca8c0b1bb7bc6eb761da3133318122d0f","observation_id":"778a610e-bda4-495d-a79b-d180f7739efe","resolution":{"observed_at":"2026-08-07T13:48:10.244869Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-07T13:46:13.585941Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21146","last_updated":"2025-05-27T12:57:37Z","snapshot_observed_at":"2026-08-07T13:32:02.084262Z","submitted_at":"2025-05-27T12:57:37Z","title":"IKMo: Image-Keyframed Motion Generation with Trajectory-Pose Conditioned Motion Diffusion Model","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T13:46:13.585941Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2505.21146"},"observation_digest":"sha256:8f46c52a535b7db109016d6c0d10a4c94369e1106e8ae5942848e5de8e5617f8","observation_id":"05b0f365-19fe-4908-9250-23e6dc6569f5","resolution":{"observed_at":"2026-08-07T13:46:13.585941Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-07T12:06:27.907738Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.03191","last_updated":"2025-05-31T11:02:24Z","snapshot_observed_at":"2026-08-07T12:01:41.223719Z","submitted_at":"2025-05-31T11:02:24Z","title":"Multimodal Generative AI with Autoregressive LLMs for Human Motion Understanding and Generation: A Way Forward","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-07T12:06:27.907738Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2506.03191"},"observation_digest":"sha256:9a62c2ae872b4efa885eef96fe5a78930907ca2903d5ea04b9fbdb5f4866f547","observation_id":"8ce634ec-3c2a-48dc-938a-0d2ab47f9bc8","resolution":{"observed_at":"2026-08-07T12:06:27.907738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-06T20:29:48.674994Z","title":"Motion- llm: Understanding human behaviors from human motions and videos","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02591","last_updated":"2025-07-23T07:25:27Z","snapshot_observed_at":"2026-08-07T06:17:01.457380Z","submitted_at":"2025-07-03T12:55:16Z","title":"AuroraLong: Bringing RNNs Back to Efficient Open-Ended Video Understanding","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T20:29:48.674994Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2507.02591"},"observation_digest":"sha256:f7f125b5055fe3408143ce03f5cc2a990b05d37c4e14003d4512d1c565e366cd","observation_id":"54b94232-4936-4c5b-83ae-3360eb5bea48","resolution":{"observed_at":"2026-08-06T20:29:48.674994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-06T17:22:08.784420Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.11102","last_updated":"2025-07-15T08:52:28Z","snapshot_observed_at":"2026-08-06T17:14:03.669426Z","submitted_at":"2025-07-15T08:52:28Z","title":"KptLLM++: Towards Generic Keypoint Comprehension with Large Language Model","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T17:22:08.784420Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2507.11102"},"observation_digest":"sha256:dd9b46b4b7dde71236a3305b012346ae7dc3a62ed43af94d3833627104a9483d","observation_id":"41e14b63-2c7c-442a-a5df-bdb37a9e620a","resolution":{"observed_at":"2026-08-06T17:22:08.784420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-06T15:33:45.839617Z","title":"Motionllm: Understanding human behaviors from human motions and videos.arXiv preprint arXiv:2405.20340, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15597","last_updated":"2025-07-21T13:19:09Z","snapshot_observed_at":"2026-08-06T15:25:42.125132Z","submitted_at":"2025-07-21T13:19:09Z","title":"Being-H0: Vision-Language-Action Pretraining from Large-Scale Human Videos","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-06T15:33:45.839617Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2507.15597"},"observation_digest":"sha256:bdeee985d29c4dcc2ccc0fb25ac67c490baff70d4ce684ac29565b0f702143a1","observation_id":"f7cd73b6-a9aa-41d4-97c9-144cc1f960c5","resolution":{"observed_at":"2026-08-06T15:33:45.839617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-05T21:54:59.349922Z","title":"Motionllm: Understanding human behaviors from human motions and videos.arXiv preprint arXiv:2405.20340, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.07863","last_updated":"2025-08-11T11:26:10Z","snapshot_observed_at":"2026-08-06T19:31:41.610911Z","submitted_at":"2025-08-11T11:26:10Z","title":"Being-M0.5: A Real-Time Controllable Vision-Language-Motion Model","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T21:54:59.349922Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2508.07863"},"observation_digest":"sha256:6055d9516d0397a55401b15090f45e64a2dac99a32121377891fa3b12d90a0de","observation_id":"c0d45bf0-6888-4ddc-82c2-125238f37983","resolution":{"observed_at":"2026-08-05T21:54:59.349922Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-05T19:03:09.468537Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.13692","last_updated":"2025-08-19T09:52:04Z","snapshot_observed_at":"2026-08-08T00:06:34.475015Z","submitted_at":"2025-08-19T09:52:04Z","title":"HumanPCR: Probing MLLM Capabilities in Diverse Human-Centric Scenes","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-05T19:03:09.468537Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2508.13692"},"observation_digest":"sha256:989d99d8f5feef5c93def3b2a8db6ac793c712d04111621a3aeaab76861adebd","observation_id":"5a7c04f5-6443-4215-8b94-e45c166ae1ed","resolution":{"observed_at":"2026-08-05T19:03:09.468537Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-05T12:35:04.892673Z","title":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., Uszkoreit, J., and Houlsby, N","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2509.01471","last_updated":"2025-09-01T13:39:14Z","snapshot_observed_at":"2026-08-07T14:51:32.636999Z","submitted_at":"2025-09-01T13:39:14Z","title":"Hierarchical Motion Captioning Utilizing External Text Data Source","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-05T12:35:04.892673Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2509.01471"},"observation_digest":"sha256:840d82cc45401094d71021d2318ae44b3306cf5fbb71f27133d609fb4c4ae856","observation_id":"65a20522-797f-474f-bae5-05e15825bd9c","resolution":{"observed_at":"2026-08-05T12:35:04.892673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":"2405.20340","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-07-04T03:39:29.096987Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":"8fefb27e-4d05-4ce7-8de7-944336c1288d","year":2024},"citing_paper":{"arxiv_id":"2511.18373","last_updated":"2026-04-11T05:44:20Z","snapshot_observed_at":"2026-08-07T02:29:17.852730Z","submitted_at":"2025-11-23T09:43:44Z","title":"MASS: Motion-Aware Spatial-Temporal Grounding for Physics Reasoning and Comprehension in Vision-Language Models","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-17T05:55:11.495430Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2511.18373"},"observation_digest":"sha256:c79ff61adf23b27749a1484fd4399469340ed5fab07c1804146d8c15c510d2bd","observation_id":"05ee7cd5-197e-498a-9d01-c2db73514fbb","resolution":{"observed_at":"2026-05-17T05:59:08.668703Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-08-03T05:29:27.699878Z","title":"Motionllm: Understanding human behaviors from human motions and videos.arXiv preprint arXiv:2405.20340, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.02401","last_updated":"2026-07-07T06:51:43Z","snapshot_observed_at":"2026-08-03T05:29:25.861855Z","submitted_at":"2026-02-02T17:59:01Z","title":"Superman: Unifying Skeleton and Vision for Human Motion Perception and Generation","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-03T05:29:27.699878Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2602.02401"},"observation_digest":"sha256:329163062bc0257d8efd0f8e5cf43d1e11f688b607666595b5e9531f0a6ca930","observation_id":"a3a45041-6704-43b5-98d6-654eeb4c829f","resolution":{"observed_at":"2026-08-03T05:29:27.699878Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-07-13T20:17:37.389126Z","title":"arXiv preprint arXiv:2405.20340 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.22282","last_updated":"2026-06-27T20:51:08Z","snapshot_observed_at":"2026-07-13T20:17:36.181558Z","submitted_at":"2026-03-23T17:59:48Z","title":"UniMotion: A Unified Framework for Motion-Text-Vision Understanding and Generation","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-13T20:17:37.389126Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2603.22282"},"observation_digest":"sha256:495c8327802df90f9e11f1de015f5a2129ac6b3a705640ef50f82c936b72ee9b","observation_id":"ca3d3740-6eab-443e-9fd9-2dfdf2a49c07","resolution":{"observed_at":"2026-07-13T20:17:37.389126Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":"2405.20340","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-07-04T03:39:29.096987Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":"8fefb27e-4d05-4ce7-8de7-944336c1288d","year":2024},"citing_paper":{"arxiv_id":"2604.21926","last_updated":"2026-04-23T17:59:16Z","snapshot_observed_at":"2026-08-03T03:34:11.234757Z","submitted_at":"2026-04-23T17:59:16Z","title":"Seeing Without Eyes: 4D Human-Scene Understanding from Wearable IMUs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-09T21:59:00.442135Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2604.21926"},"observation_digest":"sha256:75f66af7feb27823ea16a038d0f26a5ce1aafc47724133bd5415cdef99d0055b","observation_id":"76407d53-ee9b-4346-81d1-5900af9b42db","resolution":{"observed_at":"2026-05-11T14:21:07.101462Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":"2405.20340","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-07-04T03:39:29.096987Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":"8fefb27e-4d05-4ce7-8de7-944336c1288d","year":2024},"citing_paper":{"arxiv_id":"2604.23264","last_updated":"2026-04-25T12:16:37Z","snapshot_observed_at":"2026-08-06T05:53:51.992744Z","submitted_at":"2026-04-25T12:16:37Z","title":"MotionHiFlow: Text-to-motion via hierarchical flow matching","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-08T08:28:42.524111Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2604.23264"},"observation_digest":"sha256:1484e25f1f012e69681700299d6b4588764dbefbd1c967c9c4849c845561359e","observation_id":"5735b75f-2e32-4b48-8fbb-c2bace7baf84","resolution":{"observed_at":"2026-05-11T20:36:10.922429Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":"2405.20340","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-07-04T03:39:29.096987Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":"8fefb27e-4d05-4ce7-8de7-944336c1288d","year":2024},"citing_paper":{"arxiv_id":"2605.14070","last_updated":"2026-05-13T19:47:07Z","snapshot_observed_at":"2026-07-06T23:25:34.835136Z","submitted_at":"2026-05-13T19:47:07Z","title":"WirelessSenseLLM: Zero-Shot Human Activity Understanding by Bridging Wireless Signals and Human Language","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T02:26:59.318141Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2605.14070"},"observation_digest":"sha256:ef416e0843ee60fab4ccd2181e94a161ac405ca186333354acfd7a0f5c260a8f","observation_id":"4d25b8dc-0f5e-4756-b49e-222cac66069f","resolution":{"observed_at":"2026-05-15T02:28:30.749765Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":"2405.20340","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-07-04T03:39:29.096987Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":"8fefb27e-4d05-4ce7-8de7-944336c1288d","year":2024},"citing_paper":{"arxiv_id":"2606.04773","last_updated":"2026-06-03T11:53:57Z","snapshot_observed_at":"2026-08-07T06:09:02.341844Z","submitted_at":"2026-06-03T11:53:57Z","title":"NextMotionQA: Benchmarking and Judging Human Motion Understanding with Vision-Language Models","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-06-28T06:55:22.332372Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2606.04773"},"observation_digest":"sha256:b78f27c1629986005659631bf3a298bd4aebee22b8822bd0737f80d31e9aea3b","observation_id":"b8cca54f-72ac-49c7-bfe2-839abf946976","resolution":{"observed_at":"2026-07-02T07:26:46.201459Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos","version":1},"cited_work":{"arxiv_id":"2405.20340","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.20340","snapshot_observed_at":"2026-07-04T03:39:29.096987Z","title":"Motionllm: Understanding human behaviors from human motions and videos","venue":null,"work_id":"8fefb27e-4d05-4ce7-8de7-944336c1288d","year":2024},"citing_paper":{"arxiv_id":"2606.20888","last_updated":"2026-06-18T19:31:56Z","snapshot_observed_at":"2026-08-04T06:10:35.180878Z","submitted_at":"2026-06-18T19:31:56Z","title":"Fine-grained Human Motion Understanding with Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-26T17:55:49.866744Z"},"links":{"cited_paper":"/paper/2405.20340","citing_paper":"/paper/2606.20888"},"observation_digest":"sha256:6590db597b3c43c60c72e746b5d97a68aafcd68b39d31e8a25a2e5a92f2a8d93","observation_id":"871a9618-42eb-4be6-9d14-4a500309b68d","resolution":{"observed_at":"2026-07-04T03:39:29.100697Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2405.20340/citation-record","integrity":"/paper/2405.20340/integrity","json":"/paper/2405.20340/citation-record.json","paper":"/paper/2405.20340"},"outbound":[],"paper":{"arxiv_id":"2405.20340","last_updated":"2024-05-30T17:59:50Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T18:22:57.400982Z","submitted_at":"2024-05-30T17:59:50Z","title":"MotionLLM: Understanding Human Behaviors from Human Motions and Videos"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 17 inbound Pith citation observations for arXiv:2405.20340."}