{"as_of":"2026-08-06T20:23:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5b5c7d0818b9a8e74520b6d817bd917f1ff8fe74df78f6f5414a91df6ae6451d","coverage":[{"denominator":77,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":77,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T13:40:26.545640Z","state":"measured"},{"denominator":77,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":77,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2604.13581/citation-record","integrity":"/paper/2604.13581/integrity","json":"/paper/2604.13581/citation-record.json","paper":"/paper/2604.13581"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":"2303.08774","doi":"10.1002/tea.20265","metadata_source":"pith","pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GPT-4 Technical Report","venue":"cs.CL","work_id":"b928e041-6991-4c08-8c81-0359e4097c7b","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:1f5bb9e2280a4ea9f64edfe24b5c459310460c83e2a416ff90cfb89ec7aba2a5","observation_id":"8d1d1d6e-93c2-4c81-9783-401fd5aedf2f","resolution":{"observed_at":"2026-05-10T13:45:28.713183Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":"2502.13923","doi":"10.48550/arxiv.2502.13923","metadata_source":"pith","pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2.5-VL Technical Report","venue":"cs.CV","work_id":"69dffacb-bfe8-442d-be86-48624c60426f","year":2025},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:9c8256804d587f4f84fd85e414bd410bbe9864d8e6d251e6c7b80847f186a399","observation_id":"f68afbbe-6aa2-4d86-b2f0-48961a960021","resolution":{"observed_at":"2026-05-10T13:45:28.717111Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:19:13.082554+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:19:13.082554+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Multi-hmr: Multi-person whole-body hu- man mesh recovery in a single shot","venue":null,"work_id":"c0626ece-8203-4020-aa46-d97ba9df6d37","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:1134e2b590f9601be2d4cd4f6d23085eb1cda94028f1bbfaebd5521c279086dd","observation_id":"027a544d-dd9f-4b14-8a71-e10c1117650f","resolution":{"observed_at":"2026-05-18T22:52:53.312533Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Keep it smpl: Automatic estimation of 3d human pose and shape from a sin- gle image","venue":null,"work_id":"40958131-314c-431f-8cea-b270579239df","year":2016},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:56491d7bb14a58951680c3903428c2cd5ddc2849155cf7b1276d4e057d291d06","observation_id":"99c7a567-b929-4a2c-81d0-9b799a322d4b","resolution":{"observed_at":"2026-05-18T22:52:53.298657Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Executing your commands via motion diffusion in latent space","venue":null,"work_id":"517e19c1-a054-4e0b-b392-1e7ca48ccac8","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:eb8c0749deab3228c2b5b3af88b561f7adfecbc9d8dec6ae31e79e2c5284b014","observation_id":"8af1b78a-b3ad-421b-a240-59448d9704b2","resolution":{"observed_at":"2026-05-18T22:52:53.337394Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"991216d8-ba38-43e6-9ec4-083bf1554320","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:fd335de781d7b2b1780b84cce56caefe77b75aa0c8e1eeb44d2b020f324245fa","observation_id":"9399b534-5bf0-4049-8df8-b83d59b46d69","resolution":{"observed_at":"2026-05-18T22:52:53.334564Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2108.02938","last_updated":"2021-09-15T04:03:34Z","snapshot_observed_at":"2026-07-06T11:36:03.810011Z","submitted_at":"2021-08-06T04:43:13Z","title":"ILVR: Conditioning Method for Denoising Diffusion Probabilistic Models","version":2},"cited_work":{"arxiv_id":"2108.02938","doi":"10.48550/arxiv.2108.02938","metadata_source":"arxiv_reference","pith_arxiv_id":"2108.02938","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ilvr: Conditioning method for denoising diffusion probabilistic models","venue":"arXiv (Cornell University)","work_id":"cfb1b920-fb4e-4617-b335-ddfe3d6f7f84","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2108.02938","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:09663bf391cb283761cd65482af81836510fd6f4e1bd7aa5632c1f878f6db8b2","observation_id":"be0fbf85-80de-42db-94b7-cb5e6a3e6b51","resolution":{"observed_at":"2026-05-10T13:45:28.721750Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Interaction transformer for human reaction generation.IEEE Transactions on Multimedia, 25: 8842–8854","venue":null,"work_id":"9fea5ac7-0cb6-4520-b004-d870ecff4cf2","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:b93cfe1ff4813e236faa7e6f71c019be86c8e41402f979a05475d9d51a80acf2","observation_id":"6b2cc349-84b5-4565-ae7d-5e000cc18717","resolution":{"observed_at":"2026-05-18T22:52:53.252337Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Huang, Siyu Tang, Dimitris Tzionas, and Michael J","venue":null,"work_id":"ef8902d1-ecd1-47f4-b66b-a052fd214a44","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:2f8784c2a487685edf16a02b1f38dac7d7a39a7a829cfb5c99208b1bbae3ff86","observation_id":"1fcc9f00-0a99-46df-a0d1-46659e5fecb1","resolution":{"observed_at":"2026-05-18T22:52:53.248963Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Improving diffusion models for inverse problems using manifold constraints.Advances in Neural Information Pro- cessing Systems, 35:25683–25696","venue":null,"work_id":"fec136f0-22ec-4ba7-943b-980d4719ad12","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:9c798d8e36414ba244c2262f40d8357f1ee4d1cf1490942070526f75ec5809e8","observation_id":"5bdf9610-c1c6-4bb9-9964-112f9ecd6d01","resolution":{"observed_at":"2026-05-18T22:52:53.265210Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Capturing closely interacted two-person motions with reaction priors","venue":null,"work_id":"5de95ad5-f69a-4bf4-a469-8ad342108480","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:edcb31b06b62f7dce13fc79b9fc3b74bb95856f2fae42c07b4ad00699432f0e6","observation_id":"f52c2d8f-de2c-4a81-9fdb-e515b4723e23","resolution":{"observed_at":"2026-05-18T22:52:53.399045Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Diffpose: Spatiotemporal diffusion model for video-based human pose estimation","venue":null,"work_id":"b69f796a-ff18-4765-b636-c1eca346aa02","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:82e1d8b4e1120e2bd9c9f5de2f98203be6e93fa1d35259005dc886390da8388b","observation_id":"4a0e78ce-6470-4d1d-9720-8e9dcfa387d2","resolution":{"observed_at":"2026-05-18T22:52:53.242545Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Three- dimensional reconstruction of human interactions","venue":null,"work_id":"4bfcb163-e142-46fe-871e-a2d09403070f","year":2020},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:d7b0745698b3cbfe7bcecc39a87708d04bab4aaf6df72712c2a0a136a1be3422","observation_id":"ed59af49-6d42-4d85-88b2-d6d7a4433753","resolution":{"observed_at":"2026-05-18T22:52:53.252748Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Remips: Physically consistent 3d reconstruction of multiple interacting people under weak supervision.Advances in Neural Information Processing Systems, 34:19385–19397","venue":null,"work_id":"13cf5077-baa4-4b5d-a4c2-7bedc830a597","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:83b66be5f1b0bb555b960d2dc4f13bf29bb6182b644de8b9eabb699f4dd38992","observation_id":"085f1491-1cdb-4709-bcec-e5447295ce69","resolution":{"observed_at":"2026-05-18T22:52:53.245773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The potential of human pose estimation for motion capture in sports: a validation study","venue":null,"work_id":"46e53d02-a69c-49a8-92db-ee33ca8ab7ca","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:9bbbde406441ab39cc58ec4f07d7f36c2ece469f32bac4c4c7e3bb2407d43dd7","observation_id":"195a239f-1555-4235-a877-8b9fb5470a89","resolution":{"observed_at":"2026-05-18T22:52:53.274385Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.08430","last_updated":"2021-08-06T03:22:14Z","snapshot_observed_at":"2026-07-06T11:30:06.143581Z","submitted_at":"2021-07-18T12:55:11Z","title":"YOLOX: Exceeding YOLO Series in 2021","version":2},"cited_work":{"arxiv_id":"2107.08430","doi":null,"metadata_source":"pith","pith_arxiv_id":"2107.08430","snapshot_observed_at":"2026-07-04T19:20:06.601231Z","title":"YOLOX: Exceeding YOLO Series in 2021","venue":"cs.CV","work_id":"112b3cd9-8fe6-49fe-bbaa-90a3f46045c7","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2107.08430","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:de9fbc93f74b32efe48efaeb8db7b72d9359d536a0413b6a71d47bb595d65338","observation_id":"b11c6fa1-46c9-4f57-98f4-bdbc24fad827","resolution":{"observed_at":"2026-05-13T10:31:31.716715Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Remos: 3d motion- conditioned reaction synthesis for two-person interactions","venue":null,"work_id":"2cd5d5ce-929f-413e-8ccf-6470da0ecc07","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:887291f249938b80ffa9dc20ebe5db13a2f56178d5a378cd5b74aaf1ae9cb99e","observation_id":"231ad2ff-c78f-4b64-aa1a-52800ffbe519","resolution":{"observed_at":"2026-05-18T22:52:53.256187Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Humans in 4d: Recon- structing and tracking humans with transformers","venue":null,"work_id":"e2bc5796-4afc-4efb-917f-6fbe36f9d273","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:ac55e2b9ba731cdf48de82fb736a63b10e2372d0fe5fbe11575063b08b579d6d","observation_id":"38bd7a86-da8c-4ee0-9cdb-24b21279af68","resolution":{"observed_at":"2026-05-18T22:52:53.275257Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Reconstructing groups of people with hypergraph relational reasoning.2023 IEEE/CVF International Conference on Computer Vision (ICCV), pages 14827–14837","venue":null,"work_id":"2e7879de-a6d9-4966-b14f-469eb1a6c93f","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:7c41367eae0714a724168ee7817c4d470b901671a83820d87380d975950f6bc6","observation_id":"9ef78a10-c21b-4a2c-8fa1-53c940ef19e4","resolution":{"observed_at":"2026-05-18T22:52:53.302307Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Closely interactive human reconstruction with proxemics and physics-guided adaption","venue":null,"work_id":"ee9d5dd5-c463-4abd-bb28-60c5dc78fcda","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:642a1d8323789c946d17c78601643e2e5cbea00031cb3c6c213bfe02d50da72e","observation_id":"ad06032d-f147-4ecd-bf6f-4f33a1ae9da0","resolution":{"observed_at":"2026-05-18T22:52:53.201592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Motiongpt: Human motion as a foreign language","venue":null,"work_id":"cb115770-161d-4570-8e59-25915826dc64","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:21b8fb9b3213557d4e3212b77e377c4c5288b2ba45400e080e379de4e01a1018","observation_id":"1cc66f67-1d1d-4a84-91ea-836ae669d962","resolution":{"observed_at":"2026-05-18T22:52:53.305833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Coherent reconstruction of multiple humans from a single image","venue":null,"work_id":"b394d07c-950b-43d1-a749-6e6b85ef65bd","year":2020},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:4cdfed0ed443c714f6d52aca2b6f45c4408ae74e0ef7ca38489d2bec296c0909","observation_id":"53bbd18a-63a5-4ef6-a1ec-df1f83f6f189","resolution":{"observed_at":"2026-05-18T22:52:53.278429Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ultralytics yolov8","venue":null,"work_id":"28670f33-662c-48a7-a262-6d24e74db2ed","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:d3bf141315d087676580d05bd7036b60465b726fbf60d9eabf3557ab625956a7","observation_id":"066b7571-312b-4570-89a4-a5b042201a16","resolution":{"observed_at":"2026-05-18T22:52:53.262395Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"End-to-end recovery of human shape and pose","venue":null,"work_id":"c0fdc382-45f0-406f-8526-6288b5cf521f","year":2018},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:8eec004aa16027be26c6fdecf6a3d88d8270731f0955adb9faedc81866b41d3b","observation_id":"9542a235-ce89-420d-a0ea-5810ffb7a09e","resolution":{"observed_at":"2026-05-18T22:52:53.198411Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Maximizing parallelism in the construction of bvhs, octrees, and k-d trees","venue":null,"work_id":"76f24150-e3da-408f-854d-c5738a31a625","year":2012},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:8df9f7820d1f0a762009dae91208942f837bfd6c759d0db93b48a8f4f0831aa5","observation_id":"b12592ca-acd1-47b4-9133-6583a036e339","resolution":{"observed_at":"2026-05-18T22:52:53.218802Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Harmony4d: A video dataset for in-the-wild close human interactions.Advances in Neural Information Processing Systems, 37:107270–107285","venue":null,"work_id":"252477e7-b15f-46c2-a0e0-1872e3d7e427","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:82590c0148e863cec25f6855ea0680cb0958506d5558e33356bc76944708a718","observation_id":"d373bf78-d032-4f09-a33f-6e18637413b6","resolution":{"observed_at":"2026-05-18T22:52:53.262199Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Vibe: Video inference for human body pose and shape es- timation","venue":null,"work_id":"2226a6e1-801c-4637-b095-03722f8f4eb8","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:f56fd7a828bfe320b37eb8e80dea85c8975bb1bf11882f0650eeafaeec2c5516","observation_id":"1be506f4-b1e7-48af-b1a3-875b4e58eceb","resolution":{"observed_at":"2026-05-18T22:52:53.327774Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pare: Part attention regressor for 3d human body estimation","venue":null,"work_id":"11eb6adc-0e2e-4162-8f7f-4e38ac6587ee","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:eaaafdb1517eb0bee010bfdbc6e00cfd89a905a637a2ffb635be867c6b01df8a","observation_id":"397ebd9f-6868-4ce0-8471-43b3f5c6f621","resolution":{"observed_at":"2026-05-18T22:52:53.408926Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Hybrik: A hybrid analytical-neural inverse kinematics solution for 3d human pose and shape estimation","venue":null,"work_id":"950e5020-b88e-4b52-b7db-fbcbadbedcd0","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:6361d3495e3441ed386d610e8d9581f32e7ac44a0a1c1490a5fdf4a85b143c18","observation_id":"5d222bc6-8106-4185-be12-148b47ad5986","resolution":{"observed_at":"2026-05-18T22:52:53.388334Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Cliff: Carrying location information in full frames into human pose and shape estimation","venue":null,"work_id":"b2db85c6-69ad-4442-969e-dbf448998a7b","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:4aafbd4380098272babcf5adff3d062942dd07aa80c5cf5fc70abaf8af818feb","observation_id":"919d9be6-2563-45cc-ad83-7859a9afb4b7","resolution":{"observed_at":"2026-05-18T22:52:53.374975Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Intergen: Diffusion-based multi-human motion gener- ation under complex interactions.International Journal of Computer Vision, 132(9):3463–3483","venue":null,"work_id":"be912c09-566b-4eb3-a25b-99b4713c8f41","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:9aac6f12d186f6fc4e5cabc596dfbfe8056bffd21601a7dc5abdc1f70d361c8e","observation_id":"1c450a65-936d-49dd-9f9d-f3c5ad1a672c","resolution":{"observed_at":"2026-05-18T22:52:53.384967Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Motions as queries: One-stage multi- person holistic human motion capture","venue":null,"work_id":"8f4842a2-8aa4-402e-8e18-1e80e39f44b4","year":2025},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:ba5844860f7d3d966b1f0ec17892cb37fc074c43e08093fbff6a4c282ac04fee","observation_id":"d878b0cd-3039-43b3-995c-6b56990b421b","resolution":{"observed_at":"2026-05-18T22:52:53.378604Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.08983","last_updated":"2024-02-05T13:39:03Z","snapshot_observed_at":"2026-07-06T17:01:44.940193Z","submitted_at":"2023-12-14T14:29:52Z","title":"Interactive Humanoid: Online Full-Body Motion Reaction Synthesis with Social Affordance Canonicalization and Forecasting","version":3},"cited_work":{"arxiv_id":"2312.08983","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.08983","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2312.08983 (2023) 3","venue":null,"work_id":"334ebcc1-e711-441c-aeaf-ca0e752054bd","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2312.08983","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:2af1e870dab78638cbbf13f2ed83ab9abb0c3548bc9458cb31670ef15ebd6ab4","observation_id":"de078a7b-efad-4bd0-8f8b-0247e94272dd","resolution":{"observed_at":"2026-05-10T13:45:28.743951Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.05541","last_updated":"2024-03-23T04:54:21Z","snapshot_observed_at":"2026-07-06T16:59:11.649584Z","submitted_at":"2023-12-09T11:18:45Z","title":"DPoser: Diffusion Model as Robust 3D Human Pose Prior","version":2},"cited_work":{"arxiv_id":"2312.05541","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.05541","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dposer: Diffusion model as robust 3d human pose prior.arXiv preprint arXiv:2312.05541","venue":null,"work_id":"78f77c43-ff5d-447e-8720-973b750e76c0","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2312.05541","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:5b01170a767d207701085d3f776f47beff567cb3094f4be441364995a47d5c99","observation_id":"3d731841-a0d4-48a7-b325-e4c4b430a2d2","resolution":{"observed_at":"2026-05-10T13:45:28.731725Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Autotrackanything","venue":null,"work_id":"6640fd66-9743-4c2a-b0f6-9c8f6138eefe","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:fbf0802f4e32458a4807de938acbb6da7f0a774917164872ae02e749dd882830","observation_id":"59d7e6dc-177a-46a4-95dd-fc530fabb610","resolution":{"observed_at":"2026-05-18T22:52:53.394391Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Generative proxemics: A prior for 3d social interaction from images","venue":null,"work_id":"81dcdbc5-89f5-45ee-9735-2fea086e5366","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:db2d2906514b91d539f8cf36620d260ac989f65cfa56ccfb48018c42899b8fc5","observation_id":"24a0c873-b6d2-4446-bce2-e6493e7ff415","resolution":{"observed_at":"2026-05-18T22:52:53.366798Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Richter, and Vladlen Koltun","venue":null,"work_id":"f431da5e-70d2-4fd3-8419-f9f2df6994de","year":2025},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:94bb74990be28718f6eaf2927af10fd8219058cd65d3aae327d5cc47cc6d1ccf","observation_id":"15595afc-2073-4a69-a174-bdae940fb6a4","resolution":{"observed_at":"2026-05-18T22:52:53.381815Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T03:36:44.476283Z","title":"Improved denoising diffusion probabilistic models","venue":null,"work_id":"32db0798-4e73-487b-9a8b-499b596bfd8c","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:e6558abfc5673c4bb4516c13a5eeb0be0f1eae110fb226b4338ce012b0211527","observation_id":"b3a0bc52-f2a9-40b4-9a4a-8567ff6ba741","resolution":{"observed_at":"2026-05-18T22:52:53.321340Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Expressive body capture: 3d hands, face, and body from a single image","venue":null,"work_id":"ec9895b4-4395-4e62-bfa9-49381ec9c25b","year":2019},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:e04c76fa8a6aeef03a319170b514f4d9459e89456a7cb4f04312ddbc31417699","observation_id":"7f1fe315-32d5-485e-8684-ac318ebd56f2","resolution":{"observed_at":"2026-05-18T22:52:53.315815Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learning transferable visual models from natural language supervi- sion","venue":null,"work_id":"25b9941d-2ba0-4162-9b2a-57cd4e691ef5","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:854a567569777ba856f28f22d0230fc9724953c40f8608153ab73d4773eb3192","observation_id":"15ff25fb-9d3d-4edc-860c-80c7707d1d88","resolution":{"observed_at":"2026-05-18T22:52:53.330900Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T21:20:09.034420Z","title":"Humor: 3d human motion model for robust pose estimation","venue":null,"work_id":"d11205ff-50f6-4b84-959d-a6fad64b841e","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:8b4382d33db45e536397cc8140bbcb08a88176b1e0be354371555c1ce3713196","observation_id":"cc44dc07-9303-422d-b909-22505dfb46ff","resolution":{"observed_at":"2026-05-18T22:52:53.344979Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Diffhpe: Robust, coherent 3d human pose lifting with diffu- sion","venue":null,"work_id":"b5a78a62-6dd7-48ec-b21c-ffd9d8195f32","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:de16416ee1da8153f940d80d567873bdf74c20ce88a524269d04f9fa1650cca1","observation_id":"e74f3ca2-216f-4f75-a0ae-20c0083f41ed","resolution":{"observed_at":"2026-05-18T22:52:53.341979Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Phasemp: Robust 3d pose estimation via phase-conditioned human motion prior","venue":null,"work_id":"8b7e97da-5d10-4e19-bbd6-aa93754d1779","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:2d121afc10acba2faf3b0e5e884ea367147c4183d9b0ffe9d0cbf957e129f9f8","observation_id":"0ca26665-9ee5-43ff-a839-56739bd0bf57","resolution":{"observed_at":"2026-05-18T22:52:53.283581Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Sat-hmr: Real-time multi-person 3d mesh estimation via scale-adaptive tokens","venue":null,"work_id":"632ef519-1dbf-4b1f-b0c1-b9a1adf95b04","year":2025},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:7e76db50d3b3b05d3e1b9e4f8c0c95420d505b6ae3b1acfed8c309d5b676ae09","observation_id":"6788fb61-8a43-4dcf-8c68-7938be0ab16d","resolution":{"observed_at":"2026-05-18T22:52:53.291371Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.03689","last_updated":"2025-05-15T14:35:43Z","snapshot_observed_at":"2026-07-06T18:10:33.673157Z","submitted_at":"2024-05-06T17:59:36Z","title":"Pose Priors from Language Models","version":2},"cited_work":{"arxiv_id":"2405.03689","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.03689","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pose priors from language models","venue":null,"work_id":"49f90447-7dcd-4712-bd89-4abfa9426b2f","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2405.03689","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:27c7d17e77495bfd84b7c18959950a7506a0cc3375f5c00332f818dba5f4b864","observation_id":"53c8238d-916f-4de6-8aa2-643dc52eefab","resolution":{"observed_at":"2026-05-10T13:45:28.758533Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Monocular, one-stage, regression of multiple 3d people","venue":null,"work_id":"0328ecc2-852a-466e-8312-24bbd5998a6f","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:b192b4f859dd0169970ac35cc74c84303637c1fc0a46d2b65d3d1d3674a61d68","observation_id":"2e474906-302e-436e-8cec-d74050c0d340","resolution":{"observed_at":"2026-05-18T22:52:53.301934Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Putting people in their place: Monocular regression of 3d people in depth","venue":null,"work_id":"473f4fc1-2336-4fe3-b3b4-df795d13f75c","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:e21e6cda6eadea41b356a0ff6a28a27e24470cf23307a8955afd916ac68cf8b8","observation_id":"c85e3fff-52b2-4734-8b94-c46edfbc78ce","resolution":{"observed_at":"2026-05-18T22:52:53.280900Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Role-aware interac- tion generation from textual description","venue":null,"work_id":"cbddfcdc-4812-40fa-a8ee-6b2e0634edd3","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:0c47db790123d3513aa4146d1bd9fc7021359f51cff67b7d695008e4120e5021","observation_id":"d9f39284-2b2f-42ec-b981-631d1017fbb7","resolution":{"observed_at":"2026-05-18T22:52:53.363506Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Human motion diffusion model","venue":null,"work_id":"22aa6daa-ca67-40a3-ae85-07373b283cb9","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:28d78ddfddf45eaedef287b364ce74083146658c1f187d37e7ddf8e8106068ec","observation_id":"3da8615b-d35e-46c3-8fa6-8c8e04e0a868","resolution":{"observed_at":"2026-05-18T22:52:53.351674Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Multiphys: multi- person physics-aware 3d motion estimation","venue":null,"work_id":"a980e08d-c356-489b-9e11-00d87b5bd7ad","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:c9ac021eda63ed9113e5b20e43b4b70a5fcf5dfa61d8afdc75ff5f5cb9993f72","observation_id":"1948523b-56c3-4059-98e7-27095c7b15ed","resolution":{"observed_at":"2026-05-18T22:52:53.357131Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ai-based pose estimation of human operators in manu- facturing environments","venue":null,"work_id":"ec4484fe-f078-4537-838e-89573f30bca1","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:83050c5e016dd42b61bb1d468c83130de7929548e914c832bd840ecd3d55c7fd","observation_id":"5ac2ede6-9eca-4ac3-805c-338a62e230f7","resolution":{"observed_at":"2026-05-18T22:52:53.360206Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Recovering accurate 3d 10 human pose in the wild using imus and a moving camera","venue":null,"work_id":"c25469c9-2e9d-4b94-baa6-f34a2f24061a","year":2018},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:b3090f0baef910d6f2e93b626dbb4d2840a9e737af7d3b6eefba5a39356c9959","observation_id":"daa37d63-7b30-4b6b-a579-4221a8955e46","resolution":{"observed_at":"2026-05-18T22:52:53.268195Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Tlcontrol: Trajectory and language control for human motion synthesis","venue":null,"work_id":"2be09dff-62c7-4c90-846f-5e6a0da29509","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:cb0cc5c4cad7cd31eef6db54b7dc5dd9b576020e60cfa37e801cd592675bf376","observation_id":"7c454b44-dede-4217-9bfa-d99857998936","resolution":{"observed_at":"2026-05-18T22:52:53.271393Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Black, and Muhammed Kocabas","venue":null,"work_id":"e4d40b2d-cdd0-44e2-bf77-c0d8a4343033","year":2025},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:1acdd8a0ecbf92171facec06c63d7d5209a1a692e2a31ea6ed8f939ad333f391","observation_id":"83785512-473b-42ec-a49d-3317b11fc39c","resolution":{"observed_at":"2026-05-18T22:52:53.348079Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Inter- control: Generate human motion interactions by controlling every joint.CoRR","venue":null,"work_id":"3554d014-fcd5-4485-b75f-b2b38a451370","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:b176c31cd6795f90c6c257aeef10b490672f3574984e07409d229851f43de539","observation_id":"d58bee86-da79-4845-9633-4d850f28837e","resolution":{"observed_at":"2026-05-18T22:52:53.294935Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Crowd3d: Towards hundreds of peo- ple reconstruction from a single image","venue":null,"work_id":"be37ffad-b556-421e-af2c-9b006bbdbe86","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:b7571de204e661f6761bda6ade5ec6f9d43171003fe94c023d675227b4fe3677","observation_id":"a2122ad1-a533-429e-bc30-fd7dad5287b2","resolution":{"observed_at":"2026-05-18T22:52:53.309275Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"46b04cb2-2140-4501-b0a1-e7509a2f8195","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:4b0de179f57caac5d89b5f32554e030a0e7c8977305068f5503ebc20d297b76f","observation_id":"4bf87ddf-fdb0-4fbd-8b55-57afe73e0dde","resolution":{"observed_at":"2026-05-18T22:52:53.324587Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Occluded human pose estimation based on part-aware discrete diffusion priors.Knowledge-Based Systems, 315:113272","venue":null,"work_id":"f1494b90-504a-4427-8bc6-91d2de358015","year":2025},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:20c2c3d67d02cff3674524d65014a179ae177327b5a3a2161c778850fa03da26","observation_id":"c80ef291-b9cb-4055-9d53-76f31efc21dc","resolution":{"observed_at":"2026-05-18T22:52:53.391374Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08580","last_updated":"2024-04-14T22:23:18Z","snapshot_observed_at":"2026-07-06T16:32:07.186600Z","submitted_at":"2023-10-12T17:59:38Z","title":"OmniControl: Control Any Joint at Any Time for Human Motion Generation","version":2},"cited_work":{"arxiv_id":"2310.08580","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08580","snapshot_observed_at":"2026-07-09T22:06:35.163127Z","title":"Omnicontrol: Control any joint at any time for human motion generation","venue":"cs.CV","work_id":"b2dc1929-b877-4acf-ae34-659d6de94c52","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2310.08580","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:11fabd48c49bdd92e365fc18bbcb8635ad727c3b8104aa7e64243e9ea8944937","observation_id":"b0503eef-c9a8-4f88-bc9d-cb25b56b346d","resolution":{"observed_at":"2026-05-10T13:45:28.763411Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.03836","last_updated":"2025-02-06T07:42:00Z","snapshot_observed_at":"2026-07-06T20:32:00.585460Z","submitted_at":"2025-02-06T07:42:00Z","title":"Adapting Human Mesh Recovery with Vision-Language Feedback","version":1},"cited_work":{"arxiv_id":"2502.03836","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.03836","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Adapting human mesh recovery with vision-language feedback","venue":null,"work_id":"5267b5c9-78b0-43e0-aff2-4b0fcecc3202","year":2025},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2502.03836","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:005e287396bec3d2ebf697bb96ddd9c87ab12d4774b98c716232b89fe6a38674","observation_id":"97dfb901-8290-4d87-8418-aaf0e46ee057","resolution":{"observed_at":"2026-05-10T13:45:28.735611Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ghum & ghuml: Generative 3d human shape and articulated pose models","venue":null,"work_id":"a8889b95-a809-4aee-b6c3-94345096ef77","year":2020},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:4b36d824c8abb80447c1aef9c98ad27815228c6dd0609ee87d3aa2805c6b45df","observation_id":"68f2d89f-c9df-47de-ac8d-f69ec13edd37","resolution":{"observed_at":"2026-05-18T22:52:53.242822Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Regennet: Towards human action-reaction synthesis","venue":null,"work_id":"5443f298-c09d-4f07-b619-ca1f00779c5f","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:90cf9a50552299d38161f91f267ead78b5ba6b75dc5926bf230b86cfaf7b1809","observation_id":"af619c5b-c207-43a1-8904-7b4de1431c29","resolution":{"observed_at":"2026-05-18T22:52:53.259373Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Spatial tempo- ral graph convolutional networks for skeleton-based action recognition","venue":null,"work_id":"83692069-d9c7-4fbf-ab8c-eaef7e721526","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:875ead35b258e8b8e976c3712df5a74c28fea71ec54e0a718ec2359f64950e92","observation_id":"ac71c589-1814-4daa-9c07-c41c8dff2a7e","resolution":{"observed_at":"2026-05-18T22:52:53.236533Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Hi4d: 4d instance segmentation of close human interaction","venue":null,"work_id":"353c0050-aee5-4eed-a7aa-ed9e30df5443","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:77db3b36a5d47ec89da77a5fc8ad4183a31b082c006fdcd1a2e10374e03fe77f","observation_id":"bed42e8e-b8f9-4fe9-8f80-56d5307c1916","resolution":{"observed_at":"2026-05-18T22:52:53.255947Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04875","last_updated":"2018-07-12T07:55:09Z","snapshot_observed_at":"2026-08-02T03:55:49.327426Z","submitted_at":"2017-09-14T16:54:41Z","title":"Spatio-Temporal Graph Convolutional Networks: A Deep Learning Framework for Traffic Forecasting","version":4},"cited_work":{"arxiv_id":"1709.04875","doi":"10.3390/s16020157","metadata_source":"pith","pith_arxiv_id":"1709.04875","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Spatio-Temporal Graph Convolutional Networks: A Deep Learning Framework for Traffic Forecasting","venue":"cs.LG","work_id":"a51e6daa-daa3-4459-8ac0-e0016a30aec5","year":2017},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/1709.04875","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:e9b30f294e2a5728bb2751c9012022086c1c1c23a0157bb90eb459864605b93a","observation_id":"92596ee1-058e-47ba-a46f-9ac21993e8a3","resolution":{"observed_at":"2026-05-10T13:45:28.739912Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Glamr: Global occlusion-aware human mesh recovery with dynamic cameras","venue":null,"work_id":"f03054fc-91a5-4163-bfc1-d65b1d669544","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:0a41dcdba01223b039ec2680880252311850e381ed6a75a365d3aa814205478d","observation_id":"f95df2aa-29cf-4937-aedf-98934fae142c","resolution":{"observed_at":"2026-05-18T22:52:53.239582Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Monocular 3d pose and shape estimation of multiple people in natural scenes-the importance of multiple scene constraints","venue":null,"work_id":"f152b066-c633-46c3-8ece-007501b5c6c9","year":2018},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:e2a65546d677aad39fee29502ab57b185ca824308fb4499d5d77edcd67ede8b5","observation_id":"8772aa84-e8fd-4b30-97ac-93575caa6030","resolution":{"observed_at":"2026-05-18T22:52:53.208595Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Smoothnet: A plug-and-play network for refining human poses in videos","venue":null,"work_id":"71b7d490-2efe-4455-b4f2-ed0743ce96a1","year":2022},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:5478270cf6f4e75e8974062dd6d13894209dc874a5bc3867d388eb53909de872","observation_id":"50a9228d-539b-499b-b1a4-8341c8e775cb","resolution":{"observed_at":"2026-05-18T22:52:53.204869Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.14289","last_updated":"2023-07-01T07:26:22Z","snapshot_observed_at":"2026-07-06T15:46:29.060519Z","submitted_at":"2023-06-25T16:37:25Z","title":"Faster Segment Anything: Towards Lightweight SAM for Mobile Applications","version":2},"cited_work":{"arxiv_id":"2306.14289","doi":null,"metadata_source":"pith","pith_arxiv_id":"2306.14289","snapshot_observed_at":"2026-07-04T21:00:09.639501Z","title":"Faster Segment Anything: Towards Lightweight SAM for Mobile Applications","venue":"cs.CV","work_id":"d159afc6-0f47-4693-a3c0-908e661ff652","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2306.14289","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:091acd2d37ea5083e931836ef8f7a50038e2adb9d5860ac4cd0dfe5d3690db4e","observation_id":"fd7e8873-2c56-40b7-a0f1-7ab7e1a60286","resolution":{"observed_at":"2026-05-17T22:41:43.533652Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Adding conditional control to text-to-image diffusion models","venue":null,"work_id":"4dac20c4-1e6b-4024-aa92-e9e83038794f","year":2023},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:f9ff884350a57ae42fa1402ae3fccc193988bc190c6e9413c4f213a01d1032c8","observation_id":"6ccb73f4-e8d8-42a1-bb1a-3f4a7b840d47","resolution":{"observed_at":"2026-05-18T22:52:53.259131Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05260","last_updated":"2025-04-14T10:43:50Z","snapshot_observed_at":"2026-07-06T19:29:09.714725Z","submitted_at":"2024-10-07T17:58:22Z","title":"DartControl: A Diffusion-Based Autoregressive Motion Model for Real-Time Text-Driven Motion Control","version":3},"cited_work":{"arxiv_id":"2410.05260","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05260","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dartcon- trol: A diffusion-based autoregressive motion model for real-time text-driven motion control","venue":null,"work_id":"c6ea9a28-2e42-4120-9a88-7623f9511734","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"cited_paper":"/paper/2410.05260","citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:52d765ce0be956e8f4c5d2342dcf09a13853d2c4a6478c0b01bbc14918cf515e","observation_id":"f7b2b1a4-d85b-4057-b3bc-cfc4ff65e429","resolution":{"observed_at":"2026-05-10T13:45:28.749075Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"3d human pose estimation with spatial and temporal transformers","venue":null,"work_id":"dea28e02-fb55-442d-be2e-5bdc655dd94e","year":2021},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:12c4ee9bf29aa0d422cfda73299f04192d5e555683f9159ce8988e1b3c5e68df","observation_id":"3deb9596-ea89-435a-bf82-7f9022e11db5","resolution":{"observed_at":"2026-05-18T22:52:53.281424Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dpmesh: Exploiting diffusion prior for occluded human mesh recovery","venue":null,"work_id":"f1455d45-9ead-4fb0-b608-0004c7a00cff","year":2024},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:c9841a4a572d7c180fb96619866b9df540dddda93b7e58412c75752b34578419","observation_id":"3f8ba371-e660-4747-bce0-85ee3d7de9b0","resolution":{"observed_at":"2026-05-18T22:52:53.229380Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Implementation details Our model was implemented using PyTorch and trained on an NVIDIA RTX 3090 GPU","venue":null,"work_id":"2639a7bf-438a-400a-b444-8a69041e76bd","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:fd777bac973a7c7099ec04f6d9c0a92d5ae0a499daddb1cf8aa60e2db963078a","observation_id":"18763daa-a75c-412e-8755-11449c5f9000","resolution":{"observed_at":"2026-05-18T22:52:53.403191Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ablation on Geometry Optimizer Geometry Optimizer focuses on processing 3D joint posi- tions to provide geometric guidance information","venue":null,"work_id":"afafc80e-db99-4039-a4e7-bbc4e288542f","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:ce5325767db7672243e5931e9899e466623bff129d1cb00dc03a6bc5f9c9d0b1","observation_id":"1a085c2d-95df-43a2-a496-8342a6e73c3a","resolution":{"observed_at":"2026-05-18T22:52:53.191100Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"1973.8109","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Reconstruction under VLM Limitations Based on our user study, the text descriptions generated by the VLM are, on average, superior to those produced by human annotators","venue":null,"work_id":"bdf71017-ed46-4e4f-a9a5-5dd9553b7b7c","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:604c35fd1533d014c67df5ccc72b5dcada36530b3cfe4e4ecfb626776e77c352","observation_id":"765d0f5b-1a08-4ff7-86d6-a42b384589c0","resolution":{"observed_at":"2026-05-10T13:45:28.709131Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"LLMs are used only for light text polishing and grammar fixes","venue":null,"work_id":"5d2239f7-12ac-4577-8ce9-b40efd334c36","year":null},"citing_paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:26.545640Z"},"links":{"citing_paper":"/paper/2604.13581"},"observation_digest":"sha256:46173d7482ca61d5b8795af36a99b372a561c1e0e95721b0e5d01ad22a44e094","observation_id":"70a13032-811f-4141-8f61-f1a3f6008a45","resolution":{"observed_at":"2026-05-18T22:52:53.194682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2604.13581","last_updated":"2026-04-15T07:41:52Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-04T04:04:04.143303Z","submitted_at":"2026-04-15T07:41:52Z","title":"SocialMirror: Reconstructing 3D Human Interaction Behaviors from Monocular Videos with Semantic and Geometric Guidance"},"reference_resolution":{"displayed":77,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":2,"verified_exact":13,"verified_fuzzy":62},"total_outbound_references":77},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 77 of 77 outbound references and 0 inbound Pith citation observations for arXiv:2604.13581."}