{"as_of":"2026-08-22T17:52:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e2a1ee7e52109c9e10d79a247d8102bfa51e6e3c25337be7dcd121257b6ec9d3","coverage":[{"denominator":151,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T12:14:46.798018Z","state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.29180/citation-record","integrity":"/paper/2607.29180/integrity","json":"/paper/2607.29180/citation-record.json","paper":"/paper/2607.29180"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.098016Z","title":"Generating Diverse and Natural","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.098016Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:b6e2edd0aa91900ee869621143a34b16d903293f4f338e254026d31bac425e72","observation_id":"7b438406-622e-4503-a925-c832dd437b40","resolution":{"observed_at":"2026-08-03T12:14:34.098016Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.16575","last_updated":"2025-07-08T20:15:01Z","snapshot_observed_at":"2026-08-19T11:18:14.567823Z","submitted_at":"2024-11-25T16:59:42Z","title":"Rethinking Diffusion for Text-Driven Human Motion Generation: Redundant Representations, Evaluation, and Masked Autoregression","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.16575","snapshot_observed_at":"2026-08-03T12:14:34.208362Z","title":"arXiv preprint arXiv:2411.16575 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.208362Z"},"links":{"cited_paper":"/paper/2411.16575","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d2831dc8ae10850ef66a276ed87f55592de8c509406df6ddb92b28e26d250dfb","observation_id":"f7878185-a5c8-4804-9b2f-4ce8eab0ecd2","resolution":{"observed_at":"2026-08-03T12:14:34.208362Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.286493Z","title":"International Conference on Learning Representations (ICLR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.286493Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:b4615968d41ae58841802912d05e55c7aa403a377d16ce7260c6b109ca70fa56","observation_id":"33034515-1f2d-4b0a-ae4a-7bb27c652348","resolution":{"observed_at":"2026-08-03T12:14:34.286493Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.368404Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.368404Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:13f82735b256ee728784d109ca63a715515722cdc99f629ee165ebdb21cb0472","observation_id":"157f4916-bfab-40cc-90dd-1aae1b99cf00","resolution":{"observed_at":"2026-08-03T12:14:34.368404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.470544Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.470544Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:cf8495f2ce16594cd014917d17b0ad1e07a8d526c32a8fd8eb165f3adf5dad9e","observation_id":"d15c0529-c260-42c9-b611-67d9671e3223","resolution":{"observed_at":"2026-08-03T12:14:34.470544Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.592007Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.592007Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:a5afa8ddb7e75c8b120b2e214d3cc5dc90801db86d704eff98b72ac5ba16a330","observation_id":"7e0d9a9e-2fcd-477b-9d1d-97aacb64a8c7","resolution":{"observed_at":"2026-08-03T12:14:34.592007Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.734473Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.734473Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:5f97b3cf3666b63e4f50ab29e2afd98fbcbd5a0d4731149fb85abcebad4d3e0e","observation_id":"37117316-0971-4226-800b-7125a6f1eb5b","resolution":{"observed_at":"2026-08-03T12:14:34.734473Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.820204Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.820204Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:283226f8bef2668a8961fabab2628b5fe217048922cbb721a905e8d1841f6191","observation_id":"53b4a98c-4d07-4f34-8750-2384492618f9","resolution":{"observed_at":"2026-08-03T12:14:34.820204Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:34.924736Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:34.924736Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:5a9c7f7de73f360ba2d5a98dc06f6a2270cd0ed0a000989cb1a63f7970529f68","observation_id":"b83b2b16-d7fb-4a4f-a40c-368159d9fa9e","resolution":{"observed_at":"2026-08-03T12:14:34.924736Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.009899Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.009899Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:4b2005a2ca28a679ebc239f22d0e4966cad539531eaf27dfacd5a31b463fc27c","observation_id":"a3ef4872-0867-46cc-92e7-38d656e00550","resolution":{"observed_at":"2026-08-03T12:14:35.009899Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.102784Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.102784Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:f47dd25315d0a4c4b9815922e74d886a846a39ec9ec3e262de1714b7efa14a90","observation_id":"662cae2a-6edf-4be2-ae67-8bab5f27a5ca","resolution":{"observed_at":"2026-08-03T12:14:35.102784Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.188936Z","title":"International Conference on Machine Learning (ICML) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.188936Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:c390a1361ceaf6458ade41947034ef8126e763a4475526374d5b6aac9fc081e8","observation_id":"886f9dfa-18a8-4544-b16b-ef3456121735","resolution":{"observed_at":"2026-08-03T12:14:35.188936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.306946Z","title":"International Conference on Learning Representations (ICLR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.306946Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d1905adf767c37aa16dbab24361afa9b5d878970311857f613f5b140a5e3268e","observation_id":"211f225e-ab74-4726-ba7f-4f570dca2513","resolution":{"observed_at":"2026-08-03T12:14:35.306946Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.488469Z","title":"International Conference on Learning Representations (ICLR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.488469Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:8a3fe238e39add49337b0ca377136095fda054879c82225c8b03de6d32cf19c7","observation_id":"c082ea95-bd4d-45e7-8dfa-c213bd79f49a","resolution":{"observed_at":"2026-08-03T12:14:35.488469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.600178Z","title":"and Boffi, Nicholas M","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.600178Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:aec64d1114c2df5242b5f93126c579b879cf1f14ae3e52d5eac51764670ca5f8","observation_id":"dcc3d041-60b3-4c48-875a-0e6db0085707","resolution":{"observed_at":"2026-08-03T12:14:35.600178Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.694314Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.694314Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:91b9ceec5d52e93728f205d11c466494fa93e125c30729d6884009a599b0e82c","observation_id":"d847ba6d-c77d-407a-afef-094dd9e37f8e","resolution":{"observed_at":"2026-08-03T12:14:35.694314Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.777342Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.777342Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:6f1e2cfcf3682d796f592b2f3686eb5361428acc42e85fe1df6771f37d03898f","observation_id":"86f8a31a-9602-460c-9a4a-9a7af696e7a8","resolution":{"observed_at":"2026-08-03T12:14:35.777342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:35.911058Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:35.911058Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:983ef82f85bc870c83b9ea2a2453488be5b9382cc943a87a4bfed0a4adb0ccd0","observation_id":"21b8e628-c141-4cbe-98b0-51fad66bda40","resolution":{"observed_at":"2026-08-03T12:14:35.911058Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:36.024637Z","title":"Transactions on Machine Learning Research (TMLR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.024637Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:db2817cf84a872cb510c95a8444545fc2534e2726ec21b4f0ab0d7f9c6f348c8","observation_id":"e76b8f1b-fd61-4db9-b289-d90dd69a1ea7","resolution":{"observed_at":"2026-08-03T12:14:36.024637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:36.159725Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.159725Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:58ddbca1eeb144386a11d3858c3963ce0197851a44e33a1620e3146b15d74172","observation_id":"dc3c5377-5adc-4680-a1d6-a51b255f7384","resolution":{"observed_at":"2026-08-03T12:14:36.159725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:36.293818Z","title":"Big data , volume=","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.293818Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:bebae544626b70152c64468a3c0b3f4d49252b6f9285c1e1969fc971cfc62d91","observation_id":"6284d003-53e7-4dbb-8dd4-9f9392e61a7a","resolution":{"observed_at":"2026-08-03T12:14:36.293818Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:36.408815Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.408815Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d2fa9b146902086a7d51957c2ce0d9f03fdab926106b7b43def4d51d77463eb3","observation_id":"82d19110-35c9-46d6-841c-6c750899e2bc","resolution":{"observed_at":"2026-08-03T12:14:36.408815Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:36.523249Z","title":"Biometrika , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.523249Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:ba6cb551c8fbcac4e078b2ce88315cda9bd435122583c2f65a2649f0dbaf3dc6","observation_id":"4bb607ac-42d9-48a9-a0a2-560c0f9e1a1e","resolution":{"observed_at":"2026-08-03T12:14:36.523249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1312.6114","last_updated":"2022-12-10T21:04:00Z","snapshot_observed_at":"2026-08-14T23:50:45.029465Z","submitted_at":"2013-12-20T20:58:10Z","title":"Auto-Encoding Variational Bayes","version":11},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.6114","snapshot_observed_at":"2026-08-03T12:14:36.641524Z","title":"arXiv preprint arXiv:1312.6114 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.641524Z"},"links":{"cited_paper":"/paper/1312.6114","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:19e963d02e4f881a223d5ca1a49216cd46ecd5e9f0c8e6c95c0e3a352c11dbd3","observation_id":"b131c2be-78fa-4566-a703-67cf5d4a1444","resolution":{"observed_at":"2026-08-03T12:14:36.641524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:36.787210Z","title":"International Journal of Computer Vision , pages=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.787210Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:f2cf135363ede5c8c78d6a93536eee1778599b441fcab1f676a98c2a01e0a5ca","observation_id":"065537f6-ca3c-4510-83dc-1add6de3845e","resolution":{"observed_at":"2026-08-03T12:14:36.787210Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:36.922135Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:36.922135Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:c58cc9745a460354d5d031e8227d1678fa6b1ba9a435a9c5b364d09280bb938c","observation_id":"918a08cc-0a6c-4947-948f-bf2c228950b4","resolution":{"observed_at":"2026-08-03T12:14:36.922135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:37.038841Z","title":"European Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.038841Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:1cc49a5170175c3c83918017c7f2b3cd879e54b61fa2b1c717bb1b9220e62140","observation_id":"2ea218bc-f90b-40a1-bb9e-6720dd7dc802","resolution":{"observed_at":"2026-08-03T12:14:37.038841Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:37.150150Z","title":"G lo V e: Global Vectors for Word Representation","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.150150Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:0f1fea6c467a9c8752f93093ecc3bb718143007aff177d60b85ab845476315ad","observation_id":"40eb7582-e671-49e0-8e1f-99ae63e162f1","resolution":{"observed_at":"2026-08-03T12:14:37.150150Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1412.3555","last_updated":"2014-12-11T06:46:53Z","snapshot_observed_at":"2026-08-20T00:41:09.240026Z","submitted_at":"2014-12-11T06:46:53Z","title":"Empirical Evaluation of Gated Recurrent Neural Networks on Sequence Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.3555","snapshot_observed_at":"2026-08-03T12:14:37.235246Z","title":"arXiv preprint arXiv:1412.3555 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.235246Z"},"links":{"cited_paper":"/paper/1412.3555","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:1557ea974ac2df6fee4c5714c297a65c3afe71b7b1e4e520a6dd04a71342cdea","observation_id":"5eb78de4-639f-4629-a1b8-78d1b8b47b35","resolution":{"observed_at":"2026-08-03T12:14:37.235246Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:37.304411Z","title":"Proceedings of the IEEE conference on computer vision and pattern recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.304411Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:5322951d3fd606fb13442db1f1e94bcf3973719ae06062897b304abab74e55c3","observation_id":"ec145275-eb5e-4a11-9138-32b70d2d8962","resolution":{"observed_at":"2026-08-03T12:14:37.304411Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:37.410499Z","title":"Advances in Neural Information Processing Systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.410499Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:65ba0ccc397ceda49cf06b684b45c89300d099bcbbad32e5df83fd55a22dff3e","observation_id":"5bd029a9-579d-4aea-80ff-8df7fe7f771e","resolution":{"observed_at":"2026-08-03T12:14:37.410499Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:37.527419Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.527419Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:58ba1b10db477241f267b0a84c8cdc16ea5ecf6214786e56b224155b58a692a4","observation_id":"69f45c97-cb7a-4cb8-a7a5-d5ae826c3f8b","resolution":{"observed_at":"2026-08-03T12:14:37.527419Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:37.732765Z","title":"Carnegie Mellon University - CMU Graphics Lab - motion capture library , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.732765Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d1d6ccf7b8d75f0ac6a52e0c334e2278128f12610f2b7548fdf2bb6038ccd2b1","observation_id":"944172ae-3233-4fe0-aa35-888f42f94f00","resolution":{"observed_at":"2026-08-03T12:14:37.732765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:37.899127Z","title":"Advances in neural information processing systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:37.899127Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:dd0ee239812f6d28065980e1ba300b2e517b201d0eb29edb0512ce1dca1aa5f6","observation_id":"f34a290f-dfc2-4f62-8046-d027d13d19be","resolution":{"observed_at":"2026-08-03T12:14:37.899127Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:38.108655Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:38.108655Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:35dba708718033b9a61b8746cfa73b276a267af03fee184676907d85100b16c5","observation_id":"dc3f5742-d4f0-42d4-b8e3-c426c5a09e74","resolution":{"observed_at":"2026-08-03T12:14:38.108655Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:38.280514Z","title":"North American Chapter of the Association for Computational Linguistics , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:38.280514Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:de3239df4f75fba76232f0fec5026b53394b5c9685644e7c331b302e2d043213","observation_id":"5e1763ea-0cc3-4676-8225-63687e55e1e3","resolution":{"observed_at":"2026-08-03T12:14:38.280514Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:38.460054Z","title":"2019 International Conference on 3D Vision (3DV) , pages=","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:38.460054Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:dd3b6404380f61de0e3b7e6a8a6ef9bef421029141df8ae9e8ce8e441554a255","observation_id":"a846d4ff-a03a-4603-a7a9-06099f0662de","resolution":{"observed_at":"2026-08-03T12:14:38.460054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:38.633727Z","title":"Proceedings of the 28th ACM International Conference on Multimedia , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:38.633727Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:a9f48bd5a38bdf60f88189f6da39451c82d3d6c2d171c6781ab5c227c42d8c0b","observation_id":"5b52b50b-356b-48a8-b8af-29be5236b7b7","resolution":{"observed_at":"2026-08-03T12:14:38.633727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.14165","last_updated":"2020-07-22T19:47:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-28T17:29:03Z","title":"Language Models are Few-Shot Learners","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.14165","snapshot_observed_at":"2026-08-03T12:14:38.721572Z","title":"arXiv preprint arXiv:2005.14165 , year=","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:38.721572Z"},"links":{"cited_paper":"/paper/2005.14165","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:49bbf8d6af106815850d73c7727dd87cb6b0e5a9151ff3e396713767a7c995ea","observation_id":"e5e0ff85-d4a7-4660-a67c-1daa1ec5c639","resolution":{"observed_at":"2026-08-03T12:14:38.721572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:38.818337Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:38.818337Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:6ff6da025397c0b892f3ba096a2911b0f81da7f59a11dc0a18585ebf57e8f483","observation_id":"e8e47630-1e29-4127-bb68-87bb3afc4629","resolution":{"observed_at":"2026-08-03T12:14:38.818337Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:38.948488Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:38.948488Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:a9d4185d23938a885b6bb0f9c862763f7421e2d7ffd253e68530f5eb1583f66b","observation_id":"ae226e74-4a6e-4b93-ba25-2fd65908e478","resolution":{"observed_at":"2026-08-03T12:14:38.948488Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.08718","last_updated":"2022-03-23T19:47:21Z","snapshot_observed_at":"2026-07-06T11:01:02.207193Z","submitted_at":"2021-04-18T05:00:29Z","title":"CLIPScore: A Reference-free Evaluation Metric for Image Captioning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.08718","snapshot_observed_at":"2026-08-03T12:14:39.057156Z","title":"arXiv preprint arXiv:2104.08718 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:39.057156Z"},"links":{"cited_paper":"/paper/2104.08718","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:fd9c3a353e6685b451432ab9c7a72500383df3c4996e536becd5269efd227033","observation_id":"46ae4114-396d-48f0-bd62-149b84140d58","resolution":{"observed_at":"2026-08-03T12:14:39.057156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.06252","last_updated":"2021-07-20T09:20:29Z","snapshot_observed_at":"2026-08-18T02:38:21.579354Z","submitted_at":"2021-07-13T17:22:42Z","title":"Dance2Music: Automatic Dance-driven Music Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.06252","snapshot_observed_at":"2026-08-03T12:14:39.221256Z","title":"arXiv preprint arXiv:2107.06252 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:39.221256Z"},"links":{"cited_paper":"/paper/2107.06252","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:9c15beae5fc7ca64c7ce313f3696d24783a428bb41cacd6c166130e6832ba2f5","observation_id":"10e455ad-ba46-4d61-a95e-51e56d590ff1","resolution":{"observed_at":"2026-08-03T12:14:39.221256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.12598","last_updated":"2022-07-26T01:42:07Z","snapshot_observed_at":"2026-08-14T06:37:15.299690Z","submitted_at":"2022-07-26T01:42:07Z","title":"Classifier-Free Diffusion Guidance","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.12598","snapshot_observed_at":"2026-08-03T12:14:39.387216Z","title":"arXiv preprint arXiv:2207.12598 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:39.387216Z"},"links":{"cited_paper":"/paper/2207.12598","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:7c6451a7abffbfa6666ea0606f55b1cd98619805a9ac655e2a182383b00ca937","observation_id":"0e506965-2f44-4923-8903-b8bd0fd1f13d","resolution":{"observed_at":"2026-08-03T12:14:39.387216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:39.499516Z","title":"Proceedings of the IEEE/CVF conference on computer vision and pattern recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:39.499516Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d824de56c4fa55cb08a2e6214c1af69ea280b7baa00f79719890ff7e4ea7a476","observation_id":"0444a04a-bf9e-4386-8ff6-8de8fd1d6e69","resolution":{"observed_at":"2026-08-03T12:14:39.499516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:39.632306Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:39.632306Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:cb9493857a5641a3d5a770a0264787a448673cfc5dc2120cf7df43b09a9db0f6","observation_id":"93e3bdb6-7d5f-49fa-b748-9255c3e87c3f","resolution":{"observed_at":"2026-08-03T12:14:39.632306Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:39.811798Z","title":"The Eleventh International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:39.811798Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:26f4b995f957c4079bb95a414476ad5ab4521366565dc46ebc99b02745cd44fd","observation_id":"3451493b-6b42-4ae2-9898-eee9779615ce","resolution":{"observed_at":"2026-08-03T12:14:39.811798Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2208.15001","last_updated":"2022-08-31T17:58:54Z","snapshot_observed_at":"2026-08-16T16:35:36.335456Z","submitted_at":"2022-08-31T17:58:54Z","title":"MotionDiffuse: Text-Driven Human Motion Generation with Diffusion Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2208.15001","snapshot_observed_at":"2026-08-03T12:14:39.901339Z","title":"arXiv preprint arXiv:2208.15001 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:39.901339Z"},"links":{"cited_paper":"/paper/2208.15001","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:ffdd2b9afdf99501d7af86e25c7c63f9d32c32f2cd0e7e81ff8a8e2175186cd8","observation_id":"1fae7c83-c503-49ab-9637-373b0bd2c1ad","resolution":{"observed_at":"2026-08-03T12:14:39.901339Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:40.013949Z","title":"Journal of Computing in Civil Engineering , volume=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:40.013949Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d53d7b3116e944053f16ad07a1f7a2754cf0296316f988ec8563393fcd8521ab","observation_id":"f5bb4546-d8ec-4ea3-913a-3ad347e7dc4a","resolution":{"observed_at":"2026-08-03T12:14:40.013949Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:40.185199Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:40.185199Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:ab5e18fec801d8e488a8811af9f60bf3438746089d1518fa0a8af49fe23baff6","observation_id":"4fa3f310-6a68-48d4-8d00-04a0d4b28f98","resolution":{"observed_at":"2026-08-03T12:14:40.185199Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.06052","last_updated":"2023-09-24T17:00:32Z","snapshot_observed_at":"2026-08-16T16:02:22.657116Z","submitted_at":"2023-01-15T09:34:42Z","title":"T2M-GPT: Generating Human Motion from Textual Descriptions with Discrete Representations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.06052","snapshot_observed_at":"2026-08-03T12:14:40.345769Z","title":"arXiv preprint arXiv:2301.06052 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:40.345769Z"},"links":{"cited_paper":"/paper/2301.06052","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:c1a9917f506fbbd520a42a28fff94cde947f3d8e0e0d29a82f1563f99efbaf92","observation_id":"0a8df728-d112-41da-89f5-79700947e755","resolution":{"observed_at":"2026-08-03T12:14:40.345769Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:40.454672Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:40.454672Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:3e8f9cebc385eeb5ab67cab9e40dc7674b600e4d0417a2052d9fe0a99056c224","observation_id":"9c7b3a6e-1ca3-43bc-8636-6259957e3cc9","resolution":{"observed_at":"2026-08-03T12:14:40.454672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:40.627654Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:40.627654Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:458dd81fedc3d1e9439709408e74b2132f25992435d3c27aa6f136932df0b6f8","observation_id":"d858e4f1-631f-4a17-b526-817997964ad8","resolution":{"observed_at":"2026-08-03T12:14:40.627654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:40.746919Z","title":"ACM SIGGRAPH 2024 Conference Papers , pages=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:40.746919Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:9eaf01b2deecd6a53ae18f5534cdfc266a7e8cdb1868685e88814ce5f98c6304","observation_id":"61f615be-f154-4dc7-b695-ae5b7f8e7cbd","resolution":{"observed_at":"2026-08-03T12:14:40.746919Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:40.915104Z","title":"ACM SIGGRAPH 2024 Conference Papers , pages=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:40.915104Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:5b9e3e6a95a9517b0dcdef6b4a785e5922c2d7e345a4c176df7c26053fc6fdff","observation_id":"3678b3d1-e2f5-4d17-8e1a-6e175f7febd8","resolution":{"observed_at":"2026-08-03T12:14:40.915104Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:41.088294Z","title":"ACM Transactions on Graphics (TOG) , volume=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:41.088294Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d263654441752952bb695c576ed2e1ef19025afa17bd4b4f6fb025f71152f600","observation_id":"8c1d0486-763e-4d23-8bbe-7a50d2e3690f","resolution":{"observed_at":"2026-08-03T12:14:41.088294Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.08159","last_updated":"2025-01-23T17:08:57Z","snapshot_observed_at":"2026-08-16T13:10:36.203586Z","submitted_at":"2024-10-10T17:41:54Z","title":"DART: Denoising Autoregressive Transformer for Scalable Text-to-Image Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.08159","snapshot_observed_at":"2026-08-03T12:14:41.172523Z","title":"arXiv preprint arXiv:2410.08159 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:41.172523Z"},"links":{"cited_paper":"/paper/2410.08159","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:ebdde76ef5d066d2138aa2a1e8200d2af4c98971eeb497a90e141af586708320","observation_id":"40481cc8-78a4-49f7-95eb-222b0ca637a5","resolution":{"observed_at":"2026-08-03T12:14:41.172523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03441","last_updated":"2024-10-04T13:56:48Z","snapshot_observed_at":"2026-08-16T13:12:34.424178Z","submitted_at":"2024-10-04T13:56:48Z","title":"CLoSD: Closing the Loop between Simulation and Diffusion for multi-task character control","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.03441","snapshot_observed_at":"2026-08-03T12:14:41.292689Z","title":"arXiv preprint arXiv:2410.03441 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:41.292689Z"},"links":{"cited_paper":"/paper/2410.03441","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:51bf62f57ea1fd211aa7f8d400dea78829e7905166293ea7eb2097afc7e522ef","observation_id":"5d6db429-3e96-45f3-a3e3-336b4abe7fbd","resolution":{"observed_at":"2026-08-03T12:14:41.292689Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:41.452315Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:41.452315Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:35dc1892dc13cd4dd8f0226e2b49e9119ce98ad5a8cf62bfe2888e6728794879","observation_id":"60d9349d-ff77-4ea9-a6d4-e986f08f4d43","resolution":{"observed_at":"2026-08-03T12:14:41.452315Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08740","last_updated":"2024-09-23T15:59:41Z","snapshot_observed_at":"2026-08-18T20:56:40.841401Z","submitted_at":"2024-01-16T18:55:25Z","title":"SiT: Exploring Flow and Diffusion-based Generative Models with Scalable Interpolant Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08740","snapshot_observed_at":"2026-08-03T12:14:41.609978Z","title":"arXiv preprint arXiv:2401.08740 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:41.609978Z"},"links":{"cited_paper":"/paper/2401.08740","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:52716ce5836634ed5c86470ca0509443434584a144a9ea58977538f6403e06ce","observation_id":"3b7fc349-f4bc-4e61-94bc-6bdff6404539","resolution":{"observed_at":"2026-08-03T12:14:41.609978Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:41.700407Z","title":"European Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:41.700407Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:c8725e63247ef6e504686cdd0104d29762643887382a8470a3d0339142f51f8c","observation_id":"0e166cb0-3867-4f1d-a115-f4f81cbff050","resolution":{"observed_at":"2026-08-03T12:14:41.700407Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.19435","last_updated":"2024-04-01T13:02:20Z","snapshot_observed_at":"2026-08-18T20:11:17.591401Z","submitted_at":"2024-03-28T14:04:17Z","title":"BAMM: Bidirectional Autoregressive Motion Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.19435","snapshot_observed_at":"2026-08-03T12:14:41.863166Z","title":"arXiv preprint arXiv:2403.19435 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:41.863166Z"},"links":{"cited_paper":"/paper/2403.19435","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:38b6af4dde0f5f1a4d74446ce1d2e6989e11d3accffaf232590fd81309a06a75","observation_id":"70b675c8-e857-43f1-aeb2-496fc509cdc0","resolution":{"observed_at":"2026-08-03T12:14:41.863166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11838","last_updated":"2024-11-01T14:45:36Z","snapshot_observed_at":"2026-08-16T13:42:07.461220Z","submitted_at":"2024-06-17T17:59:58Z","title":"Autoregressive Image Generation without Vector Quantization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11838","snapshot_observed_at":"2026-08-03T12:14:42.016370Z","title":"arXiv preprint arXiv:2406.11838 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.016370Z"},"links":{"cited_paper":"/paper/2406.11838","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:9b98158a61075653ba73e7334afff61aa8110467364d4cb40f9871eb087e2c5e","observation_id":"fea5a449-874f-42bd-b6e4-43df16ca43f4","resolution":{"observed_at":"2026-08-03T12:14:42.016370Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:42.169998Z","title":"European Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.169998Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:f6ee6a4163ee901dbc368a652cd8eec62facf9dce10f70c330cf3447a18036b5","observation_id":"0303288c-c87c-47c5-9a24-61fc663280b5","resolution":{"observed_at":"2026-08-03T12:14:42.169998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:42.286940Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.286940Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:5c574888c68d921a84d2e7c5c2de41db11186b2d3cf431d28021c3b6e6edde7e","observation_id":"8c24dc2a-64c1-4a8e-bbef-83d1203d8373","resolution":{"observed_at":"2026-08-03T12:14:42.286940Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:42.461772Z","title":"European Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.461772Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:319fac9220b8f01a397d1383708fc68d99a04741b7138507015151868af30125","observation_id":"2f4488d5-c512-4cb5-8ced-82aac819a2a2","resolution":{"observed_at":"2026-08-03T12:14:42.461772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:42.541031Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.541031Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d5ace750a78245634bc3c0c05a05f06544913db9440ac5148842724102ead918","observation_id":"02ca5dc7-661c-4978-b0c0-fb5d406b95e1","resolution":{"observed_at":"2026-08-03T12:14:42.541031Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:42.689532Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.689532Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:99caff49562ce329cf1fd15961a0f4689a60d81acebe2ff8029f4fd8dd463215","observation_id":"cec52566-0a49-42bf-a4ac-6a4c3cd23bf3","resolution":{"observed_at":"2026-08-03T12:14:42.689532Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:42.838645Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.838645Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:b94742753ca9fc509f1a60b84d28d74ee402813c80e53470d4dd5b920c7bddb4","observation_id":"bc7dff15-1630-46cf-b391-3d4c1e2dd4a1","resolution":{"observed_at":"2026-08-03T12:14:42.838645Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:42.986528Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:42.986528Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:30b97ffbceafcec725e113d27be8b1a52aa52529a2e04f6e3341538494777341","observation_id":"e5493709-af40-4c59-a8de-ce15e0daa539","resolution":{"observed_at":"2026-08-03T12:14:42.986528Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.02502","last_updated":"2022-10-05T20:19:21Z","snapshot_observed_at":"2026-08-11T15:38:14.931716Z","submitted_at":"2020-10-06T06:15:51Z","title":"Denoising Diffusion Implicit Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.02502","snapshot_observed_at":"2026-08-03T12:14:43.108656Z","title":"arXiv preprint arXiv:2010.02502 , year=","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:43.108656Z"},"links":{"cited_paper":"/paper/2010.02502","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:9ca274a5a8b3463afc6356e550099711ea7cb8d9fc07c1854d3e0525cbee777f","observation_id":"21219550-58dc-460f-bc52-4611806b9433","resolution":{"observed_at":"2026-08-03T12:14:43.108656Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.19759","last_updated":"2024-12-30T08:43:06Z","snapshot_observed_at":"2026-08-21T16:15:22.647682Z","submitted_at":"2024-04-30T17:59:47Z","title":"MotionLCM: Real-time Controllable Motion Generation via Latent Consistency Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.19759","snapshot_observed_at":"2026-08-03T12:14:43.239574Z","title":"arXiv preprint arXiv:2404.19759 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:43.239574Z"},"links":{"cited_paper":"/paper/2404.19759","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:ad1b5092ad18d23bb4d43caeb4798289b759f48d50d40e30d2afd3f7a573f8c3","observation_id":"2ba0b127-9ab4-40c2-b3e1-ef7de06e0984","resolution":{"observed_at":"2026-08-03T12:14:43.239574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.00752","last_updated":"2024-05-31T17:55:27Z","snapshot_observed_at":"2026-08-17T20:47:46.242385Z","submitted_at":"2023-12-01T18:01:34Z","title":"Mamba: Linear-Time Sequence Modeling with Selective State Spaces","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.00752","snapshot_observed_at":"2026-08-03T12:14:43.365337Z","title":"arXiv preprint arXiv:2312.00752 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:43.365337Z"},"links":{"cited_paper":"/paper/2312.00752","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:027d64646ff7476d004a2f8bf24d5b5fb651bdfe194566410f09095116853075","observation_id":"ff255d01-c22a-44cd-8e5e-149b43ed4040","resolution":{"observed_at":"2026-08-03T12:14:43.365337Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:43.488408Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:43.488408Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:28bf0f4e22c1bf2813faf9b47c8f904426b26eec0f4d914ae02f1c029e48b099","observation_id":"ee851449-d0f5-4802-9602-c05812bdae77","resolution":{"observed_at":"2026-08-03T12:14:43.488408Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06553","last_updated":"2025-07-07T05:09:32Z","snapshot_observed_at":"2026-08-20T10:40:30.650361Z","submitted_at":"2023-12-11T17:41:17Z","title":"HOI-Diff: Text-Driven Synthesis of 3D Human-Object Interactions using Diffusion Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06553","snapshot_observed_at":"2026-08-03T12:14:43.657566Z","title":"arXiv preprint arXiv:2312.06553 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:43.657566Z"},"links":{"cited_paper":"/paper/2312.06553","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:acbef42672e22b5c3f4f0e420004cc3280449f05d2fb0cba6a2a14307bd59128","observation_id":"9705e7ae-ef88-4b4d-a51e-e406836eec26","resolution":{"observed_at":"2026-08-03T12:14:43.657566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:43.800229Z","title":"European Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:43.800229Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:ef1a60291860b10bd9a201fcfe2e44472c68158fc7159bbd67d3b825baa40793","observation_id":"a3cb6662-2841-450e-a790-666826169b81","resolution":{"observed_at":"2026-08-03T12:14:43.800229Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:43.887666Z","title":"3DV , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:43.887666Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:294ca8581bcfd1ec6a37f391796187dc322c312fd0f6e1e66bda2561134f49cd","observation_id":"03f6e9b8-5b16-48a7-92ab-6e11b1517967","resolution":{"observed_at":"2026-08-03T12:14:43.887666Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.005965Z","title":"ECCV , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.005965Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:c38cbee98bf72840200e72860293e47fc408c9d0cf15332ff3a794e740327990","observation_id":"7ebda8bb-76bb-4aaa-b559-8b4b042f0dca","resolution":{"observed_at":"2026-08-03T12:14:44.005965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.148093Z","title":"and Varol, G","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.148093Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:16f285a8cb5dd04c57b504bac21832febae7f82bd5d80211a30390d6f4aff295","observation_id":"0b4a6474-087d-4fd9-9dab-8479a32f263b","resolution":{"observed_at":"2026-08-03T12:14:44.148093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.268005Z","title":"Proceedings of the 5th ACM International Conference on Multimedia in Asia , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.268005Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:23d94fb2c2d8dcf92070361887c7f47fe4b89171b81bdef7e62a7e8b309df11b","observation_id":"2bd23899-809f-4fff-9cf3-bd11e637d54b","resolution":{"observed_at":"2026-08-03T12:14:44.268005Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.408153Z","title":"ICCV , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.408153Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:0660cbf4257239c4c499a9f35fd2ecae4c22178c63b369b690cb303671a43099","observation_id":"d61fa7bc-4a94-457e-979f-1ccedebbfa3d","resolution":{"observed_at":"2026-08-03T12:14:44.408153Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.528851Z","title":"3DV , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.528851Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:f5c0067783b4ee3f9b865c41ff6bbb2a5f881e2958ba46dad7ef3a664e9a532a","observation_id":"848684b9-2d87-41bb-8c0b-77034d75dbee","resolution":{"observed_at":"2026-08-03T12:14:44.528851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.632130Z","title":"ICCV , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.632130Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:1890c7ece5403d9397c754d94f72e5cb99235b7e08d34002b21c410dd7249521","observation_id":"e35f2f3a-39bd-4231-9a55-849a81b08ecd","resolution":{"observed_at":"2026-08-03T12:14:44.632130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.719648Z","title":"CVPR , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.719648Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:2b663aea9e52dc5850164d4278703fc9d1dbe1d1980a856a9593670112907d20","observation_id":"fb1b933a-d845-43e1-84df-239741d40b10","resolution":{"observed_at":"2026-08-03T12:14:44.719648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.813692Z","title":"CVPR , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.813692Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:145bbad89ba301afe2079634f61a65b797e602efccd29bfa47200a8314aefc32","observation_id":"73aa51ed-efb2-4d30-b055-064a1ebe9fe0","resolution":{"observed_at":"2026-08-03T12:14:44.813692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:44.928366Z","title":"CVPR , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:44.928366Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:153cc76b80001e3e9d49332e8af2215b9a9661a5916e8e4d2a5b3c28cf06ef2b","observation_id":"f0055a56-85d7-4650-baf5-0e8f89fee9f2","resolution":{"observed_at":"2026-08-03T12:14:44.928366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:45.046578Z","title":"ICCV , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.046578Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:bcf47fae0945b05f72ce897f0c4ee00541b70bf3eb44a0af9e3f333c398bcd8a","observation_id":"3fb2933a-900a-4949-87a4-fcda1f30d759","resolution":{"observed_at":"2026-08-03T12:14:45.046578Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:45.162490Z","title":"ICCV , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.162490Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:287986ca31979b1b9dfe4d956ff518118bf9a9a14e8acba61388d8f495da6709","observation_id":"4289883c-6b78-48d6-9b18-f895261b9f7d","resolution":{"observed_at":"2026-08-03T12:14:45.162490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:45.277333Z","title":"TOG , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.277333Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:0fff13cef0c348778000b2300b570c662cfc21ab893d603583a3794fb312f932","observation_id":"84f0c151-12ea-4e19-aa5a-ca2d25481f4d","resolution":{"observed_at":"2026-08-03T12:14:45.277333Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.13111","last_updated":"2026-05-19T09:35:40Z","snapshot_observed_at":"2026-08-18T03:31:47.669646Z","submitted_at":"2024-12-17T17:34:52Z","title":"Motion-2-To-3: Leveraging 2D Motion Data for 3D Motion Generations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.13111","snapshot_observed_at":"2026-08-03T12:14:45.393516Z","title":"arXiv preprint arXiv:2412.13111 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.393516Z"},"links":{"cited_paper":"/paper/2412.13111","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:f527ee57342115853c9c6005f608d1abcd1d6b4aff3fe7a48277815ebde075e1","observation_id":"889c9b97-bca6-45c2-9d63-00977966ce54","resolution":{"observed_at":"2026-08-03T12:14:45.393516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:45.515549Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.515549Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:6e819c936bb716294ba97e58596e70ad98d4a4c468306989dcfa40fc0730817f","observation_id":"8414aa1e-dcd3-425a-8b70-b90692db8c49","resolution":{"observed_at":"2026-08-03T12:14:45.515549Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:45.633417Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.633417Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:31ad536be11bb561dee044bcadb63e51be9c0bd80f4dee29012d7a4d628e3e42","observation_id":"2e6c8f1c-cd83-48d3-ab1e-d07872d3f17b","resolution":{"observed_at":"2026-08-03T12:14:45.633417Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:45.741453Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.741453Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:e3aa1f9a4fd871a6365283a4851e1fbf04af3f26677e176029f3c2de87ef3721","observation_id":"cf775d2c-e4b2-428c-b099-f0869e1058f6","resolution":{"observed_at":"2026-08-03T12:14:45.741453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14559","last_updated":"2024-12-19T06:22:19Z","snapshot_observed_at":"2026-08-18T20:14:51.456327Z","submitted_at":"2024-12-19T06:22:19Z","title":"ScaMo: Exploring the Scaling Law in Autoregressive Motion Generation Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.14559","snapshot_observed_at":"2026-08-03T12:14:45.894576Z","title":"arXiv preprint arXiv:2412.14559 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:45.894576Z"},"links":{"cited_paper":"/paper/2412.14559","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:388f9b245dd9dc1162b1166335c91758ae3c61559702f3f94202f7bc9cd246cc","observation_id":"2b0ec709-cd13-4007-ae40-5d5d415739aa","resolution":{"observed_at":"2026-08-03T12:14:45.894576Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14706","last_updated":"2025-06-04T16:54:19Z","snapshot_observed_at":"2026-08-21T08:26:49.509826Z","submitted_at":"2024-12-19T10:19:43Z","title":"EnergyMoGen: Compositional Human Motion Generation with Energy-Based Diffusion Model in Latent Space","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.14706","snapshot_observed_at":"2026-08-03T12:14:46.051810Z","title":"arXiv preprint arXiv:2412.14706 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:46.051810Z"},"links":{"cited_paper":"/paper/2412.14706","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:d8fc2f0b354424d069f7156544397acfadf6359a65751e2bcd1d9c58f45dbceb","observation_id":"b37c3ff4-30c0-46f5-9b61-60680675e571","resolution":{"observed_at":"2026-08-03T12:14:46.051810Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.10790","last_updated":"2025-02-13T16:20:05Z","snapshot_observed_at":"2026-08-20T01:49:35.927685Z","submitted_at":"2024-10-14T17:56:19Z","title":"Sitcom-Crafter: A Plot-Driven Human Motion Generation System in 3D Scenes","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.10790","snapshot_observed_at":"2026-08-03T12:14:46.219736Z","title":"arXiv preprint arXiv:2410.10790 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:46.219736Z"},"links":{"cited_paper":"/paper/2410.10790","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:34f876484a397c6f8628c1e8d3ee9c0b949a597d8f02da30403b4e5ef5979039","observation_id":"595fd122-a607-4ca8-9be9-e46293f0c4ae","resolution":{"observed_at":"2026-08-03T12:14:46.219736Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.14508","last_updated":"2024-10-18T14:43:05Z","snapshot_observed_at":"2026-08-16T13:08:00.835335Z","submitted_at":"2024-10-18T14:43:05Z","title":"LEAD: Latent Realignment for Human Motion Diffusion","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.14508","snapshot_observed_at":"2026-08-03T12:14:46.346509Z","title":"arXiv preprint arXiv:2410.14508 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:46.346509Z"},"links":{"cited_paper":"/paper/2410.14508","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:36384de3bdd83aaf53ceb7b68319433c3d1d0eb4f3b2ffc4ea5058a5761f8e78","observation_id":"03887d89-1bff-4c73-aa48-a04084ead214","resolution":{"observed_at":"2026-08-03T12:14:46.346509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.10010","last_updated":"2025-03-02T07:42:20Z","snapshot_observed_at":"2026-08-16T13:09:50.215188Z","submitted_at":"2024-10-13T21:11:04Z","title":"InterMask: 3D Human Interaction Generation via Collaborative Masked Modeling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.10010","snapshot_observed_at":"2026-08-03T12:14:46.480607Z","title":"arXiv preprint arXiv:2410.10010 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:46.480607Z"},"links":{"cited_paper":"/paper/2410.10010","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:a680e68872d8c62952675d81b4e2850454cc1caaefc207c5fac1ea276a1c50e4","observation_id":"029a9177-c26c-4a5e-a1bb-92f2275bd9d3","resolution":{"observed_at":"2026-08-03T12:14:46.480607Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.13790","last_updated":"2024-10-17T17:31:24Z","snapshot_observed_at":"2026-08-16T13:08:19.384880Z","submitted_at":"2024-10-17T17:31:24Z","title":"MotionBank: A Large-scale Video Motion Benchmark with Disentangled Rule-based Annotations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.13790","snapshot_observed_at":"2026-08-03T12:14:46.647209Z","title":"arXiv preprint arXiv:2410.13790 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:46.647209Z"},"links":{"cited_paper":"/paper/2410.13790","citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:224b92521d3512d7e30f6058cd8717e6b11f36ac952e1e7a972f5e455d256ef2","observation_id":"6d046258-a29e-498f-b267-f310dfbba768","resolution":{"observed_at":"2026-08-03T12:14:46.647209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-03T12:14:46.798018Z","title":"arXiv preprint arXiv:2411.18303 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation","version":1},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-08-03T12:14:46.798018Z"},"links":{"citing_paper":"/paper/2607.29180"},"observation_digest":"sha256:6c997183243b322c8eef0c4fa1668e65ccb7952f0a4bde9e6d4639e94ddcb04e","observation_id":"f0ac23ed-25a2-4894-ac83-4bf84cad5cb0","resolution":{"observed_at":"2026-08-03T12:14:46.798018Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.29180","last_updated":"2026-07-31T09:01:25Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-20T17:04:22.882985Z","submitted_at":"2026-07-31T09:01:25Z","title":"MoRAE: Flow-Friendly Self-Supervised Latents for Text-to-Motion Generation"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":100,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":151},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 100 of 151 outbound references and 0 inbound Pith citation observations for arXiv:2607.29180."}