{"as_of":"2026-08-09T14:56:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2a45d100ad7a708f172ceadaba3fccee662bf59624d47cc473b80853da17d3ab","coverage":[{"denominator":73,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":73,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T15:47:24.413853Z","state":"measured"},{"denominator":75,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":75,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T21:09:08.624951Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-09T22:44:15.079592Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.15064","snapshot_observed_at":"2026-08-05T21:09:08.624951Z","title":"Stableanimator++: Overcoming pose misalignment and face distortion for human image animation,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.09454","last_updated":"2025-08-13T03:11:28Z","snapshot_observed_at":"2026-08-07T21:53:02.449995Z","submitted_at":"2025-08-13T03:11:28Z","title":"Animate-X++: Universal Character Image Animation with Dynamic Backgrounds","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T21:09:08.624951Z"},"links":{"cited_paper":"/paper/2507.15064","citing_paper":"/paper/2508.09454"},"observation_digest":"sha256:c6083759a996d43166642d661d2fb31739e9b3bad2c3a04f61a3a1e0a6e6e1cf","observation_id":"812ca02c-ba14-47e5-8aae-bf98f86f8ccf","resolution":{"observed_at":"2026-08-05T21:09:08.624951Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"cited_work":{"arxiv_id":"2507.15064","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.15064","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Sta- bleanimator++: Overcoming pose misalignment and face distortion for human image animation","venue":null,"work_id":"c5731544-93e6-44b6-a416-ee8dac50dbbb","year":2025},"citing_paper":{"arxiv_id":"2604.21291","last_updated":"2026-04-23T05:10:15Z","snapshot_observed_at":"2026-07-06T23:07:56.487654Z","submitted_at":"2026-04-23T05:10:15Z","title":"Exploring the Role of Synthetic Data Augmentation in Controllable Human-Centric Video Generation","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-09T22:23:04.067757Z"},"links":{"cited_paper":"/paper/2507.15064","citing_paper":"/paper/2604.21291"},"observation_digest":"sha256:db161a89c2c1853e82fc3174cda4894f6eec504bc33b00b1a898bd0d4b476e13","observation_id":"9b8aceb2-c2b7-481f-b00f-8d7a7b0b9970","resolution":{"observed_at":"2026-05-09T22:44:15.081068Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2507.15064/citation-record","integrity":"/paper/2507.15064/integrity","json":"/paper/2507.15064/citation-record.json","paper":"/paper/2507.15064"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.436629Z","title":"Disco: Disentangled control for realistic human dance generation,","venue":null,"work_id":"b70d7ef8-65ca-47b1-a414-9ebfc568c9dd","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.175388Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:5fe02e08fedbba93a7c64ce837f51190597de079d4a31461a713038cabeee4e0","observation_id":"91a67e87-cbb5-4981-9ea0-b249c65bbafb","resolution":{"observed_at":"2026-08-06T15:47:25.439596Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.427144Z","title":"Magicanimate: Temporally consistent human image animation using diffusion model,","venue":null,"work_id":"1215fa8b-e962-4119-9df4-2ea066947cc4","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.179330Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:f7f12a2f356cf07358ad0827c910a9d42cd923d12753abf1fa5c85f7636ad63d","observation_id":"86cde8f3-2459-419d-920e-c8af73febce7","resolution":{"observed_at":"2026-08-06T15:47:25.430323Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.416108Z","title":"Animate anyone: Consistent and controllable image-to-video synthesis for character animation,","venue":null,"work_id":"f8893bfd-3caf-4667-8a9b-cf1608216413","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.183183Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:f68af26028550027f7e39406493b4d3ff9804721057cfda98f95f9345429f564","observation_id":"46fa8901-0265-413a-92f8-85b0466e52f0","resolution":{"observed_at":"2026-08-06T15:47:25.419894Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.405345Z","title":"Champ: Controllable and consistent human image animation with 3d parametric guidance,","venue":null,"work_id":"65c6a5f8-0b74-443a-b056-e9c7a795c881","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.186640Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:89e6bb0f030b3500b98cc974a0d99f7ffe97812306055c3615f424a266a366d7","observation_id":"3e5bf27d-2ad2-4083-ac66-51848ee73bfb","resolution":{"observed_at":"2026-08-06T15:47:25.408860Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01188","last_updated":"2024-06-03T10:51:10Z","snapshot_observed_at":"2026-07-06T18:24:17.711867Z","submitted_at":"2024-06-03T10:51:10Z","title":"UniAnimate: Taming Unified Video Diffusion Models for Consistent Human Image Animation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01188","snapshot_observed_at":"2026-08-06T15:47:24.190367Z","title":"Unianimate: Taming unified video diffusion models for consistent human image animation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.190367Z"},"links":{"cited_paper":"/paper/2406.01188","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:cb0b80f8d671e89851e781f3abfcf3d4e540cacdc5a87daeaf65ca0a6bc6b6f0","observation_id":"611f4827-3d5b-43b0-918d-38838403f569","resolution":{"observed_at":"2026-08-06T15:47:24.190367Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19680","last_updated":"2025-06-27T10:06:13Z","snapshot_observed_at":"2026-08-04T16:03:06.733403Z","submitted_at":"2024-06-28T06:40:53Z","title":"MimicMotion: High-Quality Human Motion Video Generation with Confidence-aware Pose Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.19680","snapshot_observed_at":"2026-08-06T15:47:24.194731Z","title":"Mimicmotion: High-quality human motion video generation with confidence-aware pose guidance,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.194731Z"},"links":{"cited_paper":"/paper/2406.19680","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:5afe7a2eec686ba5a4c8bbdf311bdcd6b6f4675c4ad745f069f8b742ba818a1d","observation_id":"33c4678c-195f-465d-8f13-39242f3d972e","resolution":{"observed_at":"2026-08-06T15:47:24.194731Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.06070","last_updated":"2025-03-08T08:43:03Z","snapshot_observed_at":"2026-07-06T18:59:34.520274Z","submitted_at":"2024-08-12T11:41:18Z","title":"ControlNeXt: Powerful and Efficient Control for Image and Video Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.06070","snapshot_observed_at":"2026-08-06T15:47:24.198418Z","title":"Controlnext: Powerful and efficient control for image and video generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.198418Z"},"links":{"cited_paper":"/paper/2408.06070","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:6ab41d09000478203509042c2d286d77925da9429945b57dfd6609bcf129a07c","observation_id":"dd0412ee-944f-43d0-8b8e-06be0b9823da","resolution":{"observed_at":"2026-08-06T15:47:24.198418Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.393857Z","title":"Animate-x: Universal character image animation with enhanced motion representation,","venue":null,"work_id":"10f1fbbd-5c53-42f6-a887-2aa6b2d5b999","year":2025},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.202630Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:56787777ba64454285f2387ea0d9412c97d24938bbe6364e6146fc5df98b8c51","observation_id":"69a83fa6-de4c-49c2-937a-fae25ec9ab96","resolution":{"observed_at":"2026-08-06T15:47:25.398424Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.383616Z","title":"Multiple biological granularities network for person re-identification,","venue":null,"work_id":"d70f26cf-0818-4348-9258-1e8c08c3d51d","year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.205683Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:8f8ec0592501338fe379bbc82886dab0592e1b2ac39454d1fb3d72ba1ae1226a","observation_id":"5b34832a-c05b-41e6-a6a0-0260744893fe","resolution":{"observed_at":"2026-08-06T15:47:25.387046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.372752Z","title":"Implicit temporal modeling with learnable alignment for video recognition,","venue":null,"work_id":"a5d53abd-a968-467d-b243-abf74be17bea","year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.208894Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:dd70199dd58a75eb47f9692d8166e8619c2f6a52a10a6bd214051f8ae375b111","observation_id":"0e74353e-5c7c-47c7-a148-0ed5941075fa","resolution":{"observed_at":"2026-08-06T15:47:25.376261Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.361833Z","title":"Diffusion models beat gans on image synthesis,","venue":null,"work_id":"619b4fa3-9668-4a2d-9103-4d5ab5ba3488","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.212277Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:de457e694a93d9ae55321f319286169e977e19fb8af4d0e1e1284a556bbabf77","observation_id":"94ec5800-9dfc-44f5-adb2-e69bcae56636","resolution":{"observed_at":"2026-08-06T15:47:25.365188Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.215699Z","title":"Denoising diffusion probabilistic models,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.215699Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:7a2d5395dc43a103ee78d4bdadc6492cdb2f14ec05ada76ea3cca8bd1b8d6353","observation_id":"d6b6bf79-0620-46f1-923b-6b9e6c0c8659","resolution":{"observed_at":"2026-08-06T15:47:24.215699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.344772Z","title":"Cascaded diffusion models for high fidelity image generation,","venue":null,"work_id":"b6e35256-231d-42e4-b651-c31a1426432d","year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.219020Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:08b0c17907fb2cda93901cebd95bc61892d6713503a1aa3e36a00e8f60d9fc34","observation_id":"a939dad6-06d2-4753-ad67-ad2c464f513d","resolution":{"observed_at":"2026-08-06T15:47:25.348455Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.222522Z","title":"Score-based generative modeling through stochastic differ- ential equations,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.222522Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:9f8a2cb62b71d0e53dae82fc3c102140c933e169cac2f3950e331324c3c02d26","observation_id":"1b02a145-6e6b-43e7-905f-a37c6c7c8fee","resolution":{"observed_at":"2026-08-06T15:47:24.222522Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.225805Z","title":"Denoising diffusion implicit models,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.225805Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:a9caf71a18beac0f2230b8e5970656812f9e025ffbbb39ffbdaae11f931b3df6","observation_id":"32290632-f262-4310-bdef-c8f7cb220a1d","resolution":{"observed_at":"2026-08-06T15:47:24.225805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.228956Z","title":"High- resolution image synthesis with latent diffusion models,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.228956Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:481e214157412c3ce66f62e1ec941f4929f7d2eb54f285ad77cfd509755cf54a","observation_id":"5eb42031-c161-4068-8445-0abff23c5d89","resolution":{"observed_at":"2026-08-06T15:47:24.228956Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.313329Z","title":"Sdedit: Guided image synthesis and editing with stochastic differential equations,","venue":null,"work_id":"c7cd90b2-20df-42ae-ae5f-3161d9b02852","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.232855Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:c3159733084949963161ead514d2f147ff37a47e9f1adeb2ba6e552a762f65f9","observation_id":"899e6636-cd34-4bad-b327-66f1d67d0c63","resolution":{"observed_at":"2026-08-06T15:47:25.316598Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.302289Z","title":"Plug-and-play diffusion features for text-driven image-to-image translation,","venue":null,"work_id":"c223e06f-fd8b-48a3-b3c2-758585ac74c7","year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.236510Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:f138253f33e61a8770496bb436dfff0f96d445506a6afd2cacf30988afab3136","observation_id":"f0d8c598-8f03-41a4-8d74-666f79efea4c","resolution":{"observed_at":"2026-08-06T15:47:25.306108Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.15241","last_updated":"2024-11-12T06:08:29Z","snapshot_observed_at":"2026-07-06T19:06:41.697671Z","submitted_at":"2024-08-27T17:59:41Z","title":"GenRec: Unifying Video Generation and Recognition with Diffusion Models","version":2},"cited_work":{"arxiv_id":"2408.15241","doi":null,"metadata_source":"pith","pith_arxiv_id":"2408.15241","snapshot_observed_at":"2026-08-06T15:47:24.730283Z","title":"GenRec: Unifying Video Generation and Recognition with Diffusion Models","venue":"cs.CV","work_id":"d823a9b3-0608-4d08-a3b0-c7744b7c5bdb","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.240018Z"},"links":{"cited_paper":"/paper/2408.15241","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:16f3f2392f3f7688f4f8e155e5e42819c1e29e68991edf5ecfa365bf7a749596","observation_id":"6dc58597-4ca5-41ca-a20b-bd0bbc6140b0","resolution":{"observed_at":"2026-08-06T15:47:24.734543Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.291376Z","title":"A survey on video diffusion models,","venue":null,"work_id":"9415b3a6-fa18-4f2d-9a06-380f3fb05cd0","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.243613Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:2fb10fee303ff211997bc547fe92940f9f77eaf6ce79c805aee1c0ffa2c20496","observation_id":"dfc10fff-f074-45aa-b734-14326fe2dca9","resolution":{"observed_at":"2026-08-06T15:47:25.294826Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.281256Z","title":"Simda: Simple diffusion adapter for efficient video generation,","venue":null,"work_id":"de102064-416b-4265-9d93-21b8af03d9d0","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.246723Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:658ba0ef0a8984b385f68766e8570aaff79f9c739357bd9292d7cc370810d8d4","observation_id":"39a7802f-c379-46e1-9cd4-3f34d46f902d","resolution":{"observed_at":"2026-08-06T15:47:25.284697Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06465","last_updated":"2024-06-10T17:02:08Z","snapshot_observed_at":"2026-08-02T18:07:43.330243Z","submitted_at":"2024-06-10T17:02:08Z","title":"AID: Adapting Image2Video Diffusion Models for Instruction-guided Video Prediction","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06465","snapshot_observed_at":"2026-08-06T15:47:24.249727Z","title":"Aid: Adapting image2video diffusion models for instruction-guided video prediction,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.249727Z"},"links":{"cited_paper":"/paper/2406.06465","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:3b1c0fcc5c8abfbfeede51a5ccbe41d9c18c3c63727d5e696b5bf7e1e36f1b22","observation_id":"d69b72d3-71d1-4c9f-86a5-bd49b333c525","resolution":{"observed_at":"2026-08-06T15:47:24.249727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.253031Z","title":"Magicmotion: Controllable video generation with dense-to-sparse trajectory guidance,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.253031Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:af00892141a04942bf31a8617fccb3b4e48776beed4fdfae6b8c2e95860679ae","observation_id":"ab66451e-28af-4d10-833b-73ff8e054f39","resolution":{"observed_at":"2026-08-06T15:47:24.253031Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.06721","last_updated":"2023-08-13T08:34:51Z","snapshot_observed_at":"2026-07-06T16:05:39.158819Z","submitted_at":"2023-08-13T08:34:51Z","title":"IP-Adapter: Text Compatible Image Prompt Adapter for Text-to-Image Diffusion Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.06721","snapshot_observed_at":"2026-08-06T15:47:24.255974Z","title":"Ip-adapter: Text compatible image prompt adapter for text-to-image diffusion models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.255974Z"},"links":{"cited_paper":"/paper/2308.06721","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:cbe649f62e37440842e48d193c9922c974d8c8fb64e8dc5bb3cef70190f15b1b","observation_id":"98a3a3c7-2c88-4ace-83a5-e46ddd785683","resolution":{"observed_at":"2026-08-06T15:47:24.255974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.07519","last_updated":"2024-02-02T16:15:22Z","snapshot_observed_at":"2026-08-06T23:08:20.817133Z","submitted_at":"2024-01-15T07:50:18Z","title":"InstantID: Zero-shot Identity-Preserving Generation in Seconds","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.07519","snapshot_observed_at":"2026-08-06T15:47:24.259245Z","title":"Instantid: Zero-shot identity-preserving generation in seconds,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.259245Z"},"links":{"cited_paper":"/paper/2401.07519","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:823ba01373e29237a87a64b9309295528a1dd6233934152000555af113541551","observation_id":"51f0bd20-46b9-4591-ae45-1f7446b29bf1","resolution":{"observed_at":"2026-08-06T15:47:24.259245Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16771","last_updated":"2024-12-28T17:42:44Z","snapshot_observed_at":"2026-08-07T05:22:26.276397Z","submitted_at":"2024-04-25T17:23:43Z","title":"ConsistentID: Portrait Generation with Multimodal Fine-Grained Identity Preserving","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.16771","snapshot_observed_at":"2026-08-06T15:47:24.262485Z","title":"Consistentid: Portrait generation with multimodal fine-grained identity preserving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.262485Z"},"links":{"cited_paper":"/paper/2404.16771","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:864f6ce54eb71cbd35bb3d7b3bd78349c2054019af760507a6189fa5d6f2d81b","observation_id":"082a7975-f46e-4536-b03a-685ce5d5a53f","resolution":{"observed_at":"2026-08-06T15:47:24.262485Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.271217Z","title":"Pulid: Pure and lightning id customization via contrastive alignment,","venue":null,"work_id":"7443158c-ff14-46be-acab-00041e0413f9","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.265829Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:5588b7b846aeb8b7839c71834115ca3358fbd94821efdd608c1cdff338bb93b5","observation_id":"53a27cab-d78a-4c3f-a534-7ce920135086","resolution":{"observed_at":"2026-08-06T15:47:25.274530Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.262016Z","title":"Facefusion,","venue":null,"work_id":"ba59a540-95bf-4b64-acf7-bd351cd60739","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.268759Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:a725138933ed531af9d76913d34d2bd2195ed8a4425b26950bba167429f30393","observation_id":"e23671c9-afee-4d8b-90de-b7d120e173c8","resolution":{"observed_at":"2026-08-06T15:47:25.265016Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.252076Z","title":"Towards real-world blind face restoration with generative facial prior,","venue":null,"work_id":"87d8fcdd-d741-4045-80f1-f456bf9deaa8","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.271655Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:99df5ab970576c202ce0c9eaae5ef27beccff366daaea458eb8fd40ace308a10","observation_id":"a0c155b8-2ee2-4014-9b8d-8abd29e25edc","resolution":{"observed_at":"2026-08-06T15:47:25.255343Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.275159Z","title":"Towards robust blind face restoration with codebook lookup transformer,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.275159Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:16a60dadd6bd803b5e097e33c1b9d33b81cb32ca9e3dd638d78f78517ea4f41d","observation_id":"7526104a-1f94-4f8a-b3fc-f137686bb8bd","resolution":{"observed_at":"2026-08-06T15:47:24.275159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.234845Z","title":"Motioneditor: Editing video motion via content-aware diffusion,","venue":null,"work_id":"e7f41351-f6b7-46ef-a1f9-7a2a5f75c85b","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.278386Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:2518d3f9a9fc94582714f41125a3cb396d3ab6f3afdf5dbc131ce80432b3e03b","observation_id":"5b792294-1c12-4a56-99ee-3ebb2c04bfd0","resolution":{"observed_at":"2026-08-06T15:47:25.238269Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.281600Z","title":"Arcface: Additive angular margin loss for deep face recognition,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.281600Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:d76c754d693e46ab241415567e73c42b1530d2d584a79a335dcfef2097f11a06","observation_id":"c7b5eca7-4a2b-479e-93f5-f71bbcf82666","resolution":{"observed_at":"2026-08-06T15:47:24.281600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.284695Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.284695Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:ce5f64dbc09c17968ae5e7a9975e8a72bae3bdcf30ed4ae515d30fa86177f4aa","observation_id":"f6d6d2e7-3a38-41df-852f-ef70e019cca3","resolution":{"observed_at":"2026-08-06T15:47:24.284695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.103560Z","title":"Bardi, I","venue":null,"work_id":"1440e493-40d1-49ea-8dd2-9229ed0297d9","year":1997},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.287550Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:e725f602ee4cef066e38640b95f28508525a11292e311a0c04a3e91eea936e62","observation_id":"968e6d78-5478-4ad2-b50e-6e0cfc304649","resolution":{"observed_at":"2026-08-06T15:47:25.107663Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.092812Z","title":"Stochastic hamilton–jacobi–bellman equations,","venue":null,"work_id":"7544ac71-416b-42f7-a679-d9177975935e","year":1992},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.291131Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:d59a17c4635df53ac139625747d08f0c9dc589f2014b4ed202576e88537fc284","observation_id":"1b5c981a-c766-464e-b2f6-92eee3d8c5c2","resolution":{"observed_at":"2026-08-06T15:47:25.096341Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17697","last_updated":"2024-11-27T07:39:20Z","snapshot_observed_at":"2026-08-06T12:03:20.133598Z","submitted_at":"2024-11-26T18:59:22Z","title":"StableAnimator: High-Quality Identity-Preserving Human Image Animation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17697","snapshot_observed_at":"2026-08-06T15:47:24.294368Z","title":"Sta- bleanimator: High-quality identity-preserving human image animation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.294368Z"},"links":{"cited_paper":"/paper/2411.17697","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:b07c3eaef47777eb1c00676a237687f546d5670357155159c1c38437d2bda8c0","observation_id":"d7acabfa-108d-42a4-9331-43fa61e70b37","resolution":{"observed_at":"2026-08-06T15:47:24.294368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.082025Z","title":"Improved denoising diffusion proba- bilistic models,","venue":null,"work_id":"78f79c63-c22f-4cd4-b784-51e16f21378a","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.298461Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:9b7545ef669c215b8445aa5046d141351d69f25a7382b9507a730fcb31e685d4","observation_id":"c5c60b8b-80fa-437a-924e-67a06ba17096","resolution":{"observed_at":"2026-08-06T15:47:25.085335Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2208.01626","last_updated":"2022-08-02T17:55:41Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-08-02T17:55:41Z","title":"Prompt-to-Prompt Image Editing with Cross Attention Control","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2208.01626","snapshot_observed_at":"2026-08-06T15:47:24.301791Z","title":"Prompt-to-prompt image editing with cross attention control,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.301791Z"},"links":{"cited_paper":"/paper/2208.01626","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:abcaf55730cdff2923ceab7d603a73d982f01e7f0473a810e037ccda042f27ec","observation_id":"69092112-06d7-4f28-86ed-f20e4a608c7d","resolution":{"observed_at":"2026-08-06T15:47:24.301791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.14792","last_updated":"2022-09-29T13:59:46Z","snapshot_observed_at":"2026-07-06T13:57:47.051387Z","submitted_at":"2022-09-29T13:59:46Z","title":"Make-A-Video: Text-to-Video Generation without Text-Video Data","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.14792","snapshot_observed_at":"2026-08-06T15:47:24.305043Z","title":"Make-a-video: Text-to-video generation without text-video data,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.305043Z"},"links":{"cited_paper":"/paper/2209.14792","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:96608d3cdcf69531b228dd644b6156bad58f8e6e0dfe6e38617c33b3033d53bc","observation_id":"149b3ae6-f103-4208-a507-1caba5ddf359","resolution":{"observed_at":"2026-08-06T15:47:24.305043Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.308407Z","title":"Animatediff: Animate your personalized text-to- image diffusion models without specific tuning,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.308407Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:bb684988d2e2b63b02ce91c565d1a77f662b5a33f47957b0e4da4279ab596c44","observation_id":"545b43c9-2eb0-4ca8-8bf8-2ee5bae34c3c","resolution":{"observed_at":"2026-08-06T15:47:24.308407Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.059856Z","title":"Tune-a-video: One-shot tuning of image diffusion models for text-to-video generation,","venue":null,"work_id":"50bd590c-24d1-4151-9004-cdbb756afac4","year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.311550Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:bf40f17eef7878f745ba414ef4be2d0688a4042834e1825e8bb4c0e8d77ebae7","observation_id":"7304a1af-00c7-4038-9587-1cd81c521709","resolution":{"observed_at":"2026-08-06T15:47:25.066972Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04468","last_updated":"2024-01-09T10:12:52Z","snapshot_observed_at":"2026-08-08T21:00:07.723397Z","submitted_at":"2024-01-09T10:12:52Z","title":"MagicVideo-V2: Multi-Stage High-Aesthetic Video Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04468","snapshot_observed_at":"2026-08-06T15:47:24.314663Z","title":"Magicvideo-v2: Multi-stage high-aesthetic video generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.314663Z"},"links":{"cited_paper":"/paper/2401.04468","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:5a743bd8e4bd510acfc01b2847ac3732685a64a17553031cad5298e12970095b","observation_id":"25e7b372-9242-490d-a6e6-1ffa4d75578f","resolution":{"observed_at":"2026-08-06T15:47:24.314663Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.318727Z","title":"Video generation models as world simulators,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.318727Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:38cebf131a6e1335cadf0cae6361e463fce3fb67fcafc0cbef03927d583c0603","observation_id":"b690f38c-cb11-4a59-990f-474a1d7b407d","resolution":{"observed_at":"2026-08-06T15:47:24.318727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20325","last_updated":"2024-05-30T17:57:30Z","snapshot_observed_at":"2026-08-05T00:06:43.226806Z","submitted_at":"2024-05-30T17:57:30Z","title":"MotionFollower: Editing Video Motion via Lightweight Score-Guided Diffusion","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.20325","snapshot_observed_at":"2026-08-06T15:47:24.321901Z","title":"Motionfollower: Editing video motion via lightweight score-guided diffusion,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.321901Z"},"links":{"cited_paper":"/paper/2405.20325","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:5d1fef2d1655ff7ccd096eff9cf2bdc237dd5aa80ee51355337e4f89f3f739f2","observation_id":"6fae8a15-7a2c-4b6a-9949-492d32168ac4","resolution":{"observed_at":"2026-08-06T15:47:24.321901Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.325271Z","title":"Scalable diffusion models with transformers,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.325271Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:2ec0554dca4923c8828dc622d84746dc7350af0372b0cca7fe6a0c69147a64bf","observation_id":"59386276-a10b-4105-97aa-f31969421527","resolution":{"observed_at":"2026-08-06T15:47:24.325271Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.10157","last_updated":"2021-09-14T21:20:06Z","snapshot_observed_at":"2026-07-06T11:02:11.424356Z","submitted_at":"2021-04-20T17:58:03Z","title":"VideoGPT: Video Generation using VQ-VAE and Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.10157","snapshot_observed_at":"2026-08-06T15:47:24.328409Z","title":"Videogpt: Video gener- ation using vq-vae and transformers,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.328409Z"},"links":{"cited_paper":"/paper/2104.10157","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:11f827422d5d36e5e362b54c3b610e2bee8b19bba7756ffd49531fd2803d1b05","observation_id":"50ff5aae-5bef-4adc-86f2-aab0c33d4baf","resolution":{"observed_at":"2026-08-06T15:47:24.328409Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.034600Z","title":"Magvit: Masked generative video transformer,","venue":null,"work_id":"d4aa4dd1-70ad-4029-8ef6-aca5978c2496","year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.331646Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:6ea0cce2611b8dca9d564d0c70467b7e7a32e82720fe2b6677ccd7feafcca8f8","observation_id":"d2ba86d3-3a09-41d2-9fa9-4a36febf1c14","resolution":{"observed_at":"2026-08-06T15:47:25.038544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03048","last_updated":"2025-05-01T09:40:21Z","snapshot_observed_at":"2026-08-02T13:01:06.918463Z","submitted_at":"2024-01-05T19:55:15Z","title":"Latte: Latent Diffusion Transformer for Video Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03048","snapshot_observed_at":"2026-08-06T15:47:24.334771Z","title":"Latte: Latent diffusion transformer for video generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.334771Z"},"links":{"cited_paper":"/paper/2401.03048","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:0621e787e6538822897d18bffc395f4419c70d6bb5eb34d1c5cf18c20164391f","observation_id":"4bac6244-472d-4d0e-96cb-d946ee2aeb3b","resolution":{"observed_at":"2026-08-06T15:47:24.334771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04233","last_updated":"2024-05-07T11:52:49Z","snapshot_observed_at":"2026-08-09T05:07:43.521969Z","submitted_at":"2024-05-07T11:52:49Z","title":"Vidu: a Highly Consistent, Dynamic and Skilled Text-to-Video Generator with Diffusion Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04233","snapshot_observed_at":"2026-08-06T15:47:24.337873Z","title":"Vidu: a highly consistent, dynamic and skilled text-to-video generator with diffusion models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.337873Z"},"links":{"cited_paper":"/paper/2405.04233","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:96e31971c15cbbd0d2604332672133a43ca8b3faf22aba664510ebddca87d53a","observation_id":"19fcc1b4-f978-469b-bae8-ef51b579d778","resolution":{"observed_at":"2026-08-06T15:47:24.337873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.15868","last_updated":"2022-05-29T19:02:15Z","snapshot_observed_at":"2026-07-06T13:15:58.303738Z","submitted_at":"2022-05-29T19:02:15Z","title":"CogVideo: Large-scale Pretraining for Text-to-Video Generation via Transformers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.15868","snapshot_observed_at":"2026-08-06T15:47:24.340967Z","title":"Cogvideo: Large- scale pretraining for text-to-video generation via transformers,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.340967Z"},"links":{"cited_paper":"/paper/2205.15868","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:c5bb2ac2dd2cefda803c8ed872ac544447e50bb2d6aa07011449fb995053245a","observation_id":"734e55c8-0ca9-47b3-97ba-4848e9c0df06","resolution":{"observed_at":"2026-08-06T15:47:24.340967Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.03603","last_updated":"2025-03-11T08:14:25Z","snapshot_observed_at":"2026-08-03T00:44:01.942521Z","submitted_at":"2024-12-03T23:52:37Z","title":"HunyuanVideo: A Systematic Framework For Large Video Generative Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.03603","snapshot_observed_at":"2026-08-06T15:47:24.344323Z","title":"Hunyuanvideo: A systematic framework for large video generative models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.344323Z"},"links":{"cited_paper":"/paper/2412.03603","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:6e3cacb000d92a5023a7afe3191a76104f9892e80b999c0275915c168887414c","observation_id":"313456ff-a76a-4d03-a41e-91612cc2c8cd","resolution":{"observed_at":"2026-08-06T15:47:24.344323Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.15127","last_updated":"2023-11-25T22:28:38Z","snapshot_observed_at":"2026-08-07T21:47:08.589400Z","submitted_at":"2023-11-25T22:28:38Z","title":"Stable Video Diffusion: Scaling Latent Video Diffusion Models to Large Datasets","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.15127","snapshot_observed_at":"2026-08-06T15:47:24.347204Z","title":"Stable video diffusion: Scaling latent video diffusion models to large datasets,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.347204Z"},"links":{"cited_paper":"/paper/2311.15127","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:877ae92001ebcf299170e2b09cac1dd9d75ee3cd5442619c2aa2213046a75262","observation_id":"9abe239a-4641-4e4f-b275-793dddf924e4","resolution":{"observed_at":"2026-08-06T15:47:24.347204Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.024233Z","title":"First order motion model for image animation,","venue":null,"work_id":"65c109fa-db5c-4bcb-b7ca-f78e43da41a6","year":2019},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.350325Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:9e4a84bba385392ba7f1032325c2005d468f71f407f10fc5f69a38857a0b361c","observation_id":"653a5f77-4718-4ebc-a20a-c6dd0d0bf640","resolution":{"observed_at":"2026-08-06T15:47:25.027626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.014024Z","title":"Motion representations for articulated animation,","venue":null,"work_id":"fab8dcc9-19da-449a-87ea-f04d0a62198b","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.353262Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:8c764b2df744eb1521b157006b9973cb5e3d0677a79234e1880fd29c25077f19","observation_id":"bb512882-df4a-47c8-84ce-cae2cebc085b","resolution":{"observed_at":"2026-08-06T15:47:25.017329Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:25.003047Z","title":"Few-shot human motion transfer by personalized geometry and texture modeling,","venue":null,"work_id":"4d332e13-fbf7-4baf-b3a0-0cf4c9b7e216","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.356218Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:2cfbd472fa39579c159312065e36678fea4032360e0263bc579316e4c687aa18","observation_id":"fc9a871c-067a-4579-bea4-0943d2cc6eb4","resolution":{"observed_at":"2026-08-06T15:47:25.006666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.990710Z","title":"Generative adversarial networks,","venue":null,"work_id":"7965a287-b2ef-41b2-a55f-daa4baafde73","year":2020},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.359293Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:2828290affa49028a8bac44d8481ba9fcb33e58ccd35929f0e712bc12ac63314","observation_id":"62a9eed1-5b51-442b-827b-006c71cecfb8","resolution":{"observed_at":"2026-08-06T15:47:24.994730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.979172Z","title":"Transformers are SSMs: Generalized models and efficient algorithms through structured state space duality,","venue":null,"work_id":"a1fc638e-9a2d-461e-ba1d-8a8035eb9ce8","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.362276Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:2ac1a8be9cbd4330aa628f8569596890618b3589237199394c412d752c19b2fe","observation_id":"dab0004b-325e-4f61-b703-13ee461e962e","resolution":{"observed_at":"2026-08-06T15:47:24.982737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.968072Z","title":"Lora: Low-rank adaptation of large language models,","venue":null,"work_id":"b53a6a9a-48a0-499d-b374-d37576f65787","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.365530Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:37e7ab70309730bea36bd5e3716ebdcfd38aba22874523489c604906ba4ba62f","observation_id":"c1907a3c-2166-4eac-bb51-ef37715e1acf","resolution":{"observed_at":"2026-08-06T15:47:24.971823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.957939Z","title":"Photomaker: Customizing realistic human photos via stacked id embedding,","venue":null,"work_id":"3be88204-2a23-4155-b3d4-a812b7268845","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.368806Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:c2621ae687ddddb3c5909f1aa0286818e27b4053deb27fa026dce4eb83823226","observation_id":"2756b243-d59f-4b9b-ae9f-d539d162965c","resolution":{"observed_at":"2026-08-06T15:47:24.961237Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.02663","last_updated":"2023-12-06T12:23:36Z","snapshot_observed_at":"2026-08-03T22:37:55.724576Z","submitted_at":"2023-12-05T11:02:45Z","title":"FaceStudio: Put Your Face Everywhere in Seconds","version":2},"cited_work":{"arxiv_id":"2312.02663","doi":null,"metadata_source":"pith","pith_arxiv_id":"2312.02663","snapshot_observed_at":"2026-08-06T15:47:24.470207Z","title":"FaceStudio: Put Your Face Everywhere in Seconds","venue":"cs.CV","work_id":"f9886d5d-e795-410d-aa87-f1f41974468c","year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.371894Z"},"links":{"cited_paper":"/paper/2312.02663","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:49d013b75638748740ff5e31b6586a925a44801cd4194334580bba2a0efa1174","observation_id":"8b257221-6ef0-4a09-8dfd-185965bda203","resolution":{"observed_at":"2026-08-06T15:47:24.475883Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.931211Z","title":"Auto-encoding variational bayes,","venue":null,"work_id":"2d319c45-c492-4b3c-b60a-c8b72caa868f","year":2014},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.375036Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:50aef31ae20c27da83e856d9f152d48403cffeb4e0bca10187eb4f4b35a1dd39","observation_id":"03e59538-02ac-4466-9cb7-92dd726bec02","resolution":{"observed_at":"2026-08-06T15:47:24.944696Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.902447Z","title":"Effective whole-body pose estimation with two-stages distillation,","venue":null,"work_id":"2b003836-2094-4106-ae49-ce1e8c697bbb","year":2023},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.378158Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:1c7de5720ff3015a855be03462d5aaed2b2fba4a15d50ff2c2e6b320fbc44afd","observation_id":"2ff94b09-8529-4e37-b61e-23f3de66f198","resolution":{"observed_at":"2026-08-06T15:47:24.916078Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.381097Z","title":"Method for registration of 3-d shapes,","venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.381097Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:328c2ccc8db4e571e190b9912d68bfb23bdac5a18424bb14fe174b7688a88790","observation_id":"2fbcd60d-3dec-4575-b21d-497a066e023a","resolution":{"observed_at":"2026-08-06T15:47:24.381097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.864970Z","title":"Generative modeling with phase stochastic bridges,","venue":null,"work_id":"41cbc85c-349d-4232-81c0-c6ab0c44bda7","year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.384627Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:6b88d5c456810fc51617c3525126a16550a26c6db8b528a4b0e09c269e38c1c9","observation_id":"cdd08634-e2e6-4ba6-a133-f9aaa88a720f","resolution":{"observed_at":"2026-08-06T15:47:24.876167Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.837416Z","title":"Elucidating the design space of diffusion-based generative models,","venue":null,"work_id":"6ddfbfde-67fd-4310-8d5c-9a49a6e35090","year":2022},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.387609Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:f90617fd551d66ed57ff882b85d43391a77469abdabcb97014bc8d65ca3ca1cf","observation_id":"f130e83d-6e63-4a6e-9a66-39330be61fe4","resolution":{"observed_at":"2026-08-06T15:47:24.846355Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.826346Z","title":null,"venue":null,"work_id":"5e5f6359-f2c0-49b6-b20b-2672489eb6ac","year":2004},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.390899Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:a99a1856e3868f8a84819d1f4591abf22c25d7f1767a5cef9d3fc4cb2af7b13c","observation_id":"a6bb6481-4350-42e2-b1cc-cf54fa921634","resolution":{"observed_at":"2026-08-06T15:47:24.829637Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.815080Z","title":null,"venue":null,"work_id":"b7c359a8-6ee5-4f21-b424-cc749cad7540","year":2012},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.394331Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:6177f2561b225fc4bcff1a556cb769cd3fc73060b32e3ef9c9a81e1ea28611b9","observation_id":"4f9e3f32-b491-439e-a7bd-dc2a3bb552d9","resolution":{"observed_at":"2026-08-06T15:47:24.818724Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.803388Z","title":"Tweedie’s formula and selection bias,","venue":null,"work_id":"12e57117-4f17-4d59-854a-efee71d130d3","year":2011},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.397470Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:11fd7b0305b3f266d4719367094522a8b56e22ee4922b0fa81017d39cc1622ae","observation_id":"5d77bdb0-1845-4feb-979c-f62b5c1394e4","resolution":{"observed_at":"2026-08-06T15:47:24.807110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.790466Z","title":"Learning high fidelity depths of dressed humans by watching social media dance videos,","venue":null,"work_id":"4d77e28e-57a0-4171-abb9-56910bb1c3fe","year":2021},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.400731Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:d697506dddf84efaf48094845060008818c35e8bf30dbf896089fe1322d3cb7e","observation_id":"75799053-daff-479c-af76-d08fb95ff8e5","resolution":{"observed_at":"2026-08-06T15:47:24.794024Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.779256Z","title":"Image quality metrics: Psnr vs. ssim,","venue":null,"work_id":"43fc0665-ca20-40dd-8b33-1dbdf47a950c","year":2010},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.404144Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:d048074897fd4db2f8fb84facd32bc16d3f4195f69d4bfad0f2a5093941b1755","observation_id":"9789f167-44c5-4749-aa25-82758f52ee0e","resolution":{"observed_at":"2026-08-06T15:47:24.783200Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03168","last_updated":"2025-02-28T14:39:17Z","snapshot_observed_at":"2026-08-08T10:55:01.205999Z","submitted_at":"2024-07-03T14:41:39Z","title":"LivePortrait: Efficient Portrait Animation with Stitching and Retargeting Control","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.03168","snapshot_observed_at":"2026-08-06T15:47:24.407215Z","title":"Liveportrait: Efficient portrait animation with stitching and retargeting control,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.407215Z"},"links":{"cited_paper":"/paper/2407.03168","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:c1d1c9ae25f754fa7949e98677a5d783f1575db385e782e3bd27a0ff5aa821a9","observation_id":"f9f70bdc-493b-4841-bde0-aad9efee4fdc","resolution":{"observed_at":"2026-08-06T15:47:24.407215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T15:47:24.410707Z","title":"Cotracker: It is better to track together,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.410707Z"},"links":{"citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:e3dbe1184f868914e86e3a6147d3cfb1f96dc6586da71b6b2de249fa903c9e4a","observation_id":"66ff15d7-1ca2-461d-9df1-3c72c1a3ed32","resolution":{"observed_at":"2026-08-06T15:47:24.410707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.06072","last_updated":"2025-03-26T08:33:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-12T11:47:11Z","title":"CogVideoX: Text-to-Video Diffusion Models with An Expert Transformer","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.06072","snapshot_observed_at":"2026-08-06T15:47:24.413853Z","title":"Cogvideox: Text-to-video diffusion models with an expert transformer,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-06T15:47:24.413853Z"},"links":{"cited_paper":"/paper/2408.06072","citing_paper":"/paper/2507.15064"},"observation_digest":"sha256:9ef4521d7292953bb8b692925a07b7a8e58361a1fc97edee5982458eb945e18d","observation_id":"80ece95e-a724-490f-8ab1-ca3c62170032","resolution":{"observed_at":"2026-08-06T15:47:24.413853Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.15064","last_updated":"2025-07-20T17:59:26Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T09:47:36.038947Z","submitted_at":"2025-07-20T17:59:26Z","title":"StableAnimator++: Overcoming Pose Misalignment and Face Distortion for Human Image Animation"},"reference_resolution":{"displayed":73,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":35,"verified_exact":2,"verified_fuzzy":36},"total_outbound_references":73},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 73 of 73 outbound references and 2 inbound Pith citation observations for arXiv:2507.15064."}