{"as_of":"2026-08-09T07:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:14485b7116395c013d3e8ec6b2b096b22d2a49c46d0729c24ee5e0df4822ddf0","coverage":[{"denominator":22,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T16:57:42.654116Z","state":"measured"},{"denominator":25,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":25,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T14:42:22.214222Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-29T19:03:51.552172Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06061","snapshot_observed_at":"2026-08-05T14:42:22.214222Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.21016","last_updated":"2025-08-28T17:18:31Z","snapshot_observed_at":"2026-08-06T03:08:28.577840Z","submitted_at":"2025-08-28T17:18:31Z","title":"Inference-Time Alignment Control for Diffusion Models with Reinforcement Learning Guidance","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-05T14:42:22.214222Z"},"links":{"cited_paper":"/paper/2502.06061","citing_paper":"/paper/2508.21016"},"observation_digest":"sha256:cbb1026c9aaab7c99ffb7051b6129d0bbc3038884b4061ddebb8c6ad2689bb90","observation_id":"5b73787a-0011-4fb3-8fcf-a88a0e1b9001","resolution":{"observed_at":"2026-08-05T14:42:22.214222Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06061","snapshot_observed_at":"2026-08-04T10:44:08.654870Z","title":"9 Justin Fu, Katie Luo, and Sergey Levine","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.09222","last_updated":"2026-06-01T11:49:06Z","snapshot_observed_at":"2026-08-07T20:08:20.794829Z","submitted_at":"2025-10-10T10:08:10Z","title":"FM-IRL: Flow-Matching for Reward Modeling and Policy Regularization in Reinforcement Learning","version":3},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-04T10:44:08.654870Z"},"links":{"cited_paper":"/paper/2502.06061","citing_paper":"/paper/2510.09222"},"observation_digest":"sha256:ad33260d599c34576aa83069c6d252a940fcc64e31c431711d93c5f89cc49fbb","observation_id":"b909be6c-2528-444c-96c2-60f7217c6854","resolution":{"observed_at":"2026-08-04T10:44:08.654870Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"cited_work":{"arxiv_id":"2502.06061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06061","snapshot_observed_at":"2026-06-29T19:03:51.552172Z","title":"Fu, J., Luo, K., and Levine, S","venue":null,"work_id":"b06c69ea-5595-4dbb-b757-664c444a8c52","year":null},"citing_paper":{"arxiv_id":"2605.27095","last_updated":"2026-06-01T11:46:42Z","snapshot_observed_at":"2026-08-06T05:19:12.135756Z","submitted_at":"2026-05-26T14:38:03Z","title":"Adversarial Dual On-Policy Distillation from Expressive Teacher","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-29T18:55:23.777484Z"},"links":{"cited_paper":"/paper/2502.06061","citing_paper":"/paper/2605.27095"},"observation_digest":"sha256:fdb0a07a25096e37fa68ef88e3257b757c5a6e8af8e7dba7f5543ca649ccabb6","observation_id":"5a44649f-859a-42bf-97b9-d208d985208c","resolution":{"observed_at":"2026-06-29T19:03:51.553716Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.06061/citation-record","integrity":"/paper/2502.06061/integrity","json":"/paper/2502.06061/citation-record.json","paper":"/paper/2502.06061"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:43.059667Z","title":"left of”, “on top of","venue":null,"work_id":"24828b5b-bd22-4ccc-9eae-603ae3d92270","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.556146Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:2330d8d868a5bf02d181f9dc448b741739386fc483956415f377e206b885d84a","observation_id":"add414f3-fad3-488d-a656-799906ed1732","resolution":{"observed_at":"2026-08-08T16:57:43.064886Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:43.042229Z","title":"To address this challenge, we introduce W2 regularization, which effectively prevents over-optimization and policy collapse (Lemma 1)","venue":null,"work_id":"ee133764-5445-400f-94ed-db84373a4edf","year":2023},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.562201Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:fb67cbb6ff1a95ebca4f3739fb1b97ce8a0962d27fb26cf52d6d417539bdffa0","observation_id":"4047a014-4084-48fa-bfc8-b790c9635c23","resolution":{"observed_at":"2026-08-08T16:57:43.048103Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:43.076748Z","title":"left of”, “on top of","venue":null,"work_id":"3c798d95-7085-4710-a21a-6799d4f4cdcc","year":2023},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.550430Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:8e27101f4594a9c9a7b12ce41ca968dcfa0254415b6f87fca888652bdc7aa23d","observation_id":"c6adba18-3860-4df4-88a8-87d193498cc8","resolution":{"observed_at":"2026-08-08T16:57:43.081852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:43.022742Z","title":"a cat in the sky","venue":null,"work_id":"06608cae-9b4f-45ea-824d-e17be32aaa65","year":2021},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.567605Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:f6fdda9699ad9f51bd745846af5c4a331cf76bcd2714d668cadab6388c2412cf","observation_id":"8848c0cf-3223-4696-b1c9-b9e4d71afab0","resolution":{"observed_at":"2026-08-08T16:57:43.029134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.998595Z","title":"a train on top of a surfboard,","venue":null,"work_id":"874a2c32-7911-44f3-8f7c-58b8bae54f35","year":2023},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.573522Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:9cc0dca79c006d04c005975025192c86c9a0915a3d697b7a0a659a4afb20b264","observation_id":"f7cdfa02-8a77-4434-8a4c-3dc8e23792bd","resolution":{"observed_at":"2026-08-08T16:57:43.006766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.978905Z","title":"3) Bottom row: Alpha Clip reward showcases consis- tent performance even with different text-image alignment rewards, validating the reward-agnostic nature/property of our approach","venue":null,"work_id":"7d3543f6-ef77-4581-8e92-f9eb8a98f82f","year":2023},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.578823Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:838e43887c57a5605f6128628d877736d73134c2b4fdda4dcc525f5da960ca3d","observation_id":"b9e49203-b038-4437-af6c-c5399e991a44","resolution":{"observed_at":"2026-08-08T16:57:42.984986Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.962786Z","title":null,"venue":null,"work_id":"4146da3e-d6b6-4b76-84cd-e52ef7485928","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.584564Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:43858c6d30493e60c747ce9c951831fd4b13aa7c682effdfef2fbcafd59ab793","observation_id":"7e8aa6a2-b43a-4b18-8778-2e6f50948a7e","resolution":{"observed_at":"2026-08-08T16:57:42.968182Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.945415Z","title":null,"venue":null,"work_id":"48f4c911-9c81-4cd5-b32e-ce9963262e47","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.589732Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:8c99db5b2f14c28d4982bb2eadf39fa66090132aef89810ea2a01297a70608a5","observation_id":"9f8a3694-bc36-4e72-90d6-3ead5fe9bcbd","resolution":{"observed_at":"2026-08-08T16:57:42.950919Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.928642Z","title":"on top\",","venue":null,"work_id":"decfd566-7930-4721-8287-1d00bb2ed555","year":2025},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.595160Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:9a4a84b2461d4c5deba5f27b92066db9c4c3aa891ff999528423c61a3dfdfcf1","observation_id":"d28eda71-a0da-4321-847d-cbecdb6d56fc","resolution":{"observed_at":"2026-08-08T16:57:42.933892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.873474Z","title":"Given that w (x1) > 0 for all x1 ∈ Xand attains its maximum at x∗ 1, we define: ϵ (x1) = w (x1) w (x∗","venue":null,"work_id":"546af60a-9780-4e07-ae85-86f1867c01db","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.612598Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:8a20397ddc360c48c50bd849ee40ab8adc3b26c3f280f9e388f600fbf26f63fc","observation_id":"4066ae4b-8442-4f8a-aa12-ffbb3ade7660","resolution":{"observed_at":"2026-08-08T16:57:42.879540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.855948Z","title":null,"venue":null,"work_id":"aab42bd8-50f7-4d84-8334-633176888ff6","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.617699Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:22834b497f7c8da70e51e8577810231408d88b19d7dd5b9c0f9db9da663dca4c","observation_id":"a396bb3a-3861-4a00-84dc-28a8157f6d16","resolution":{"observed_at":"2026-08-08T16:57:42.861835Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.838920Z","title":"Then, we can rewrite qN θ (x1) using ϵ(x1) as: qN θ (x1) = [w (x∗ 1)]N ϵ (x1)N q (x1) ZN (73) Then for x1 ̸= x∗ 1, we can have: ϵ(x1)N → 0 as N → ∞since ϵ(x1) < 1","venue":null,"work_id":"f94b4e7a-5e33-4d06-9d22-46fa504a1dda","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.623385Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:3c363e3b21456bda1390cb6fed8120d14241720327a12b039b1e5716e8a066d1","observation_id":"440d2089-c959-4adc-bfa9-168ca71fe817","resolution":{"observed_at":"2026-08-08T16:57:42.844221Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.820886Z","title":"And we can have the normalization constant as follows: ZN = Z X w (x1)N q (x1) dx1 = [w (x∗ 1)]N Z X ϵ (x1)N q (x1) dx1","venue":null,"work_id":"1f889927-937c-446d-9bd4-c729ea827040","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.628335Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:4ae2e05fc2dc47b52b13cfdefc8107a7c3126d23cf74ddd786c516773b9b9669","observation_id":"12896a55-1a87-4433-b125-2cd1f4707661","resolution":{"observed_at":"2026-08-08T16:57:42.826271Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.801626Z","title":"Then, we can have the limit behavior: For x1 ̸= x∗ 1 : qN θ (x1) = [w (x∗ 1)]N ϵ (x1)N q (x1) [w (x∗ 1)]N q (x∗ 1) = ϵ (x1)N q (x1) q (x∗","venue":null,"work_id":"de9fc354-693a-4f30-a052-01552d7d628e","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.633170Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:8d77ee69ff0504af9b612ee13066a5019b7cbeed13f12c1e1237b987001c5340","observation_id":"630b8d02-6064-4967-aa5d-b271f77edd9c","resolution":{"observed_at":"2026-08-08T16:57:42.807617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.784427Z","title":"(75) For x1 = x∗ 1 : qN θ (x∗","venue":null,"work_id":"125215b6-26d6-44b6-aebc-0f5e55009ba9","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.637958Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:72da007e5107af5c9051c725e4a2042c2d2cb51d8ab5d60d6ca071b1dd438abf","observation_id":"bfea99bd-4bc6-456c-aebf-059103fa33c8","resolution":{"observed_at":"2026-08-08T16:57:42.789504Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.767585Z","title":null,"venue":null,"work_id":"49aad384-18cc-4e12-97d0-630f177a89a7","year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.642634Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:1124cc7f8d0f3e9b5a563d5735135edd2ab04b659106fabd75a6cab2ee3896b4","observation_id":"68dbfe68-f857-4469-8984-b3681be2010d","resolution":{"observed_at":"2026-08-08T16:57:42.772687Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.750893Z","title":"cat\") = pclip (x,","venue":null,"work_id":"e10a652f-1ef3-4551-9e75-3c5984b92ac3","year":2024},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.647660Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:ccfab8b29f81e9a4ad6d5f7e0808cb49063a9b8e7724df88b2e6a9256c8623a9","observation_id":"96598b70-5c5e-4947-b81d-82bee4ab30e4","resolution":{"observed_at":"2026-08-08T16:57:42.756262Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.732000Z","title":null,"venue":null,"work_id":"9ac7dc95-90f5-4d72-b3e2-42e179145d39","year":2024},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.654116Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:54a81d0c247dbed254e514f0cdda4ea75086bd882782fd87f1e5148470a5999e","observation_id":"e95af9ff-2d57-4371-b120-ad8d9cf1bf3e","resolution":{"observed_at":"2026-08-08T16:57:42.738323Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.891252Z","title":"In this paper, we introduce two methods to handle the Overoptimization and ease the mode collapse risk in online RW-CFM algorithms","venue":null,"work_id":"53a853fe-caf3-4cbe-b11d-de33e74869c7","year":2024},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":2014,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.606891Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:397a450b504db9a4bd95bf0e18ddc264e0bacedcec2e00376d5175fa22d8adc5","observation_id":"e5a1e71e-039e-4ec2-a381-5cf4f4b03a79","resolution":{"observed_at":"2026-08-08T16:57:42.897400Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1612.00796","last_updated":"2017-01-25T13:01:51Z","snapshot_observed_at":"2026-08-04T14:09:02.234596Z","submitted_at":"2016-12-02T19:18:37Z","title":"Overcoming catastrophic forgetting in neural networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1612.00796","snapshot_observed_at":"2026-08-08T16:57:42.537531Z","title":"James Kirkpatrick, Razvan Pascanu, Neil C","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.537531Z"},"links":{"cited_paper":"/paper/1612.00796","citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:074bb0d77689e297c444a25f67d9d7ad5f1236e06fefb47cb280e7860cb5db66","observation_id":"b87d7f19-3b44-4ef1-834d-3d443c4d9783","resolution":{"observed_at":"2026-08-08T16:57:42.537531Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:57:42.910358Z","title":null,"venue":null,"work_id":"a57b4180-2a60-4f70-90d7-d77713571cad","year":2024},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.600862Z"},"links":{"citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:7d688339c9a096384b78e420e564233bea136ad2adf4ac5b745b86392ddf59b1","observation_id":"c132e1f4-b409-48e7-a311-447bfe713566","resolution":{"observed_at":"2026-08-08T16:57:42.916713Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.08448","last_updated":"2024-05-14T09:12:30Z","snapshot_observed_at":"2026-08-05T00:57:12.421643Z","submitted_at":"2024-05-14T09:12:30Z","title":"Understanding the performance gap between online and offline alignment algorithms","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.08448","snapshot_observed_at":"2026-08-08T16:57:42.544295Z","title":"URL https://doi.org/10.48550/arXiv","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-08T16:57:42.544295Z"},"links":{"cited_paper":"/paper/2405.08448","citing_paper":"/paper/2502.06061"},"observation_digest":"sha256:ecce7f1923f7e22e5543a36680523a4d194d2ac0374d390d71b306e9684c27bf","observation_id":"1435c790-34aa-44ff-b440-baa7ed828949","resolution":{"observed_at":"2026-08-08T16:57:42.544295Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.06061","last_updated":"2025-02-09T22:45:15Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-08T16:50:18.012795Z","submitted_at":"2025-02-09T22:45:15Z","title":"Online Reward-Weighted Fine-Tuning of Flow Matching with Wasserstein Regularization"},"reference_resolution":{"displayed":22,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":1,"unresolved":7,"verified_exact":0,"verified_fuzzy":14},"total_outbound_references":22},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 22 of 22 outbound references and 3 inbound Pith citation observations for arXiv:2502.06061."}