{"as_of":"2026-08-21T13:12:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:04c6201909a52d1fcedfd895ee87ebd99b8c9a9b608d112f46420a505cbac623","coverage":[{"denominator":15,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T23:32:58.224956Z","state":"measured"},{"denominator":18,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":18,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-30T16:20:17.727753Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T10:39:45.633252Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"cited_work":{"arxiv_id":"2505.04560","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.04560","snapshot_observed_at":"2026-07-04T10:39:45.633252Z","title":"Abkd: Pursuing a proper allocation of the probability mass in knowledge distillation viaα-β- divergence,","venue":null,"work_id":"c9731a53-23a3-4f12-b961-a6a90ba0dc67","year":2025},"citing_paper":{"arxiv_id":"2605.28869","last_updated":"2026-05-22T08:22:31Z","snapshot_observed_at":"2026-08-14T09:14:37.658084Z","submitted_at":"2026-05-22T08:22:31Z","title":"Balancing Multimodal Learning through Label Space Reshaping","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-30T16:20:17.727753Z"},"links":{"cited_paper":"/paper/2505.04560","citing_paper":"/paper/2605.28869"},"observation_digest":"sha256:3fc7a9e34b579091a586d696ea17b3d040fc0375a26a7e52ddea07a37d04d066","observation_id":"1b28718b-df86-446b-9849-7abe7f5aa53e","resolution":{"observed_at":"2026-06-30T16:24:55.682605Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"cited_work":{"arxiv_id":"2505.04560","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.04560","snapshot_observed_at":"2026-07-04T10:39:45.633252Z","title":"Abkd: Pursuing a proper allocation of the probability mass in knowledge distillation viaα-β- divergence,","venue":null,"work_id":"c9731a53-23a3-4f12-b961-a6a90ba0dc67","year":2025},"citing_paper":{"arxiv_id":"2606.03091","last_updated":"2026-06-04T07:55:34Z","snapshot_observed_at":"2026-08-04T16:58:01.578812Z","submitted_at":"2026-06-02T03:26:25Z","title":"BAHSD: Bridging the Long-tail Gap via Adaptive Distillation in Black-box Sequential Recommendation","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-28T08:41:42.421943Z"},"links":{"cited_paper":"/paper/2505.04560","citing_paper":"/paper/2606.03091"},"observation_digest":"sha256:e91772d5479d76fa11aee2cf621c16a9f744f475e120c304f8459094ff6193a9","observation_id":"1a7d23ba-bd5a-4dc7-97bb-009202c9d384","resolution":{"observed_at":"2026-07-02T04:56:39.148582Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"cited_work":{"arxiv_id":"2505.04560","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.04560","snapshot_observed_at":"2026-07-04T10:39:45.633252Z","title":"Abkd: Pursuing a proper allocation of the probability mass in knowledge distillation viaα-β- divergence,","venue":null,"work_id":"c9731a53-23a3-4f12-b961-a6a90ba0dc67","year":2025},"citing_paper":{"arxiv_id":"2606.23124","last_updated":"2026-06-22T10:09:52Z","snapshot_observed_at":"2026-08-18T00:03:21.338695Z","submitted_at":"2026-06-22T10:09:52Z","title":"PRIDE: Privileged Information-enhanced Distillation for Empathetic Dialogue Generation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-26T08:37:10.149519Z"},"links":{"cited_paper":"/paper/2505.04560","citing_paper":"/paper/2606.23124"},"observation_digest":"sha256:9d7803dacc20e54040c5c679e3b906124ae7331e57f775a165cbc31b144fd951","observation_id":"fdb004bb-8a7a-4662-87d3-2d34822da537","resolution":{"observed_at":"2026-07-04T10:39:45.635129Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.04560/citation-record","integrity":"/paper/2505.04560/integrity","json":"/paper/2505.04560/citation-record.json","paper":"/paper/2505.04560"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.465833Z","title":"Ify1 andy2 are such thatδ1 <q t(y1) =q t(y2)≤p(y 1) (whereδ1 >0 ), andp(y1)≥p(y 2) +ζ , it holds that∆α1 t (y1,y 2)≥∆ α2 t (y1,y 2)","venue":null,"work_id":"9e6cd627-1e77-481e-b7cd-0d48e70b1bea","year":2024},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.166154Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:3b629a1f92b40fb14f873a820d83064b34a0a48d38b7457380889231f47d3783","observation_id":"e8eca142-0432-4e5f-88d3-3695167bdccd","resolution":{"observed_at":"2026-08-15T23:32:58.470546Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.434770Z","title":"If y1 andy2 are such that p(y1)< qt(y1) =q t(y2)≤1−δ 2, but p(y1)≥p(y 2) +ζ , then, ∆r t (y1,y 2)> ∆f t (y1,y 2)","venue":null,"work_id":"0aa159d4-9ac1-4e8b-b412-cf443cd24db5","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.176549Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:30cf038e74767ac77d5a69b5ca48311c384cb05e0a48494fea30888b36735768","observation_id":"3459fa8b-950d-472b-add9-9ec75fc2c870","resolution":{"observed_at":"2026-08-15T23:32:58.439676Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.480322Z","title":"The proof is in App","venue":null,"work_id":"912b3d33-4525-425c-9bff-f9b0c6be49a5","year":2024},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.160308Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:9bfb9636bb94b0f0b27d41bf09b81a687c3ef130f4ebca417a17766aa568c33e","observation_id":"40a9f78b-4124-4f77-b89b-79ef193912ba","resolution":{"observed_at":"2026-08-15T23:32:58.484984Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.402777Z","title":"Ify1 andy2 are such thatqt(y2) +ζ≤q t(y1)≤1−δ 2, andc0·qt(y2)<p(y 1) =p(y 2)<c 1·qt(y1), wherec0 and c1 are constants withc 0 >1andc 1 <1, then,∆ r t (y1,y 2)>∆ f t (y1,y 2)","venue":null,"work_id":"7afacb4b-4aa0-4b82-9ad2-acc8c8de42a3","year":2024},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.187085Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:eee7baf0c9b627214cfa81de9f86e09ff24d56da176c19b051616ffb97cf7e5f","observation_id":"15827d17-3b23-4519-8f41-a73cd3f0edf7","resolution":{"observed_at":"2026-08-15T23:32:58.408040Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.450334Z","title":"Ify1 andy2 are such that δ1 < qt(y1) =q t(y2)< p(y 1) (whereδ1 >0,δ 2 >0 ), but p(y1)≥p(y 2) +ζ , then, ∆r t (y1,y 2)>∆ f t (y1,y 2)","venue":null,"work_id":"f4ad5824-acf8-424e-82ac-596dd4df7394","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.171514Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:b976fd9c3c3de5c8abaee1ce2c21714e025ab583e6bb2749ab2942bc9860eae7","observation_id":"06e7c6dd-2d6e-4242-82fe-da6b0055050d","resolution":{"observed_at":"2026-08-15T23:32:58.455258Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.418239Z","title":"If y1 andy2 are such that qt(y2) +ζ≤q t(y1)≤1−δ 2, and p(y1) =p(y 2)> c 0·qt(y1), wherec0 is a positive constant>1, then,∆ r t (y1,y 2)>∆ f t (y1,y 2)","venue":null,"work_id":"606b2d23-008f-40fb-874d-cd869056afed","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.181774Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:72c2c2953d41a24f9a97ed9a869f03f0b3aa75f96c8d284cd02470770a1d3f8e","observation_id":"993e4e09-b336-42d0-904d-2d9a9e0e616d","resolution":{"observed_at":"2026-08-15T23:32:58.424195Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.387604Z","title":"(36) Simplifying: ∂q(i) ∂f(i) =q(i)(1−q(i)).(37)","venue":null,"work_id":"524d9a6b-d667-48e2-a93d-361a0311bff9","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.192961Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:17b00cd8d0a5ef71e1dda5339068482a2df847cdacb066ad39273e5a35a4872b","observation_id":"6c9bd6a8-c1d2-427a-a8f0-f75fc1e60589","resolution":{"observed_at":"2026-08-15T23:32:58.392427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.369957Z","title":"q(y) 1−α (p(y)α−q(y) α) +q(y) X k q(k) 1−α(q(k)α−p(k)α) !#! .(47) Taking the logarithm on both sides, we get: logqα t+1(y) qt(y) =η","venue":null,"work_id":"4d7812dd-0bb7-498b-ba4a-e7e0fdde4169","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.198307Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:6afc753c95e180926700744734e95de2865cf5eff472468c526a9029e961f7c7","observation_id":"51cee0a9-0c3c-44c1-afcb-d633cab0a4d2","resolution":{"observed_at":"2026-08-15T23:32:58.376293Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.353727Z","title":null,"venue":null,"work_id":"25039d66-a69b-4644-8423-151205de13b5","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.203821Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:e044b4441f7378aedd0a9f6062b96303aa5ed58d5b22b8381b7664de10489e46","observation_id":"de502133-de94-428d-88da-c3bd310a0d66","resolution":{"observed_at":"2026-08-15T23:32:58.358555Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.338026Z","title":"Moreover, for allp(y1)> c0, it holds that f(p(y 1))<0","venue":null,"work_id":"fd0ed837-c2ad-4b07-84a0-4ede87a768d5","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.209075Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:d874a31642308a0de11d28311883c16515eb9ae9261dd0f5980afab0c53be43a","observation_id":"39e346fc-a2ca-4425-b15f-080d5475d328","resolution":{"observed_at":"2026-08-15T23:32:58.343184Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.321387Z","title":"X k q(k) 1−α (q(k)α−p(k)α) # . (121) Consideringf(α)≜∆ α(y1,y 2)/η, taking the derivative with respect toαyields: f′(α) = 1 α2","venue":null,"work_id":"499848fd-6c8c-49a4-9d3e-f7d49edd8103","year":null},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.214000Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:e7e7b00b5415f69ecd978a4b572d79f32f8bd7efd37fdb0880b2a77f96b8be46","observation_id":"0438db78-f617-420a-8792-1ac5e1ef2d5a","resolution":{"observed_at":"2026-08-15T23:32:58.327081Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.307588Z","title":"Proof.First, we prove Case 1","venue":null,"work_id":"dad4ccd5-e886-454f-867b-46807f8950cb","year":2024},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.218923Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:d68f2ea79807cdda35cabe883a768d50c83f6ffdaf3846f43859aaf5fe30d63d","observation_id":"5d60ce0a-2813-4cb4-b209-52fc067ae3c2","resolution":{"observed_at":"2026-08-15T23:32:58.311644Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15190","last_updated":"2023-07-27T20:39:06Z","snapshot_observed_at":"2026-08-19T20:26:32.298786Z","submitted_at":"2023-07-27T20:39:06Z","title":"f-Divergence Minimization for Sequence-Level Knowledge Distillation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15190","snapshot_observed_at":"2026-08-15T23:32:58.154253Z","title":"emnlp-main.340","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":340,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.154253Z"},"links":{"cited_paper":"/paper/2307.15190","citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:9aa34abad8019b199feca58537a171ece392403d81855d6d715d8fdadb2d013e","observation_id":"a09273f2-c889-4b91-a6b9-5d9cd8f94e3f","resolution":{"observed_at":"2026-08-15T23:32:58.154253Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1606.07947","last_updated":"2016-09-22T01:17:12Z","snapshot_observed_at":"2026-08-14T21:51:10.458695Z","submitted_at":"2016-06-25T18:16:39Z","title":"Sequence-Level Knowledge Distillation","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.07947","snapshot_observed_at":"2026-08-15T23:32:58.148277Z","title":"findings-emnlp.364/","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":364,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.148277Z"},"links":{"cited_paper":"/paper/1606.07947","citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:9056415c65bfa73ae14edbc21123f7610fda4faf069c8021260cbc236cb24264","observation_id":"839dc3fe-7bde-4790-bc74-54c7b8c0fe46","resolution":{"observed_at":"2026-08-15T23:32:58.148277Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:32:58.291151Z","title":"a photo of a{classname}","venue":null,"work_id":"55612677-6013-4366-805b-b578e835c1b9","year":2022},"citing_paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence","version":3},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-15T23:32:58.224956Z"},"links":{"citing_paper":"/paper/2505.04560"},"observation_digest":"sha256:904edb053faa2bf48c48fd0ca44c8583e1c9d24e38526b30059176bd3593aa15","observation_id":"f7c65cdb-5602-425b-ae37-afdcda639d35","resolution":{"observed_at":"2026-08-15T23:32:58.296972Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.04560","last_updated":"2025-06-03T12:33:27Z","latest_version":3,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-20T16:46:00.785920Z","submitted_at":"2025-05-07T16:48:49Z","title":"ABKD: Pursuing a Proper Allocation of the Probability Mass in Knowledge Distillation via $\\alpha$-$\\beta$-Divergence"},"reference_resolution":{"displayed":15,"state_counts":{"malformed_identifier":2,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":2,"verified_exact":0,"verified_fuzzy":11},"total_outbound_references":15},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 15 of 15 outbound references and 3 inbound Pith citation observations for arXiv:2505.04560."}