{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2026:62RBERZYJCR4O3E7DLPJ662LHD","short_pith_number":"pith:62RBERZY","canonical_record":{"source":{"id":"2606.29986","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AR","submitted_at":"2026-06-29T09:00:38Z","cross_cats_sorted":["cs.DC"],"title_canon_sha256":"6c19495373dc8ccc630ece9e607f6240abefe0d1e3d42a7fc2c00a7217e2e28c","abstract_canon_sha256":"4031bd5abf64c40c68cd973d9e45b5ce1b6852a7d5ad20a1c6163c68d9c4d4be"},"schema_version":"1.0"},"canonical_sha256":"f6a212473848a3c76c9f1ade9f7b4b38f2e154506e65cb16ae36285fc18217ec","source":{"kind":"arxiv","id":"2606.29986","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2606.29986","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"arxiv_version","alias_value":"2606.29986v1","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2606.29986","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"pith_short_12","alias_value":"62RBERZYJCR4","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"pith_short_16","alias_value":"62RBERZYJCR4O3E7","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"pith_short_8","alias_value":"62RBERZY","created_at":"2026-06-30T02:17:44Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2026:62RBERZYJCR4O3E7DLPJ662LHD","target":"record","payload":{"canonical_record":{"source":{"id":"2606.29986","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AR","submitted_at":"2026-06-29T09:00:38Z","cross_cats_sorted":["cs.DC"],"title_canon_sha256":"6c19495373dc8ccc630ece9e607f6240abefe0d1e3d42a7fc2c00a7217e2e28c","abstract_canon_sha256":"4031bd5abf64c40c68cd973d9e45b5ce1b6852a7d5ad20a1c6163c68d9c4d4be"},"schema_version":"1.0"},"canonical_sha256":"f6a212473848a3c76c9f1ade9f7b4b38f2e154506e65cb16ae36285fc18217ec","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-06-30T02:17:44.214935Z","signature_b64":"uTAU/FZgnHiYv5cwNGrOtSNUdjgJrpdDY7vyEFZCjLDSzDPV45le/H21pgjI3PCriP0p6Skfjat3A4WsRBOiCg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"f6a212473848a3c76c9f1ade9f7b4b38f2e154506e65cb16ae36285fc18217ec","last_reissued_at":"2026-06-30T02:17:44.214348Z","signature_status":"signed_v1","first_computed_at":"2026-06-30T02:17:44.214348Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2606.29986","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-06-30T02:17:44Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"BSUHagQdREiZX/ggjg9km0UeM1ZrwhVqIwcIGWjrDO5sdororDH+9/9MXPShkDqDHIWqRtl02nHXo53SaIRDDQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-30T12:55:52.012243Z"},"content_sha256":"18f041168f7e7900b0f69d5e64cad697de72c75adeb9ae6a2d0fe5aaeccd0c5f","schema_version":"1.0","event_id":"sha256:18f041168f7e7900b0f69d5e64cad697de72c75adeb9ae6a2d0fe5aaeccd0c5f"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2026:62RBERZYJCR4O3E7DLPJ662LHD","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"HBM Is Not All You Need: Efficient Disaggregated LLM Serving across Memory-heterogeneous Accelerators","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.DC"],"primary_cat":"cs.AR","authors_text":"James Yen, Mingyuan Xia, Yun Wang, Zhengwei Qi, Zhixiang Wei","submitted_at":"2026-06-29T09:00:38Z","abstract_excerpt":"LLM inference comprises a compute-bound prefill phase and a memory-bound decode phase, and recent systems disaggregate them onto separate hardware. Yet today's datacenter GPUs rely on costly HBM whose bandwidth sits almost entirely idle during prefill. LLM serving across memory-heterogeneous accelerators (MemHA) pairs GDDR-based accelerators for prefill with HBM-based GPUs for decode, promising lower cost without sacrificing performance. Pushed to its most economical form, MemHA serving is inherently cross-vendor, since the best-suited chip for each phase may come from a different vendor. This"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2606.29986","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2606.29986/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-06-30T02:17:44Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"FBpqqYn0m47CZO8pyMl8QZFJCiiV2xAX6haaYHuWj6blAGY78YyKjogZXntFD/lTmUTSfG4ZtpxofHSVtsetDQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-30T12:55:52.012631Z"},"content_sha256":"7e24357c63fa5de1cc812ccf24052e5f7f8d56d26018a0012a77bcee4432c25e","schema_version":"1.0","event_id":"sha256:7e24357c63fa5de1cc812ccf24052e5f7f8d56d26018a0012a77bcee4432c25e"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/62RBERZYJCR4O3E7DLPJ662LHD/bundle.json","state_url":"https://pith.science/pith/62RBERZYJCR4O3E7DLPJ662LHD/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/62RBERZYJCR4O3E7DLPJ662LHD/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-06-30T12:55:52Z","links":{"resolver":"https://pith.science/pith/62RBERZYJCR4O3E7DLPJ662LHD","bundle":"https://pith.science/pith/62RBERZYJCR4O3E7DLPJ662LHD/bundle.json","state":"https://pith.science/pith/62RBERZYJCR4O3E7DLPJ662LHD/state.json","well_known_bundle":"https://pith.science/.well-known/pith/62RBERZYJCR4O3E7DLPJ662LHD/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2026:62RBERZYJCR4O3E7DLPJ662LHD","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"4031bd5abf64c40c68cd973d9e45b5ce1b6852a7d5ad20a1c6163c68d9c4d4be","cross_cats_sorted":["cs.DC"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AR","submitted_at":"2026-06-29T09:00:38Z","title_canon_sha256":"6c19495373dc8ccc630ece9e607f6240abefe0d1e3d42a7fc2c00a7217e2e28c"},"schema_version":"1.0","source":{"id":"2606.29986","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2606.29986","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"arxiv_version","alias_value":"2606.29986v1","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2606.29986","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"pith_short_12","alias_value":"62RBERZYJCR4","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"pith_short_16","alias_value":"62RBERZYJCR4O3E7","created_at":"2026-06-30T02:17:44Z"},{"alias_kind":"pith_short_8","alias_value":"62RBERZY","created_at":"2026-06-30T02:17:44Z"}],"graph_snapshots":[{"event_id":"sha256:7e24357c63fa5de1cc812ccf24052e5f7f8d56d26018a0012a77bcee4432c25e","target":"graph","created_at":"2026-06-30T02:17:44Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2606.29986/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"LLM inference comprises a compute-bound prefill phase and a memory-bound decode phase, and recent systems disaggregate them onto separate hardware. Yet today's datacenter GPUs rely on costly HBM whose bandwidth sits almost entirely idle during prefill. LLM serving across memory-heterogeneous accelerators (MemHA) pairs GDDR-based accelerators for prefill with HBM-based GPUs for decode, promising lower cost without sacrificing performance. Pushed to its most economical form, MemHA serving is inherently cross-vendor, since the best-suited chip for each phase may come from a different vendor. This","authors_text":"James Yen, Mingyuan Xia, Yun Wang, Zhengwei Qi, Zhixiang Wei","cross_cats":["cs.DC"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AR","submitted_at":"2026-06-29T09:00:38Z","title":"HBM Is Not All You Need: Efficient Disaggregated LLM Serving across Memory-heterogeneous Accelerators"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2606.29986","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:18f041168f7e7900b0f69d5e64cad697de72c75adeb9ae6a2d0fe5aaeccd0c5f","target":"record","created_at":"2026-06-30T02:17:44Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"4031bd5abf64c40c68cd973d9e45b5ce1b6852a7d5ad20a1c6163c68d9c4d4be","cross_cats_sorted":["cs.DC"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AR","submitted_at":"2026-06-29T09:00:38Z","title_canon_sha256":"6c19495373dc8ccc630ece9e607f6240abefe0d1e3d42a7fc2c00a7217e2e28c"},"schema_version":"1.0","source":{"id":"2606.29986","kind":"arxiv","version":1}},"canonical_sha256":"f6a212473848a3c76c9f1ade9f7b4b38f2e154506e65cb16ae36285fc18217ec","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"f6a212473848a3c76c9f1ade9f7b4b38f2e154506e65cb16ae36285fc18217ec","first_computed_at":"2026-06-30T02:17:44.214348Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-06-30T02:17:44.214348Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"uTAU/FZgnHiYv5cwNGrOtSNUdjgJrpdDY7vyEFZCjLDSzDPV45le/H21pgjI3PCriP0p6Skfjat3A4WsRBOiCg==","signature_status":"signed_v1","signed_at":"2026-06-30T02:17:44.214935Z","signed_message":"canonical_sha256_bytes"},"source_id":"2606.29986","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:18f041168f7e7900b0f69d5e64cad697de72c75adeb9ae6a2d0fe5aaeccd0c5f","sha256:7e24357c63fa5de1cc812ccf24052e5f7f8d56d26018a0012a77bcee4432c25e"],"state_sha256":"b410bc5e69d88e49b4c6412a29b5e0b8df6e0297932f03befcd899a35b13b865"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"u2jP1AL27peQfbMAkNifbta71Eb0rTZETVpkm+uqlK9/ksxxMjHENeBBYxONFkDiPUN6wjmbpi03EH4d+5auAA==","signed_message":"bundle_sha256_bytes","signed_at":"2026-06-30T12:55:52.014569Z","bundle_sha256":"9af6c9447ca334a2befb6cba5554c456c22c88b467c6d88d965bd2cc114da808"}}