{"as_of":"2026-08-19T21:23:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:82c722440d732891fe66adc721410eeff7e957f1bc6eb56865a81baa00702707","coverage":[{"denominator":47,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":47,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-14T15:42:46.891891Z","state":"measured"},{"denominator":47,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":47,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.09789/citation-record","integrity":"/paper/2607.09789/integrity","json":"/paper/2607.09789/citation-record.json","paper":"/paper/2607.09789"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Recent improvements of the particle and heavy ion transport code system—PHITS version 3.33","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:e9eae83933c5f9cbbe1edcc853d0187655270b9601827de388fd506d8d1423de","observation_id":"a2ec28d7-6270-4dd0-ad33-bc8e518c2b07","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"MCNP version 6.2 release notes","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:ba868917223ab94eb83ca01b021e2ee0ea832fa600f89bc7ceb2f2fa542182b1","observation_id":"02ce4a4b-15f0-4585-b5da-7486e31b0fec","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"GEANT4—a simulation toolkit","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:172d3a989aea23644718b6bdd146cd66a607f7c4eeaf3e810659d0fc6a1e6788","observation_id":"293cf11b-73e8-4b70-844b-e12dde132c5c","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"FLUKA: A multi-particle transport code (program version 2005)","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:60670045c058b5a1a8f957c5be7797b2e3a42d3f3c59d7791514a590c178fc53","observation_id":"9985ea9a-af1b-4075-89c8-2bea58df31a7","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-08-08T11:58:24.516369Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Evaluating large language models trained on code [Arxiv:2107.03374 [cs.lg]]; 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:0262a740f1f18673d6efda4710672b1cd9d3da2a914834153bdf4d15004e6a24","observation_id":"e73e0443-d4dd-457b-ba88-fb90ae8f41e3","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"SWE-bench: Can language models resolve real-world GitHub issues? In: The Twelfth International Conference on Learning Representations (ICLR); 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:eb2446eca49cb025dd3709ec1d5690fdb3b197468be22cb898fe6d530a191371","observation_id":"70ce1b01-f659-4acd-a1f2-7c2ae610a75a","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"DS-1000: A natural and reliable benchmark for data science code generation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:ad3c39bc5d0c5e226f0dd1768a9f9043c900b7fd7a22087aa2ebec707fc2bddb","observation_id":"a223f2aa-f381-4e21-a30c-4fe3f3136a0c","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Invited paper: VerilogEval: Evaluating large language models for Verilog code generation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:9413230078a816d3f4427cbf493e830526a6c7343ab7aa7a559edb0d1a996b9c","observation_id":"5369fd8b-2fd5-402d-bcdb-0abf8cf82d51","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"BigCodeBench: Benchmarking code generation with diversefunctioncallsandcomplexinstructions.In:TheThirteenthInternationalConference on Learning Representations (ICLR); 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:b33deb36e5a5a44efc93afe61fc4f1fd295c52e57e0f82f7ab688fd603e13531","observation_id":"3883d4ff-e314-4fe0-b238-3c68cd20ddf9","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"SciCode: A research coding benchmark curated by scientists","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:07b69735fca2fcb2b009d50a6e797f10fe0a80a90e4a68995250bd34394782d4","observation_id":"aa598f56-4dc0-4f3f-9aff-e2ffed2d9ebc","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Automating Monte Carlo simulations in nuclear engineering with domain knowledge-embedded large language model agents","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:706e202bae3d64b6f6fa283413d649ebc253520e4a2ef827ce943c89f979d8f7","observation_id":"3f6db513-5eb9-4a35-bcd9-d5a82776a6fe","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"A self-correcting multi-agent LLM framework for language-based physics simulation and explanation","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:1201ea8ba9ec92dd3cb25af55c3faf8500058c4aa79e29aa37898c5f0d0414bb","observation_id":"0eb98e7c-860f-4f57-8b52-c03e7e3daef8","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Evaluating the performance of large language 19 models for geometry and simulation file generation in physics-based simulations","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:f64a2bc0437906addf6e5596ea060ee8f098a9069826cda1d774f25e0556f802","observation_id":"9c8c1575-97e9-4691-aea3-c6f60217f7e0","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"OpenFOAMGPT: A retrieval-augmented large language model (LLM) agent for OpenFOAM-based computational fluid dynamics","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:e89d7359d2e5e7cea6c1fad266b1f78468800f05b101472185472d9bf1fc65df","observation_id":"ac238909-decc-4a99-a2c3-24d9646e5dd7","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21320","last_updated":"2024-08-07T04:34:11Z","snapshot_observed_at":"2026-08-18T18:42:16.318156Z","submitted_at":"2024-07-31T04:01:08Z","title":"MetaOpenFOAM: an LLM-based multi-agent framework for CFD","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21320","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"MetaOpenFOAM: An LLM-based multi-agent framework for CFD [Arxiv:2407.21320 [physics.flu-dyn]]; 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2407.21320","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:13f921744def17a77009063159f7a6a69b6deef738494df0a38f970e61dbe8c2","observation_id":"26279d89-94a5-495a-8f10-c49953b0e576","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.20374","last_updated":"2026-04-26T17:39:58Z","snapshot_observed_at":"2026-08-13T19:20:38.694867Z","submitted_at":"2025-09-19T22:21:26Z","title":"CFDLLMBench: A Benchmark Suite for Evaluating Large Language Models in Computational Fluid Dynamics","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.20374","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"CFDLLMBench: A benchmark suite for evaluating large language models in computational fluid dynamics [Arxiv:2509.20374 [cs.lg]]; 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2509.20374","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:a2b6fe27d839d135fb66dbd809237d500e0ccf31a62c679ae2df4216de999ff9","observation_id":"db51d044-e521-49b4-9d6e-8d9f7dd0fb6b","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"MooseAgent: A LLM based multi-agent framework for automating MOOSE simulation ; 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:d8cf3fb1626848c589a0b1dd74b8fb97d36a83205278d62d2066379278837773","observation_id":"47064579-345f-4d28-9dce-f5a7b2285cb4","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"AutoSAM: an agentic framework for automating input file generation for the SAM code with multi-modal retrieval-augmented generation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:90ca629716d0504a5851ed5c7e3ec18255d894f64263668506d163343b9991f3","observation_id":"ade3ced7-7bf8-49f6-8c9c-45359b431a77","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:9d970ad35cdd69e81af0e797188bc3955f35e52ac9da654f9671058e3070f3c5","observation_id":"fcf21966-a402-4969-a9e8-e9c892463ce2","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.11014","last_updated":"2025-07-15T06:14:45Z","snapshot_observed_at":"2026-08-18T18:59:24.869424Z","submitted_at":"2025-07-15T06:14:45Z","title":"SIMCODE: A Benchmark for Natural Language to ns-3 Network Simulation Code Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.11014","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"SIMCODE: A benchmark for natural language to ns-3 network simulation code generation ; 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2507.11014","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:62062687bdc580855e616c59b063157e17799513ffd56ae760f3af171dff8b53","observation_id":"2b838f9c-ea72-49c4-99a7-ddd74b777c1e","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2603.20630","last_updated":"2026-05-21T19:14:59Z","snapshot_observed_at":"2026-08-06T21:24:24.530237Z","submitted_at":"2026-03-21T03:58:40Z","title":"Evaluating LLM-generated code for domain-specific languages: molecular dynamics with LAMMPS","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2603.20630","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Evaluating LLM-generated code for domain- specific languages: Molecular dynamics with LAMMPS ; 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2603.20630","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:3bfc7ddd8e9648d22956261907ecc6ff8321950efac02bae077093be93da5b7c","observation_id":"44b8a582-24c6-44dd-aa3c-ad2baa6257cd","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"GRACE: an agentic AI for particle physics experiment design and simulation ; 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:aa4b6a167b2401e98a27a49504ab01ea8c6b2f863c7bc9cabed25a4c5a9c2a2d","observation_id":"3d8dc4d3-787a-4b7f-9c6c-51e4f7f86467","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.19863","last_updated":"2025-06-26T22:36:10Z","snapshot_observed_at":"2026-08-18T00:30:16.226603Z","submitted_at":"2025-06-10T09:28:18Z","title":"Exploring the Capabilities of the Frontier Large Language Models for Nuclear Energy Research","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.19863","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Exploring the capabilities of the frontier large language models for nuclear energy research ; 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2506.19863","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:753f053cd92bc5604873848ad308141c300074d698b63a2ecc7120abb47fa199","observation_id":"5511e7fb-efc0-4a9e-8209-02e513c962db","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.05987","last_updated":"2023-02-18T18:27:49Z","snapshot_observed_at":"2026-08-16T16:46:24.365290Z","submitted_at":"2022-07-13T06:47:51Z","title":"DocPrompting: Generating Code by Retrieving the Docs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.05987","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"DocPrompting: Generating code by retrieving the docs","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2207.05987","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:a777dce6f5c522c120db9cc2d55bc43c176cdda436478881e4d107f7a2b080c4","observation_id":"35839785-28fd-486b-9ecb-22dda7be9f28","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11925","last_updated":"2024-07-03T09:16:27Z","snapshot_observed_at":"2026-08-16T13:42:22.029356Z","submitted_at":"2024-06-17T08:34:57Z","title":"DocCGen: Document-based Controlled Code Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11925","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"DocCGen: Document-based controlled code generation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2406.11925","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:fd764ed56ec00e74e8aaa57c705c0f3f00a3e77c70e078e08eaeb348b0a47191","observation_id":"ed989457-4c9a-47d9-8abd-7668b2dd119f","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"MetaGPT: Meta programming for a multi-agent collabo- rative framework","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:268ddf38b3467286b15f22bd8126908afa2308ef2ba60934e36e07ea1de9689f","observation_id":"1e1b44e7-7e3e-4872-ad10-01ae6563dd02","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"AutoGen: Enabling next-gen LLM applications via multi-agent conversation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:ca8a2b9e46512ef1fb0f488624803111ebf6a664a2437a251712feff07b37d58","observation_id":"15c7861c-7add-4412-a054-9af9e2760ced","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"SciAgents: Automating scientific discovery through bioinspired multi-agent intelligent graph reasoning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:83ba0bf98ea712ec305a27f9c901aff113f23c62322a332fe996d9e13d2834c2","observation_id":"6f42b1bd-8b7f-4ad2-8848-b9f1feb20171","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"SWE-agent: Agent-computer interfaces enable automated software engineering","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:3b18cc15b25bc33a44aa1ef9d17e4ca0ea37fc2e1af3e9accb1a9a8673949ae1","observation_id":"b28d99c9-199b-44fd-a2bd-2f5f8ad4f558","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"ReAct: Synergizing reasoning and acting in language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:e233cdfc990470490df7d4caa75373d3895acefe058c71a7b79d6eee3a69016f","observation_id":"aab0b677-b066-4aac-b79c-f88ed43ce24d","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Self-refine: Iterative refinement with self-feedback","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:00903369fb82025e47c49f17f4ca17bfff0778a16ece396887d0f71878b91aac","observation_id":"77893e61-e976-4871-8ceb-b54da612f8a7","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Reflexion: Language agents with verbal reinforcement 20 learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:7266b942282e1e0b5e1c487651b3155ba3161b2a60dbd7b0bb2bc0d5b0a00459","observation_id":"3650f9e5-6de9-45a2-a462-2c478eb25457","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Large language models cannot self-correct reasoning yet","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:692c18f16a902fd7abba5457f61960cc237ef1e627a61c53fefc2344541604cd","observation_id":"bddc33ad-7626-4c93-b482-0c8266373cea","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Is self-repair a silver bullet for code generation? In: The Twelfth International Conference on Learning Representations (ICLR); 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:98ce3afee3f39d6de4843d0196f79763e3ff1cc6a817f9b9477290705e426ff7","observation_id":"40e6902e-b75a-44bd-ba15-268027f2cd00","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"When can LLMs actually correct their own mistakes? A critical survey of self-correction of LLMs","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:3d49116fb940b5c9c979558fa40ec85727de94746b7f0c8852dbb4bea985b067","observation_id":"4bdb19ce-56ff-4bde-b79f-e1e2f2497bac","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"LiveCodeBench: Holistic and contamination free evaluation of large language models for code","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:88a761304fb5fe7ec8526566a60baa22887357d0af6cf4b865b06bc1fafb9c9c","observation_id":"65e5bce6-6bc3-47e8-a30c-e4b7f8017a84","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2009.10297","last_updated":"2020-09-27T04:07:11Z","snapshot_observed_at":"2026-08-12T14:26:33.940158Z","submitted_at":"2020-09-22T03:10:49Z","title":"CodeBLEU: a Method for Automatic Evaluation of Code Synthesis","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.10297","snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"CodeBLEU: a method for automatic evaluation of code synthesis [Arxiv:2009.10297 [cs.se]]; 2020","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"cited_paper":"/paper/2009.10297","citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:ebf2e88af4c89215e0ed24bd6799e72bef0597f94a7f43a4ef6c857da48a61fd","observation_id":"bf91b133-af1e-49c7-ae73-c58328577741","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Codex CLI: Command-line coding agent [https://github.com/openai/codex]; 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:03882cd50a6b483d5a79e2a4611d3e3b905b2af9c9acb45a9e8fcbfc72713c58","observation_id":"6ca96384-62a1-4cea-aa1d-2d3c332f1651","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"OpenAI Agents SDK [https://github.com/openai/openai-agents-python]","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:66d084daf9fb7fe0dbe09e68493faceaf742233fd94997f0fac871df75010326","observation_id":"8ffd9203-55d5-4113-bb45-fc810badda2f","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:143b3f51ee467c1ca71b6106e36eecdefd500432f62ac672fa2da94f5f5c2716","observation_id":"968c8a8d-0454-4d0d-855b-0ecb0dbf56a2","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"XCOM: Photon cross section database (version 1.5) [National institute of standards and technology, gaithersburg, md]; 2010","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:342e058f076d33e6250b18f3f4328dc7c125a1a5ac18963a1530218c1155025f","observation_id":"c7f85bdd-9ed9-4a06-94d9-dbf1870a0fce","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Nuclear data sheets for A = 137","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:635bdf7aa4ca76d238404c9c474c2d8e0e6bd34bacf3269da24289b42fe4fb72","observation_id":"756cc65b-d048-4e7a-b306-b23f4059fe35","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Benchmark study of particle and heavy-ion trans- port code system using shielding integral benchmark archive and database for accelerator- shielding experiments","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:ad392e7eeedb8e6f9ef895e1dd5f037ed37b15a07269891ce989c0dbcd3ef9a1","observation_id":"b50c2e4a-5bef-46a4-9afc-9b8305f862ef","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Validation of the physical and RBE-weighted dose estimator based on PHITS coupled with a microdosimetric kinetic model for proton therapy","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:acff71b3c1176110a4701d6ca1696bc5335afae899ece5305bc2a538a4d820e0","observation_id":"6803b39b-0930-40ce-b46f-27ca02a5b544","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Improvements in the particle and heavy-ion transport code system (PHITS) for simulating neutron-response functions and detection efficiencies of a liquid organic scintillator","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:c542f4bbc880fc15c4105ba7e357162d60a191f2023d9498db5e8395fa7a12ca","observation_id":"66efbf9a-2390-4974-932a-05a3afcc2523","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"A PHITS-based computational model of a TRIGA-fueled subcritical reactor for gamma dose mapping","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:7a376c86ec425dc9fd77ebdaa84b21cb752d2eba2ffdf95dae8f05237e98ad92","observation_id":"eb0acfbc-db60-42de-8e41-4b4f0a01d4b6","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T15:42:46.891891Z","title":"Dose estimation for astronauts using dose conversion coeffi- cients calculated with the PHITS code and the ICRP/ICRU adult reference computational phantoms","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-07-14T15:42:46.891891Z"},"links":{"citing_paper":"/paper/2607.09789"},"observation_digest":"sha256:49a57b28a5793737c870a921439b277ac20a6d68abc60e45b49e6933955795c0","observation_id":"75b2fa6e-a836-43e6-bfc1-d3c6ebcfca5d","resolution":{"observed_at":"2026-07-14T15:42:46.891891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.09789","last_updated":"2026-07-08T16:35:22Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-18T21:57:50.309439Z","submitted_at":"2026-07-08T16:35:22Z","title":"PHITSBench: an execution-scored benchmark for AI-assisted PHITS radiation-transport input generation using natural language"},"reference_resolution":{"displayed":47,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":47,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":47},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 47 of 47 outbound references and 0 inbound Pith citation observations for arXiv:2607.09789."}