{"as_of":"2026-08-09T10:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b632848782003026468ff69265bfb25c54c1131c1d9cee2ed51c988e1f861347","coverage":[{"denominator":52,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":52,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T18:22:38.879800Z","state":"measured"},{"denominator":53,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":53,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T00:27:21.014118Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-05T06:16:20.379382Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"cited_work":{"arxiv_id":"2508.14764","doi":"10.48550/arxiv.2508.14764","metadata_source":"pith","pith_arxiv_id":"2508.14764","snapshot_observed_at":"2026-08-05T06:16:20.379382Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","venue":"physics.ed-ph","work_id":"c4c3277e-4086-4d3a-9a45-4393797383f1","year":2025},"citing_paper":{"arxiv_id":"2608.00748","last_updated":"2026-08-01T16:32:03Z","snapshot_observed_at":"2026-08-09T07:51:15.316574Z","submitted_at":"2026-08-01T16:32:03Z","title":"Me and My Bot: What Users Talk About in AI Companion Communities on Reddit","version":1},"reference_index":2026,"source":"pdf_text","source_observed_at":"2026-08-05T00:27:21.014118Z"},"links":{"cited_paper":"/paper/2508.14764","citing_paper":"/paper/2608.00748"},"observation_digest":"sha256:7ee56b4a728b95948c72dec5c602ac25e5e8823216bbeb0100e6a19d0d9d5cb7","observation_id":"a725a0bc-af27-40f1-a102-f2567f27ae3c","resolution":{"observed_at":"2026-08-05T00:27:21.045369Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-07T21:38:08.171017+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-07T21:38:08.171017+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2508.14764/citation-record","integrity":"/paper/2508.14764/integrity","json":"/paper/2508.14764/citation-record.json","paper":"/paper/2508.14764"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:36.079033Z","title":null,"venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.079033Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:09b32be3f9094c6517d1703d936ecb62324697c3b45979b6debb56007777e170","observation_id":"0cd3559e-7c47-4d36-a9a4-66e5ed25346b","resolution":{"observed_at":"2026-08-05T18:22:36.079033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.733757Z","title":"To find agreement between human raters and LLMs, we FIG","venue":null,"work_id":"2d54586e-f53e-4b87-9b68-b94b95bd727c","year":null},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:35.961318Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:0bffd1587f8d6a3b2d9636c20c0067b15115a25119a31f7058ef105d1c6534cb","observation_id":"fd02565e-4b8c-45ba-b8e0-e410f1fb4001","resolution":{"observed_at":"2026-08-05T18:22:39.738538Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.704147Z","title":"Honey, G","venue":null,"work_id":"bc2c8c54-3352-4b12-9fc0-a5e8d6e4975f","year":2014},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.136603Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:58f4e1eeca24b39bc791a2488a2bf4f49268091b2cf381dd1e4aeef534251f8f","observation_id":"860c4662-d38d-4cfc-b09c-f9601c390cc1","resolution":{"observed_at":"2026-08-05T18:22:39.709235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.684647Z","title":"Fischer, C","venue":null,"work_id":"3fdabf3e-cbc2-4dad-b784-91e0febc512f","year":2014},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.194604Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:32f639342d514822d7d7a260edf6470db0d87f1387573112f30f721eff0aaf47","observation_id":"14fd7bd6-97ad-4816-9e25-04ad74c54f50","resolution":{"observed_at":"2026-08-05T18:22:39.689283Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.664540Z","title":null,"venue":null,"work_id":"d5c84c3b-b9ba-4286-b115-6d75ea1de93a","year":2012},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.281607Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:d9d743769c6f68b2ac7023e5f3f6d6418803826d22198495496463134ae1afd3","observation_id":"168f0429-891a-4b22-9c06-4a66eef1d926","resolution":{"observed_at":"2026-08-05T18:22:39.668790Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.645833Z","title":"Dalal, A","venue":null,"work_id":"4c3f0517-0d2f-444e-af62-e732737f536a","year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.372775Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a68f39847d11acd273d989aa35b677cf6fd80e457e5d54a55ba5987e3b8215af","observation_id":"040f22c4-a813-4b66-9aae-1b20385d975f","resolution":{"observed_at":"2026-08-05T18:22:39.651125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.625880Z","title":"Slavit, E","venue":null,"work_id":"8cce6639-9ecc-4506-9363-fda9b53973f3","year":2019},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.454909Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:0d4ca4ccbf5b4b821f55c756cb3005069eab55b8d41140e118373eaa07b96b1e","observation_id":"ca62b77a-8bb9-4bdc-b056-1eed9bb0fd1d","resolution":{"observed_at":"2026-08-05T18:22:39.632136Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.606220Z","title":"Slavit, E","venue":null,"work_id":"ca7ab467-ac8e-4fac-a289-d34b92eaaab9","year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.533925Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a276dbcea09ba641bf0f9231c0d74a217c914a7a1e8ce41de2eabd05b64ba0a6","observation_id":"7394c228-c935-4989-b59f-9bf1af8b8b1e","resolution":{"observed_at":"2026-08-05T18:22:39.611080Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.587585Z","title":null,"venue":null,"work_id":"b461d305-f7bf-41c0-b141-32a37ab0a109","year":2020},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.616010Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:920970f683e15c3d627b44a5064ff10deacb2dd41a1bfcd4d93983c800da449d","observation_id":"dc0323ed-ec0c-4345-b45e-e458d747f176","resolution":{"observed_at":"2026-08-05T18:22:39.593548Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.571658Z","title":null,"venue":null,"work_id":"19edf2b9-376a-445d-845d-df05d09443cb","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.680585Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:c181f7a7d0a6b67de1907764ecfd53eecdd159b0197bcfba8619f4ea78b7d752","observation_id":"6efca26d-85ee-47a4-839f-1f001dd2acc7","resolution":{"observed_at":"2026-08-05T18:22:39.576100Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.554082Z","title":"Talanquer and J","venue":null,"work_id":"267639df-c276-4b89-ba51-ed10619e32c2","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.742077Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:8f119f31f3e6740d70981000f6d6517f791882362b462cdc571e994d850ba2e4","observation_id":"4e06193e-d724-4040-95eb-28b9667b0c85","resolution":{"observed_at":"2026-08-05T18:22:39.559225Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.538541Z","title":null,"venue":null,"work_id":"5fd0cff8-e9a7-49a5-9352-556bdb616e37","year":2018},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.810808Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:f2c4bcdfe85ca5a4abf67e7d12c3044dad0f7e65e6343216fb0ddb89b30ad4f8","observation_id":"cbcd871d-2e95-4ab9-9183-1c23dcaeefc9","resolution":{"observed_at":"2026-08-05T18:22:39.542863Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.523219Z","title":"Wilkinson, A","venue":null,"work_id":"d569ea19-2096-42c8-bc06-8ee8f98d6bd3","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.872197Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:455d8e4ecd70cf7429a1020c011b78c58507bff229cdc7e0c8f9a83c443f2bcd","observation_id":"1c33057e-cc00-46b8-86bb-3de6c7597b72","resolution":{"observed_at":"2026-08-05T18:22:39.527850Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.504655Z","title":"Etkina and G","venue":null,"work_id":"385a6385-3c16-4238-8706-6940b2994439","year":2014},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.963718Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:bc2592fb1518af8fdc5a87d964082a0f27d54266f9c7bcd647d967d146f89e2b","observation_id":"28256f66-6d61-45f8-b4c1-0cb65c0a880e","resolution":{"observed_at":"2026-08-05T18:22:39.509389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.488921Z","title":null,"venue":null,"work_id":"8084f776-95c9-4976-a4f7-7e2d6c1da5f6","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.004072Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:3a439c3d53b566561e54f5c75ba8da7441dad388d37aa03197a0f8ad0510fa5c","observation_id":"e2cbc877-a394-463a-a7a0-1f0ab554990a","resolution":{"observed_at":"2026-08-05T18:22:39.493736Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.475134Z","title":null,"venue":null,"work_id":"1182668d-4653-4878-9b87-ac18f07ff55b","year":1996},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.042177Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:6e4591e736410d8176a0e49489a91ec4468fa3567a426cffe81d767f71fff990","observation_id":"39f0ff9c-2fb9-4455-adcd-75612cd3dd4f","resolution":{"observed_at":"2026-08-05T18:22:39.479473Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.460637Z","title":"Erickson et al., Qualitative methods in research on teaching (Institute for Research on Teaching East Lansing, MI, 1985)","venue":null,"work_id":"f95af4a7-9f3d-4ecc-9cbf-57b6247d5824","year":1985},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.101955Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:0ece4fd6f90d9ebba9fc995059ab3e5a2c1feb751442c5129af06adb6a95d327","observation_id":"579f2ea1-48c7-4a39-94d9-4460f43e44eb","resolution":{"observed_at":"2026-08-05T18:22:39.465025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.445616Z","title":"Braun and V","venue":null,"work_id":"6f24eb33-028e-4068-bddc-43a3b49a1e46","year":2006},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.267943Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:223890f141f0da779a03c4cd892b28e9eed6e7b1f8021f7d5d06be6e784a2902","observation_id":"d602f4a1-0dfc-4d58-a647-1fbbb5566f86","resolution":{"observed_at":"2026-08-05T18:22:39.450146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.430278Z","title":null,"venue":null,"work_id":"450ad6cd-bb5f-4faa-9734-39daf30e87dc","year":2016},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.372269Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:ee4ae0d4c30654a2900a79523b74baabc6011762950d234f4d9aade24154979e","observation_id":"38c3cbe9-82a8-427b-803b-ec5ba860b07e","resolution":{"observed_at":"2026-08-05T18:22:39.435180Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.416023Z","title":null,"venue":null,"work_id":"df8a8c54-51b8-45bb-8f30-56be7dd46890","year":2015},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.488537Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:f41541c6bd874510ce26a29e4f1e54f1bd6ae12aae6f520709bd04889cc5ed09","observation_id":"c180b053-addb-4f86-a0c2-947d47547ee9","resolution":{"observed_at":"2026-08-05T18:22:39.420040Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.398323Z","title":"Saldana, The coding manual for qualitative researchers , V ol","venue":null,"work_id":"860b5039-18ee-41dd-a086-252d1c4257b8","year":2009},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.586251Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:d3d29ad547b730ac3329091fd691a4795a4d6b388d6254be91e5cdbafba160e3","observation_id":"aa043082-b914-4871-881c-ad720b52b1dd","resolution":{"observed_at":"2026-08-05T18:22:39.403882Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.379235Z","title":"Houghton, K","venue":null,"work_id":"8df4bdfd-f0ec-409c-ab9f-0cc58f6b5456","year":2015},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.677001Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:7ed8801ffbb90f403ba2e205422bd3d11d4652716cc783e815314628392a3d79","observation_id":"673e88e1-8bc9-49e9-9677-20bf4ea635f6","resolution":{"observed_at":"2026-08-05T18:22:39.384281Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.363503Z","title":"Jackson and P","venue":null,"work_id":"0917b8cc-cffb-4eb4-b86f-94c1408ad9eb","year":2019},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.773548Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:9801765ac7de3e7781fe544d2194122f9fc9c175c63da7df8ccecc18a5abc9c4","observation_id":"fee14e93-29f6-4660-b040-9a7e7dcb63d3","resolution":{"observed_at":"2026-08-05T18:22:39.368611Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.348201Z","title":"Uysal and N","venue":null,"work_id":"adceec83-68b5-43b4-8775-e1cbb2abc87a","year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.866609Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:c604dc6e7f350c2ab4d6b897d53f2c5de45c4d8edb018dadd7ac356031a944ca","observation_id":"4fb51306-6fbd-400d-84aa-9c1814f7e45c","resolution":{"observed_at":"2026-08-05T18:22:39.352912Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.331781Z","title":null,"venue":null,"work_id":"fb8694a4-2d5b-43d6-8545-f78da414407b","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.001074Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:1cf88961dec4550240f003faefeecbab4ba6931f8b69f498b64c53f11c68ee99","observation_id":"d045538f-c993-46f5-9e2c-424cdeba319a","resolution":{"observed_at":"2026-08-05T18:22:39.336619Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.316896Z","title":null,"venue":null,"work_id":"d88b3676-7221-4030-87ea-0e4e733172bf","year":2009},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.135492Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:0ae2b7fc02e16ce1b762d6b115b1f089a841d336a7989d57302ea0756f1c94d7","observation_id":"c68b5e03-f09f-4cd2-ae36-987cacdb21ba","resolution":{"observed_at":"2026-08-05T18:22:39.321165Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:38.228019Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.228019Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:12fa1bd5f79874237144a1ad20109d33407af37806214057cabad0cdf7fdac3f","observation_id":"412135b2-2063-4713-9220-7a325364a447","resolution":{"observed_at":"2026-08-05T18:22:38.228019Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.291277Z","title":null,"venue":null,"work_id":"9afad1db-c7e5-42ff-92ec-cbc643f0b857","year":1975},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.319873Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:d341b5ade2503d566dd8514a520b95bef812bcdde067ce51e442b2861696588f","observation_id":"55eb4a2e-a401-48d8-a9a0-5af891485742","resolution":{"observed_at":"2026-08-05T18:22:39.295854Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.275827Z","title":null,"venue":null,"work_id":"252dbb48-33e3-43e8-ac3a-d38d39cb041c","year":null},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.482804Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:9d5471ea0c46a36781c513a966051597198d294bd971de17ce75f5cfb8c99788","observation_id":"acca8eee-682e-423e-bf31-19db028dbc75","resolution":{"observed_at":"2026-08-05T18:22:39.280200Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11882","last_updated":"2023-05-09T19:55:50Z","snapshot_observed_at":"2026-07-06T15:29:50.743434Z","submitted_at":"2023-05-09T19:55:50Z","title":"Exploring the Efficacy of ChatGPT in Analyzing Student Teamwork Feedback with an Existing Taxonomy","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11882","snapshot_observed_at":"2026-08-05T18:22:38.608842Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.608842Z"},"links":{"cited_paper":"/paper/2305.11882","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:5040348246d03fed89741e9948f8f59418d071f6f48174f275ec3d13a508119e","observation_id":"13eb7344-ee42-4540-a345-b76a39ec6199","resolution":{"observed_at":"2026-08-05T18:22:38.608842Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.261565Z","title":"Hitch, Artificial intelligence augmented qualitative analysis: the way of the future?, Qualitative Health Research 34, 595 (2024)","venue":null,"work_id":"c7952536-067a-450d-9d04-1de7dd0179d2","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.688426Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:e3f8081a77811f1c792b22fee79fe5dfb41aefd017c71c7fe30397b38bc499c6","observation_id":"ee7926f1-8ba1-497c-8505-f09b31213a0e","resolution":{"observed_at":"2026-08-05T18:22:39.265877Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.247946Z","title":"Tabone and J","venue":null,"work_id":"1d85c7f0-5df7-4651-94d2-c297d28c6c45","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.732854Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:bd94ae656c9f3b610c7341ced8c3a5c2bb9c8cd8f75bfd81fa8f9320b6838c0b","observation_id":"5a7d55a1-ff41-457d-b963-04c8feb8903e","resolution":{"observed_at":"2026-08-05T18:22:39.252102Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.10771","last_updated":"2024-05-28T02:26:20Z","snapshot_observed_at":"2026-08-07T04:08:37.186859Z","submitted_at":"2023-09-19T17:18:09Z","title":"Redefining Qualitative Analysis in the AI Era: Utilizing ChatGPT for Efficient Thematic Analysis","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.10771","snapshot_observed_at":"2026-08-05T18:22:38.788741Z","title":"Zhang, C","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.788741Z"},"links":{"cited_paper":"/paper/2309.10771","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:70399fb06569de389bc6c68752d00f356493533f9302ed08f378a71c82811bf4","observation_id":"a6aafd8f-42f5-4d83-9d5d-123f3d742975","resolution":{"observed_at":"2026-08-05T18:22:38.788741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.234033Z","title":null,"venue":null,"work_id":"28fd1406-9860-4709-8eeb-f32d5fb1d920","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.801848Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:c0f9d07cfbd2443cea410d272d1af565543ac5bbc88006c1c50b4d103f4c202e","observation_id":"1bf2da57-7cf4-485b-b228-bf35e0831d77","resolution":{"observed_at":"2026-08-05T18:22:39.238145Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.219712Z","title":null,"venue":null,"work_id":"99e3ebf2-fc87-4d2f-a4cd-b3e0806ecb24","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.805840Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:f72eb9495fbf9ef9614421e56ad695ca941d4e101d9f2ddf936b9454a712c321","observation_id":"a46cf742-9688-4c6d-b747-32993b68e106","resolution":{"observed_at":"2026-08-05T18:22:39.223869Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.205555Z","title":null,"venue":null,"work_id":"d46836fd-48c7-4d8d-bb80-c3548faa2dae","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.809762Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:c465eeebfa4579b47c51c7105510f07155291243ba0fe2a6b7b36d689e8f50b6","observation_id":"e63cbe10-d1dd-484a-ab06-47408f9f0435","resolution":{"observed_at":"2026-08-05T18:22:39.209920Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T18:22:38.813821Z","title":"Achiam, S","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.813821Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:3affd85bb7c2eec3007bdd71208d95febb9f9bcc3e130d35278ab3edafffa0cc","observation_id":"d01810d8-3b3e-4b67-bb33-6dc1326e7abd","resolution":{"observed_at":"2026-08-05T18:22:38.813821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.190300Z","title":null,"venue":null,"work_id":"1fca96c4-3064-47df-9ac4-2a7c65d2ec99","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.818785Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:589e5de02ce933eab72b2ed357ae223f4d98552c33ce52e6971c17b72f6869b8","observation_id":"c9889a5f-d283-4e59-9115-b98a93a5aefe","resolution":{"observed_at":"2026-08-05T18:22:39.195389Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.174935Z","title":null,"venue":null,"work_id":"b8fc65c8-71ec-4259-ae69-2cd35ddf26d3","year":2020},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.823114Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:127f241e72961bdb79b9360ff634df23057314bed69ed7c5b2d882f0cb402b15","observation_id":"eb5bfd5a-3f86-48fa-9eb3-4369bd977176","resolution":{"observed_at":"2026-08-05T18:22:39.179312Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.159834Z","title":null,"venue":null,"work_id":"4d73f6fd-6e1f-4fb5-b16d-86363ef697d4","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.828118Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:8e784af94a35d47f871e7c23173360f578331a4a555b57e644e9ad00a2cd3b27","observation_id":"061a9329-e136-4d17-b7a9-30c276c9a24d","resolution":{"observed_at":"2026-08-05T18:22:39.164102Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.144494Z","title":null,"venue":null,"work_id":"29dbf692-fe42-4cf2-ba21-9f2040d98a19","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.832151Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:b2daf5d4fb2e2ae075ff30542034700f6735fb06a67cf28b1a759cb45840efb3","observation_id":"5e445e4a-ce0b-49ea-86d2-77a01b2a7879","resolution":{"observed_at":"2026-08-05T18:22:39.149368Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05957","last_updated":"2025-03-13T14:42:31Z","snapshot_observed_at":"2026-08-07T17:21:01.739960Z","submitted_at":"2025-03-07T21:52:21Z","title":"Applying a STEM Ways of Thinking Framework for Student-generated Engineering Design-based Physics Problems","version":2},"cited_work":{"arxiv_id":"2503.05957","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.05957","snapshot_observed_at":"2026-08-05T18:22:38.949718Z","title":"Applying a STEM Ways of Thinking Framework for Student-generated Engineering Design-based Physics Problems","venue":"physics.ed-ph","work_id":"1ed41afe-076b-4b59-b9de-11b1285a057c","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.836255Z"},"links":{"cited_paper":"/paper/2503.05957","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:eb9c4a90e170642659ce565c81e46966f1ebf05c1c8add515865fb81fc94fdd8","observation_id":"93731589-e9a2-4cfe-aa43-802585d2a1be","resolution":{"observed_at":"2026-08-05T18:22:38.957039Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.128804Z","title":"Bijker, S","venue":null,"work_id":"6fe392a1-7b6f-47d9-ac0a-b91676628fd6","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.840450Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:e7fe8fd408a9c8e2d7c3fd7b672c473aa9be0dd9c1acd4cb21a6dcf413cab8cb","observation_id":"750d32c8-dbc0-486b-a058-395f0527c9ee","resolution":{"observed_at":"2026-08-05T18:22:39.133612Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.111264Z","title":null,"venue":null,"work_id":"21bae3cf-1a7b-4699-bbca-9e4418f0ad1f","year":1994},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.845001Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:aa6d48bf46f0caa57e63a433253fa233c692e6c156397f8766a6cca025c581c3","observation_id":"4264472f-d91f-45ba-ad8d-217bce3634ca","resolution":{"observed_at":"2026-08-05T18:22:39.116170Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.094608Z","title":"Mizumoto and M","venue":null,"work_id":"d1de9057-3242-4a4a-b3ff-bd11f10d7c09","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.849419Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:97dfc01961b47239e647102a24cbefce13612eda4be9070a5a628bcd6c876018","observation_id":"f228e63f-10c1-4881-89a2-33e3e30c84e5","resolution":{"observed_at":"2026-08-05T18:22:39.099592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.14735","last_updated":"2025-05-11T09:23:41Z","snapshot_observed_at":"2026-07-06T16:37:03.248953Z","submitted_at":"2023-10-23T09:15:18Z","title":"Unleashing the potential of prompt engineering for large language models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.14735","snapshot_observed_at":"2026-08-05T18:22:38.853411Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.853411Z"},"links":{"cited_paper":"/paper/2310.14735","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:1397d8c721b487aa54c995e44684867adb82eed586cb25fb8bae0562ea3fbdca","observation_id":"9ab8cb29-450f-4460-9d6e-759f20f9b023","resolution":{"observed_at":"2026-08-05T18:22:38.853411Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.078160Z","title":null,"venue":null,"work_id":"94afae27-f4ef-4d2f-aaa7-25afb6951158","year":1977},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.857964Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:79c2d827f887a399158b2317ef441570e5d1091932f9d82198ffc664e836c734","observation_id":"916401d2-7ab3-4ccd-9f9d-c2fccf468c75","resolution":{"observed_at":"2026-08-05T18:22:39.083092Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.062367Z","title":null,"venue":null,"work_id":"4596e2a2-42de-455b-ba6f-c8342239bfca","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.862085Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:26ca0c7e5d140cc16b1e78aaab05ef1c7070b3e43ef79b8b2b78fc3520b7eb93","observation_id":"c6c5f65d-8733-46b0-b53e-b45f25d46936","resolution":{"observed_at":"2026-08-05T18:22:39.067668Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.047079Z","title":null,"venue":null,"work_id":"eecdd225-1274-4221-a308-ed462f353179","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.866440Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:5d4dd70e0f67dc09cc42dd4e496d757606bf2080aa27727c255f83960bf3e741","observation_id":"671f9955-5cd0-4760-8162-6d91ee302200","resolution":{"observed_at":"2026-08-05T18:22:39.051844Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.032621Z","title":null,"venue":null,"work_id":"4f016e41-f696-43de-b845-cdd604b8d4a8","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.870936Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:da43440549e4dcbf7a5ed835d0376d675e7d7efcf594140611f5abeed57820f4","observation_id":"8ba0ae9c-e47d-46b6-b589-8a8684f917d2","resolution":{"observed_at":"2026-08-05T18:22:39.036749Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-05T18:22:38.875155Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.875155Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:927a061f2328d80774644bd6428304575126de33b405595f9873de4b72f078c3","observation_id":"f45be533-4e8b-4905-b790-667213e86783","resolution":{"observed_at":"2026-08-05T18:22:38.875155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.017511Z","title":"Tschisgale, P","venue":null,"work_id":"23858bb2-61ca-4319-aeab-e7e0c2ee1637","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.879800Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:03568ecdb3f7a202862bf3c70222031750526319ad8de6dc7029e2ace6b5080f","observation_id":"8c9748e9-2495-4725-86ad-04e6627ce478","resolution":{"observed_at":"2026-08-05T18:22:39.022025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","latest_version":2,"primary_category":"physics.ed-ph","snapshot_observed_at":"2026-08-09T00:15:51.869634Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis"},"reference_resolution":{"displayed":52,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":31,"verified_exact":1,"verified_fuzzy":20},"total_outbound_references":52},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 52 of 52 outbound references and 1 inbound Pith citation observation for arXiv:2508.14764."}