{"as_of":"2026-08-09T19:11:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bf51cc4557a06abee17ec6e52f7aef2cfdccfa07a7454004e461db1e88e0976a","coverage":[{"denominator":52,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":52,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T18:22:38.879800Z","state":"measured"},{"denominator":53,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":53,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T00:27:21.014118Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-05T06:16:20.379382Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"cited_work":{"arxiv_id":"2508.14764","doi":"10.48550/arxiv.2508.14764","metadata_source":"pith","pith_arxiv_id":"2508.14764","snapshot_observed_at":"2026-08-05T06:16:20.379382Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","venue":"physics.ed-ph","work_id":"c4c3277e-4086-4d3a-9a45-4393797383f1","year":2025},"citing_paper":{"arxiv_id":"2608.00748","last_updated":"2026-08-01T16:32:03Z","snapshot_observed_at":"2026-08-09T16:56:56.289112Z","submitted_at":"2026-08-01T16:32:03Z","title":"Me and My Bot: What Users Talk About in AI Companion Communities on Reddit","version":1},"reference_index":2026,"source":"pdf_text","source_observed_at":"2026-08-05T00:27:21.014118Z"},"links":{"cited_paper":"/paper/2508.14764","citing_paper":"/paper/2608.00748"},"observation_digest":"sha256:2cabb29b8c1bcdeb041796c3dc2962348086b877ba4a9a394fd3d18c3ab8f728","observation_id":"a725a0bc-af27-40f1-a102-f2567f27ae3c","resolution":{"observed_at":"2026-08-05T00:27:21.045369Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-07T21:38:08.171017+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-07T21:38:08.171017+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2508.14764/citation-record","integrity":"/paper/2508.14764/integrity","json":"/paper/2508.14764/citation-record.json","paper":"/paper/2508.14764"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:36.079033Z","title":null,"venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.079033Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:1e4e1389795f9db3fd3c02903e52fa257f6c44300d4436c5200631dba56d9785","observation_id":"0cd3559e-7c47-4d36-a9a4-66e5ed25346b","resolution":{"observed_at":"2026-08-05T18:22:36.079033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.733757Z","title":"To find agreement between human raters and LLMs, we FIG","venue":null,"work_id":"2d54586e-f53e-4b87-9b68-b94b95bd727c","year":null},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:35.961318Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:93dc472732c83f57b65bc413a497a091041a2cdf4ec6024b6ba3d013aa681da5","observation_id":"fd02565e-4b8c-45ba-b8e0-e410f1fb4001","resolution":{"observed_at":"2026-08-05T18:22:39.738538Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.704147Z","title":"Honey, G","venue":null,"work_id":"bc2c8c54-3352-4b12-9fc0-a5e8d6e4975f","year":2014},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.136603Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:20a58f971046403714e17f9d9fbd798d398426a2ea5d87cc2876b038815e9e70","observation_id":"860c4662-d38d-4cfc-b09c-f9601c390cc1","resolution":{"observed_at":"2026-08-05T18:22:39.709235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.684647Z","title":"Fischer, C","venue":null,"work_id":"3fdabf3e-cbc2-4dad-b784-91e0febc512f","year":2014},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.194604Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:33a32e1b6674d076a60761a67533fc20608cde6f4ee6a33c3fef0f71de4c2499","observation_id":"14fd7bd6-97ad-4816-9e25-04ad74c54f50","resolution":{"observed_at":"2026-08-05T18:22:39.689283Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.664540Z","title":null,"venue":null,"work_id":"d5c84c3b-b9ba-4286-b115-6d75ea1de93a","year":2012},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.281607Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:b8ff7d5fa7434f758c21ce9c962a2559a6c9878409fcb5bc629d5e50339132be","observation_id":"168f0429-891a-4b22-9c06-4a66eef1d926","resolution":{"observed_at":"2026-08-05T18:22:39.668790Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.645833Z","title":"Dalal, A","venue":null,"work_id":"4c3f0517-0d2f-444e-af62-e732737f536a","year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.372775Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:5daf41611baaf721be09853d8da9118ac59ae68501cc83127496efed597e0532","observation_id":"040f22c4-a813-4b66-9aae-1b20385d975f","resolution":{"observed_at":"2026-08-05T18:22:39.651125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.625880Z","title":"Slavit, E","venue":null,"work_id":"8cce6639-9ecc-4506-9363-fda9b53973f3","year":2019},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.454909Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:b90dff9cf8b004ecb5bf7a42f8037ebbd4735087826265aadfca542b54548c07","observation_id":"ca62b77a-8bb9-4bdc-b056-1eed9bb0fd1d","resolution":{"observed_at":"2026-08-05T18:22:39.632136Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.606220Z","title":"Slavit, E","venue":null,"work_id":"ca7ab467-ac8e-4fac-a289-d34b92eaaab9","year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.533925Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:b7f84ad78ca24c92a2191bb19c3bca703ea74e9a9c7d9b2bf511349a6c121eb0","observation_id":"7394c228-c935-4989-b59f-9bf1af8b8b1e","resolution":{"observed_at":"2026-08-05T18:22:39.611080Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.587585Z","title":null,"venue":null,"work_id":"b461d305-f7bf-41c0-b141-32a37ab0a109","year":2020},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.616010Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:985fd2b642621bf2743c67b560fa4c19379d683eef88047019fc787a400012e1","observation_id":"dc0323ed-ec0c-4345-b45e-e458d747f176","resolution":{"observed_at":"2026-08-05T18:22:39.593548Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.571658Z","title":null,"venue":null,"work_id":"19edf2b9-376a-445d-845d-df05d09443cb","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.680585Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:89fb91e82506d1b1245de7a7aeff30c92b8364348ad5c745be0a45185037f792","observation_id":"6efca26d-85ee-47a4-839f-1f001dd2acc7","resolution":{"observed_at":"2026-08-05T18:22:39.576100Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.554082Z","title":"Talanquer and J","venue":null,"work_id":"267639df-c276-4b89-ba51-ed10619e32c2","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.742077Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:98d8c291a60191ce23af97dbc73a17178d4962bac92c47847255b60adb7ff94d","observation_id":"4e06193e-d724-4040-95eb-28b9667b0c85","resolution":{"observed_at":"2026-08-05T18:22:39.559225Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.538541Z","title":null,"venue":null,"work_id":"5fd0cff8-e9a7-49a5-9352-556bdb616e37","year":2018},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.810808Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:4e10db5000869e37aa255903a8543e021a656ae34af16e5c7b642280125a8a00","observation_id":"cbcd871d-2e95-4ab9-9183-1c23dcaeefc9","resolution":{"observed_at":"2026-08-05T18:22:39.542863Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.523219Z","title":"Wilkinson, A","venue":null,"work_id":"d569ea19-2096-42c8-bc06-8ee8f98d6bd3","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.872197Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:1fc5ba21c795d7578e99ed7b1b4b205c2ecd0c2daec8c9891344d6a2b1fd537c","observation_id":"1c33057e-cc00-46b8-86bb-3de6c7597b72","resolution":{"observed_at":"2026-08-05T18:22:39.527850Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.504655Z","title":"Etkina and G","venue":null,"work_id":"385a6385-3c16-4238-8706-6940b2994439","year":2014},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:36.963718Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:1131189d61a756fde4c01953caa4b08b7405c914e13dc55fdcdab526fc04bdfe","observation_id":"28256f66-6d61-45f8-b4c1-0cb65c0a880e","resolution":{"observed_at":"2026-08-05T18:22:39.509389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.488921Z","title":null,"venue":null,"work_id":"8084f776-95c9-4976-a4f7-7e2d6c1da5f6","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.004072Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:5cbf641ae207679314958b64e44207168671e425954bd8d4f58eecfbdca829d6","observation_id":"e2cbc877-a394-463a-a7a0-1f0ab554990a","resolution":{"observed_at":"2026-08-05T18:22:39.493736Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.475134Z","title":null,"venue":null,"work_id":"1182668d-4653-4878-9b87-ac18f07ff55b","year":1996},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.042177Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:cccaeaaa50cc22f70c2d48dd6293f50f5bcd73628d51c6cfc1b0e5f40aad3dd8","observation_id":"39f0ff9c-2fb9-4455-adcd-75612cd3dd4f","resolution":{"observed_at":"2026-08-05T18:22:39.479473Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.460637Z","title":"Erickson et al., Qualitative methods in research on teaching (Institute for Research on Teaching East Lansing, MI, 1985)","venue":null,"work_id":"f95af4a7-9f3d-4ecc-9cbf-57b6247d5824","year":1985},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.101955Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:00cb2bcd21bc17f15284c53fe80ccbda25cd3095a4f6d43dd29daba15d410d94","observation_id":"579f2ea1-48c7-4a39-94d9-4460f43e44eb","resolution":{"observed_at":"2026-08-05T18:22:39.465025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.445616Z","title":"Braun and V","venue":null,"work_id":"6f24eb33-028e-4068-bddc-43a3b49a1e46","year":2006},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.267943Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:70dfa62a359cdf2f63f547931470126c8b202fc830b89f4a0a33c019c95ba72b","observation_id":"d602f4a1-0dfc-4d58-a647-1fbbb5566f86","resolution":{"observed_at":"2026-08-05T18:22:39.450146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.430278Z","title":null,"venue":null,"work_id":"450ad6cd-bb5f-4faa-9734-39daf30e87dc","year":2016},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.372269Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:77ee8e76567d28df9ec8f0cc42fb19f510496f4752bd783defdd33fd2c1453eb","observation_id":"38c3cbe9-82a8-427b-803b-ec5ba860b07e","resolution":{"observed_at":"2026-08-05T18:22:39.435180Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.416023Z","title":null,"venue":null,"work_id":"df8a8c54-51b8-45bb-8f30-56be7dd46890","year":2015},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.488537Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:226236c5918ef4b0564ee4b4d6a9ed2b150a089e4a925083c2d36c7dbd2e7f5f","observation_id":"c180b053-addb-4f86-a0c2-947d47547ee9","resolution":{"observed_at":"2026-08-05T18:22:39.420040Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.398323Z","title":"Saldana, The coding manual for qualitative researchers , V ol","venue":null,"work_id":"860b5039-18ee-41dd-a086-252d1c4257b8","year":2009},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.586251Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:2708b388989c612ff5d508db896874e8dd130b688e5270830d0f9b51a91f1966","observation_id":"aa043082-b914-4871-881c-ad720b52b1dd","resolution":{"observed_at":"2026-08-05T18:22:39.403882Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.379235Z","title":"Houghton, K","venue":null,"work_id":"8df4bdfd-f0ec-409c-ab9f-0cc58f6b5456","year":2015},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.677001Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:e9a718968384d122766e003b3a06daa30d51d06f60525894d2620c07c5f96a3f","observation_id":"673e88e1-8bc9-49e9-9677-20bf4ea635f6","resolution":{"observed_at":"2026-08-05T18:22:39.384281Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.363503Z","title":"Jackson and P","venue":null,"work_id":"0917b8cc-cffb-4eb4-b86f-94c1408ad9eb","year":2019},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.773548Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:cf894e7388581cae61e28dda47257171a0bec3f0e056cd726d1e69d30d53f5df","observation_id":"fee14e93-29f6-4660-b040-9a7e7dcb63d3","resolution":{"observed_at":"2026-08-05T18:22:39.368611Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.348201Z","title":"Uysal and N","venue":null,"work_id":"adceec83-68b5-43b4-8775-e1cbb2abc87a","year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:37.866609Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:238dfb2eb68920ae8fe15071d4778525ab8852dfe5ae19b1b3853e1884626ccf","observation_id":"4fb51306-6fbd-400d-84aa-9c1814f7e45c","resolution":{"observed_at":"2026-08-05T18:22:39.352912Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.331781Z","title":null,"venue":null,"work_id":"fb8694a4-2d5b-43d6-8545-f78da414407b","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.001074Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a3030f217570731d93148cb25c39314a5791db8b01bb8561a70d1458ea215905","observation_id":"d045538f-c993-46f5-9e2c-424cdeba319a","resolution":{"observed_at":"2026-08-05T18:22:39.336619Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.316896Z","title":null,"venue":null,"work_id":"d88b3676-7221-4030-87ea-0e4e733172bf","year":2009},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.135492Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a83935d059bd16d4f0a58a83e1b6d979fabb9932cc1705da4dd2e81c2b47db7a","observation_id":"c68b5e03-f09f-4cd2-ae36-987cacdb21ba","resolution":{"observed_at":"2026-08-05T18:22:39.321165Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:38.228019Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.228019Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:11fbcaf9d66de2616d970aa31946b12d590b07c9abe4c1ffe3a72bd4a291ac24","observation_id":"412135b2-2063-4713-9220-7a325364a447","resolution":{"observed_at":"2026-08-05T18:22:38.228019Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.291277Z","title":null,"venue":null,"work_id":"9afad1db-c7e5-42ff-92ec-cbc643f0b857","year":1975},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.319873Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:6be899b31542fb6fe46adc997a1f039b8f8741eacd3e2c7354577aed684c3b97","observation_id":"55eb4a2e-a401-48d8-a9a0-5af891485742","resolution":{"observed_at":"2026-08-05T18:22:39.295854Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.275827Z","title":null,"venue":null,"work_id":"252dbb48-33e3-43e8-ac3a-d38d39cb041c","year":null},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.482804Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:c42dab82e831321a466895185e228c461c405d31e3aed255e0a57ed44a4b4f5e","observation_id":"acca8eee-682e-423e-bf31-19db028dbc75","resolution":{"observed_at":"2026-08-05T18:22:39.280200Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11882","last_updated":"2023-05-09T19:55:50Z","snapshot_observed_at":"2026-07-06T15:29:50.743434Z","submitted_at":"2023-05-09T19:55:50Z","title":"Exploring the Efficacy of ChatGPT in Analyzing Student Teamwork Feedback with an Existing Taxonomy","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11882","snapshot_observed_at":"2026-08-05T18:22:38.608842Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.608842Z"},"links":{"cited_paper":"/paper/2305.11882","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:366dd1f90cd9585ff73a910a5a7ec2a5dd20b6971e6f2ccfc2de7e947d2f8acd","observation_id":"13eb7344-ee42-4540-a345-b76a39ec6199","resolution":{"observed_at":"2026-08-05T18:22:38.608842Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.261565Z","title":"Hitch, Artificial intelligence augmented qualitative analysis: the way of the future?, Qualitative Health Research 34, 595 (2024)","venue":null,"work_id":"c7952536-067a-450d-9d04-1de7dd0179d2","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.688426Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:bd4d7b08feb33a816b6b680d2d2074c83ca809fd5670c30fbd92adbccb3bae07","observation_id":"ee7926f1-8ba1-497c-8505-f09b31213a0e","resolution":{"observed_at":"2026-08-05T18:22:39.265877Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.247946Z","title":"Tabone and J","venue":null,"work_id":"1d85c7f0-5df7-4651-94d2-c297d28c6c45","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.732854Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:386147c79dd35d28845a9548476bd6bd7b7fbdb2951ba2f64f53cb972ee3cee6","observation_id":"5a7d55a1-ff41-457d-b963-04c8feb8903e","resolution":{"observed_at":"2026-08-05T18:22:39.252102Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.10771","last_updated":"2024-05-28T02:26:20Z","snapshot_observed_at":"2026-08-07T04:08:37.186859Z","submitted_at":"2023-09-19T17:18:09Z","title":"Redefining Qualitative Analysis in the AI Era: Utilizing ChatGPT for Efficient Thematic Analysis","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.10771","snapshot_observed_at":"2026-08-05T18:22:38.788741Z","title":"Zhang, C","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.788741Z"},"links":{"cited_paper":"/paper/2309.10771","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:30d0eb208f56059d15fd18d63a1ba0dffe440599bd538b34795d78a77724bda9","observation_id":"a6aafd8f-42f5-4d83-9d5d-123f3d742975","resolution":{"observed_at":"2026-08-05T18:22:38.788741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.234033Z","title":null,"venue":null,"work_id":"28fd1406-9860-4709-8eeb-f32d5fb1d920","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.801848Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:5252268b980b613e7d3e14af30d1c1a0cfdd9d9b8b71e2bdface0de7e475e0e6","observation_id":"1bf2da57-7cf4-485b-b228-bf35e0831d77","resolution":{"observed_at":"2026-08-05T18:22:39.238145Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.219712Z","title":null,"venue":null,"work_id":"99e3ebf2-fc87-4d2f-a4cd-b3e0806ecb24","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.805840Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:4ab41685558507eb4caed67bc45a6a8d83179c15097d56795f3f5e36a92afc79","observation_id":"a46cf742-9688-4c6d-b747-32993b68e106","resolution":{"observed_at":"2026-08-05T18:22:39.223869Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.205555Z","title":null,"venue":null,"work_id":"d46836fd-48c7-4d8d-bb80-c3548faa2dae","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.809762Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:98b8b33ef7a732a9eec76c1b1ff2af7a81ad85a96fe2b2df2a41db553d72567f","observation_id":"e63cbe10-d1dd-484a-ab06-47408f9f0435","resolution":{"observed_at":"2026-08-05T18:22:39.209920Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T18:22:38.813821Z","title":"Achiam, S","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.813821Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:cac1244cb6b6a93d320d45d2837e48a2bac812098903fb36d6f6f70171f5fc65","observation_id":"d01810d8-3b3e-4b67-bb33-6dc1326e7abd","resolution":{"observed_at":"2026-08-05T18:22:38.813821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.190300Z","title":null,"venue":null,"work_id":"1fca96c4-3064-47df-9ac4-2a7c65d2ec99","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.818785Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:32c4e85dd8f5d5ff0c64c4c8b352493cbd5b14055c0b5462f1182bab5a54f366","observation_id":"c9889a5f-d283-4e59-9115-b98a93a5aefe","resolution":{"observed_at":"2026-08-05T18:22:39.195389Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.174935Z","title":null,"venue":null,"work_id":"b8fc65c8-71ec-4259-ae69-2cd35ddf26d3","year":2020},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.823114Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:158e884ad01ed2c8a56c293c96d5acc95c829d9fbe735b0e6eecada4d4916009","observation_id":"eb5bfd5a-3f86-48fa-9eb3-4369bd977176","resolution":{"observed_at":"2026-08-05T18:22:39.179312Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.159834Z","title":null,"venue":null,"work_id":"4d73f6fd-6e1f-4fb5-b16d-86363ef697d4","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.828118Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:deafcfa1a053f135fd4cb5583439d74af408d28629979aeb6dfc289ebbbcfc71","observation_id":"061a9329-e136-4d17-b7a9-30c276c9a24d","resolution":{"observed_at":"2026-08-05T18:22:39.164102Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.144494Z","title":null,"venue":null,"work_id":"29dbf692-fe42-4cf2-ba21-9f2040d98a19","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.832151Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a3843480979e9c364f016c26382f26a16c79283f5148ca7f25e93c9958f7a600","observation_id":"5e445e4a-ce0b-49ea-86d2-77a01b2a7879","resolution":{"observed_at":"2026-08-05T18:22:39.149368Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05957","last_updated":"2025-03-13T14:42:31Z","snapshot_observed_at":"2026-08-07T17:21:01.739960Z","submitted_at":"2025-03-07T21:52:21Z","title":"Applying a STEM Ways of Thinking Framework for Student-generated Engineering Design-based Physics Problems","version":2},"cited_work":{"arxiv_id":"2503.05957","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.05957","snapshot_observed_at":"2026-08-05T18:22:38.949718Z","title":"Applying a STEM Ways of Thinking Framework for Student-generated Engineering Design-based Physics Problems","venue":"physics.ed-ph","work_id":"1ed41afe-076b-4b59-b9de-11b1285a057c","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.836255Z"},"links":{"cited_paper":"/paper/2503.05957","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a63d9bfefbb0e541099a5d5f3141214391351929413b24c5c43e414c870ce6c5","observation_id":"93731589-e9a2-4cfe-aa43-802585d2a1be","resolution":{"observed_at":"2026-08-05T18:22:38.957039Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.128804Z","title":"Bijker, S","venue":null,"work_id":"6fe392a1-7b6f-47d9-ac0a-b91676628fd6","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.840450Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:938d04adb6094c7456fedbd24e97c26f025db779982ab96b9c77a27df18bccc3","observation_id":"750d32c8-dbc0-486b-a058-395f0527c9ee","resolution":{"observed_at":"2026-08-05T18:22:39.133612Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.111264Z","title":null,"venue":null,"work_id":"21bae3cf-1a7b-4699-bbca-9e4418f0ad1f","year":1994},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.845001Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a3e6d043e8726c7b3df4c866f92813587afd7e32812963b70e0db14ab7a3e339","observation_id":"4264472f-d91f-45ba-ad8d-217bce3634ca","resolution":{"observed_at":"2026-08-05T18:22:39.116170Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.094608Z","title":"Mizumoto and M","venue":null,"work_id":"d1de9057-3242-4a4a-b3ff-bd11f10d7c09","year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.849419Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:acfb4ac77bc19d07b015ef3d65913e97c00e585aebbb8990dff16a1428b218f4","observation_id":"f228e63f-10c1-4881-89a2-33e3e30c84e5","resolution":{"observed_at":"2026-08-05T18:22:39.099592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.14735","last_updated":"2025-05-11T09:23:41Z","snapshot_observed_at":"2026-08-09T16:54:58.064408Z","submitted_at":"2023-10-23T09:15:18Z","title":"Unleashing the potential of prompt engineering for large language models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.14735","snapshot_observed_at":"2026-08-05T18:22:38.853411Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.853411Z"},"links":{"cited_paper":"/paper/2310.14735","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:4a75b009922613ac4717e7ec27daa55ee9b0a324be0b9b44f20d5b845d2a87ce","observation_id":"9ab8cb29-450f-4460-9d6e-759f20f9b023","resolution":{"observed_at":"2026-08-05T18:22:38.853411Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.078160Z","title":null,"venue":null,"work_id":"94afae27-f4ef-4d2f-aaa7-25afb6951158","year":1977},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.857964Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:7def492bec8eb2c2a110c482e66bc11b7e771e7e77afbd0f296de0957a707202","observation_id":"916401d2-7ab3-4ccd-9f9d-c2fccf468c75","resolution":{"observed_at":"2026-08-05T18:22:39.083092Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.062367Z","title":null,"venue":null,"work_id":"4596e2a2-42de-455b-ba6f-c8342239bfca","year":2024},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.862085Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:c3b410ea0619db93451817d541d338349d08de1aeed6bb3739c50a33c8e99260","observation_id":"c6c5f65d-8733-46b0-b53e-b45f25d46936","resolution":{"observed_at":"2026-08-05T18:22:39.067668Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.047079Z","title":null,"venue":null,"work_id":"eecdd225-1274-4221-a308-ed462f353179","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.866440Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:8374cf4a92f17769b6090bd3df541fad0a5ac0fc0f8fbd44df9a39f4630ce60c","observation_id":"671f9955-5cd0-4760-8162-6d91ee302200","resolution":{"observed_at":"2026-08-05T18:22:39.051844Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.032621Z","title":null,"venue":null,"work_id":"4f016e41-f696-43de-b845-cdd604b8d4a8","year":2010},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.870936Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:4f68d1939ddebcbb59cdd4dacf084b4f9bd40f83067c417a167c6d342d3283e2","observation_id":"8ba0ae9c-e47d-46b6-b589-8a8684f917d2","resolution":{"observed_at":"2026-08-05T18:22:39.036749Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-05T18:22:38.875155Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.875155Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:a1348f8b2df3284391f9f6eefeedc0b4c4b981f3005f4c10361a85fb5b1752d9","observation_id":"f45be533-4e8b-4905-b790-667213e86783","resolution":{"observed_at":"2026-08-05T18:22:38.875155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T18:22:39.017511Z","title":"Tschisgale, P","venue":null,"work_id":"23858bb2-61ca-4319-aeab-e7e0c2ee1637","year":2023},"citing_paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.879800Z"},"links":{"citing_paper":"/paper/2508.14764"},"observation_digest":"sha256:038803ddfb31696e343e004f27ef537b88c6c78cdd4c287764dabf0887c515ef","observation_id":"8c9748e9-2495-4725-86ad-04e6627ce478","resolution":{"observed_at":"2026-08-05T18:22:39.022025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.14764","last_updated":"2025-08-31T00:33:38Z","latest_version":2,"primary_category":"physics.ed-ph","snapshot_observed_at":"2026-08-09T16:57:08.884817Z","submitted_at":"2025-08-20T15:12:52Z","title":"Investigation of the Inter-Rater Reliability between Large Language Models and Human Raters in Qualitative Analysis"},"reference_resolution":{"displayed":52,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":31,"verified_exact":1,"verified_fuzzy":20},"total_outbound_references":52},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 52 of 52 outbound references and 1 inbound Pith citation observation for arXiv:2508.14764."}