{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:KX3FKXPDOV3P367MI65Z2HJAQ2","short_pith_number":"pith:KX3FKXPD","schema_version":"1.0","canonical_sha256":"55f6555de37576fdfbec47bb9d1d2086bf9f2d03f3f25ac3b6175196c6136645","source":{"kind":"arxiv","id":"2506.12937","version":2},"attestation_state":"computed","paper":{"title":"HypER: Literature-grounded Hypothesis Generation and Distillation with Provenance","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.CL"],"primary_cat":"cs.AI","authors_text":"Abraham Bernstein, Bhavana Dalvi Mishra, Chandrayee Basu, Cristina Sarasua, Peter Clark, Rosni Vasu","submitted_at":"2025-06-15T18:41:23Z","abstract_excerpt":"Large Language models have demonstrated promising performance in research ideation across scientific domains. Hypothesis development, the process of generating a highly specific declarative statement connecting a research idea with empirical validation, has received relatively less attention. Existing approaches trivially deploy retrieval augmentation and focus only on the quality of the final output ignoring the underlying reasoning process behind ideation. We present $\\texttt{HypER}$ ($\\textbf{Hyp}$othesis Generation with $\\textbf{E}$xplanation and $\\textbf{R}$easoning), a small language mod"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2506.12937","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2025-06-15T18:41:23Z","cross_cats_sorted":["cs.CL"],"title_canon_sha256":"3fee86f18145fb1e638f95d92d5918e3b4d1e7d58cd5b22aeb2dc59f7ffcdc83","abstract_canon_sha256":"5f1152d77d1541cee2f79e66319f556e6215e37e14ebeb8f622285b4123f9cad"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:57:20.529629Z","signature_b64":"KMS6lAKz6FMNJRzrYoN+sDAYavSXOE7zyd72nkdsLOgM0BiVR/68ScyrGJPhd/Eh3HuYvNS9NvumV/xKoRZIAQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"55f6555de37576fdfbec47bb9d1d2086bf9f2d03f3f25ac3b6175196c6136645","last_reissued_at":"2026-07-05T11:57:20.529138Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:57:20.529138Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"HypER: Literature-grounded Hypothesis Generation and Distillation with Provenance","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.CL"],"primary_cat":"cs.AI","authors_text":"Abraham Bernstein, Bhavana Dalvi Mishra, Chandrayee Basu, Cristina Sarasua, Peter Clark, Rosni Vasu","submitted_at":"2025-06-15T18:41:23Z","abstract_excerpt":"Large Language models have demonstrated promising performance in research ideation across scientific domains. Hypothesis development, the process of generating a highly specific declarative statement connecting a research idea with empirical validation, has received relatively less attention. Existing approaches trivially deploy retrieval augmentation and focus only on the quality of the final output ignoring the underlying reasoning process behind ideation. We present $\\texttt{HypER}$ ($\\textbf{Hyp}$othesis Generation with $\\textbf{E}$xplanation and $\\textbf{R}$easoning), a small language mod"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2506.12937","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2506.12937/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2506.12937","created_at":"2026-07-05T11:57:20.529198+00:00"},{"alias_kind":"arxiv_version","alias_value":"2506.12937v2","created_at":"2026-07-05T11:57:20.529198+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2506.12937","created_at":"2026-07-05T11:57:20.529198+00:00"},{"alias_kind":"pith_short_12","alias_value":"KX3FKXPDOV3P","created_at":"2026-07-05T11:57:20.529198+00:00"},{"alias_kind":"pith_short_16","alias_value":"KX3FKXPDOV3P367M","created_at":"2026-07-05T11:57:20.529198+00:00"},{"alias_kind":"pith_short_8","alias_value":"KX3FKXPD","created_at":"2026-07-05T11:57:20.529198+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2510.21652","citing_title":"AstaBench: Rigorous Benchmarking of AI Agents with a Scientific Research Suite","ref_index":2,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2","json":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2.json","graph_json":"https://pith.science/api/pith-number/KX3FKXPDOV3P367MI65Z2HJAQ2/graph.json","events_json":"https://pith.science/api/pith-number/KX3FKXPDOV3P367MI65Z2HJAQ2/events.json","paper":"https://pith.science/paper/KX3FKXPD"},"agent_actions":{"view_html":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2","download_json":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2.json","view_paper":"https://pith.science/paper/KX3FKXPD","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2506.12937&json=true","fetch_graph":"https://pith.science/api/pith-number/KX3FKXPDOV3P367MI65Z2HJAQ2/graph.json","fetch_events":"https://pith.science/api/pith-number/KX3FKXPDOV3P367MI65Z2HJAQ2/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2/action/timestamp_anchor","attest_storage":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2/action/storage_attestation","attest_author":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2/action/author_attestation","sign_citation":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2/action/citation_signature","submit_replication":"https://pith.science/pith/KX3FKXPDOV3P367MI65Z2HJAQ2/action/replication_record"}},"created_at":"2026-07-05T11:57:20.529198+00:00","updated_at":"2026-07-05T11:57:20.529198+00:00"}