{"metric_id":"BEMO:2000487","metric_iri":"https://w3id.org/bemo/BEMO_2000487","preferred_label":"Study Design Appropriateness","normalized_label":"study_design_appropriateness","abbreviation":"","pillar_id":"BEMO:1000002","pillar_label":"Study Design, Statistical Validity, and Causal Inference","category_id":"BEMO:1100018","category_label":"Study Design and Internal Validity","parent_class_iri":"https://w3id.org/bemo/BEMO_1100018","ontology_namespace":"BEMO","ontology_version":"0.1.0","metric_version":"1.0.0","lifecycle_status":"Candidate","deprecated":false,"replacement_id":"","created_date":"2026-08-02","modified_date":"2026-08-02","language":"en","license_status":"Pending owner approval; CC BY 4.0 recommended for OBO compatibility.","scientific_definition":"The extent to which study design is sufficient and fit for the stated biomedical inference.","what_it_measures":"Assesses study design appropriateness using evidence appropriate to study design and internal validity, distinguishing random uncertainty from systematic error.","why_it_matters":"Material weakness in study design appropriateness can change the direction, magnitude, certainty, or biological interpretation of the research conclusion.","measurement_criteria":"Clear intended use or estimand; accepted reference or criterion; prespecified acceptance thresholds; independent validation; performance across relevant conditions.","methods_of_assessment":"Structured critical appraisal; quantitative estimation with uncertainty; prespecified thresholds; independent replication; sensitivity and subgroup analyses.","units_or_scale_source":"Ordinal rubric, domain judgment, or normalized score","evidence_level_source":"Result / experiment / study / body of evidence, as applicable","applicable_study_types_source":"Randomized trials, nonrandomized intervention studies, cohort, case-control, cross-sectional studies","related_frameworks_source":"CONSORT; STROBE; RoB 2; ROBINS-I; ROBINS-E; NIH Quality Tools; JBI","closely_related_metrics_source":"Temporal Precedence; Protocol Fidelity; Outcome Ascertainment Validity","common_misinterpretations":"Treating study design appropriateness as interchangeable with overall study quality, or interpreting a favorable value as proof that all other bias and validity domains are satisfactory.","limitations":"Depends on context, operational definition, thresholds, data quality, and evaluator judgment; it should not be used as a stand-alone summary score without domain-level evidence.","scientific_importance":8,"frequency_of_use":"Common","maturity_of_metric":"Established","references_source":"https://www.consort-spirit.org/ | https://www.equator-network.org/reporting-guidelines/strobe/ | https://www.riskofbias.info/welcome/rob-2-0-tool | https://www.riskofbias.info/welcome/home/current-version-of-robins-i | https://www.riskofbias.info/welcome/robins-e-tool | https://www.nhlbi.nih.gov/health-topics/study-quality-assessment-tools | https://jbi.global/critical-appraisal-tools","metric_kind":"AtomicOrCompositeNotYetCurated","output_concept":"BEMO:0000200 Metric Assessment","scale_type":"OrdinalOrNormalizedScoreScale","output_datatype":"xsd:string_or_decimal","unit_ontology_iri":"","unit_text":"Ordinal rubric, domain judgment, or normalized score","minimum_value":null,"maximum_value":null,"null_value":null,"directionality":"ContextDependentDirection","computation_mode":"RuleBasedRubricComputation","computation_readiness":"RubricRequired","formula_status":"ContextSpecificProtocolRequired","human_readable_formula":"Apply a versioned, prespecified domain rubric or validated normalized scoring model.","machine_formula_language":"BEMO-Expression-JSON","machine_formula_expression":"{\"language\":\"BEMO-Expression-JSON\",\"operator\":\"external_protocol\",\"protocolRef\":\"REQUIRED\",\"allowed_outputs\":[\"ordinal_category\",\"normalized_score\"]}","required_inputs":"evidence_records; assessment_context; rubric_version; operational_definition","optional_inputs":"weights; thresholds; expert_adjudication","aggregation_rule":"Not specified in source; must be defined and versioned before composite use.","normalization_method":"None by default; any normalization must be justified, versioned, and validated.","missing_data_policy":"Must be declared before computation; report missingness; no silent imputation; perform sensitivity analysis when material.","uncertainty_required":true,"uncertainty_method":"Confidence interval, credible interval, bootstrap distribution, inter-rater reliability, or sensitivity analysis as scientifically applicable.","confidence_interval_required":false,"threshold_policy":"Prespecify, justify, version, and sensitivity-test thresholds; do not derive and evaluate on the same data without correction.","decision_thresholds":"","quality_control_requirements":"Source provenance; input validation; duplicate control; uncertainty reporting; independent or orthogonal validation where applicable.","calibration_requirement":"Required for normalized scores; inter-rater reliability required for human rubrics.","validation_status":"Source maturity: Established; BEMO computation profile requires independent validation.","benchmark_requirement":"A representative, versioned benchmark set is required before production use.","gold_standard_requirement":"Use an independent reference standard when one exists; document expert-adjudicated alternatives.","external_validation_required":true,"evidence_object_scope":"Result / experiment / study / body of evidence, as applicable","study_type_applicability":"Randomized trials, nonrandomized intervention studies, cohort, case-control, cross-sectional studies","domain_applicability":"Study Design and Internal Validity","required_data_sources":"Context-specific; use source references, registered studies, primary data, and authoritative biomedical resources as applicable.","computational_method_family":"structured critical appraisal or rubric scoring","software_implementation_status":"Specification generated; metric-specific reference implementation pending.","reference_implementation":"scripts/compute_metric.py --metric BEMO:2000487","api_endpoint_template":"/v1/metrics/BEMO:2000487/compute","json_schema_ref":"schemas/bemo-assessment.schema.json","shacl_shape_iri":"https://w3id.org/bemo/shapes/BEMO_2000487_AssessmentShape","provenance_model":"W3C PROV-O","source_workbook":"Pure_Biomedical_Evidence_Research_Metrics_Part_2(3).xlsx","source_sheet":"Biomedical Evidence Metrics","source_row":494,"source_record_hash":"7f9fb8e0c1ccb555a18bd7b6b07b3173d661a6707fab17a21ee632e59503b393","source_references":"https://www.consort-spirit.org/ | https://www.equator-network.org/reporting-guidelines/strobe/ | https://www.riskofbias.info/welcome/rob-2-0-tool | https://www.riskofbias.info/welcome/home/current-version-of-robins-i | https://www.riskofbias.info/welcome/robins-e-tool | https://www.nhlbi.nih.gov/health-topics/study-quality-assessment-tools | https://jbi.global/critical-appraisal-tools","curator":"Unassigned","reviewer":"Unassigned","owner":"BEMO Project","approval_status":"Draft","governance_note":"Candidate term pending scientific, ontology-engineering, and computation-method review.","editor_note":"Source content preserved verbatim; computation metadata is a generated default and must be curated before normative use.","definition_source":"Pure_Biomedical_Evidence_Research_Metrics_Part_2(3).xlsx row 494; see source_references","exact_synonyms":"","broad_synonyms":"","narrow_synonyms":"","related_synonyms":"","xrefs":"https://www.consort-spirit.org/ | https://www.equator-network.org/reporting-guidelines/strobe/ | https://www.riskofbias.info/welcome/rob-2-0-tool | https://www.riskofbias.info/welcome/home/current-version-of-robins-i | https://www.riskofbias.info/welcome/robins-e-tool | https://www.nhlbi.nih.gov/health-topics/study-quality-assessment-tools | https://jbi.global/critical-appraisal-tools","obo_subset":"bemo_study_design_statistical_validity_and_causal_inference","fair_findable":"Provisional: stable local ID assigned; public namespace registration pending.","fair_accessible":"Provisional: package is distributable; permanent public release location pending.","fair_interoperable":"Yes: OWL 2, RDF, SKOS, SHACL, JSON-LD, CSV, and OBO-style identifiers.","fair_reusable":"Provisional: rich metadata supplied; open-license owner approval pending.","confidence_in_metric_definition":"Not yet formally assessed","metric_dependency_status":"Not yet curated","dependency_notes":"Closely related metrics are represented; causal or computational dependencies require expert curation.","implementation_notes":"Exact execution requires a validated computation profile unless formula_status is GenericTemplateDefined.","test_case_status":"Generic validation template provided","example_input_ref":"examples/example_metric_input.json","example_output_ref":"examples/example_metric_output.jsonld","notes":"No source scientific statement was silently altered or replaced.","url_id":"BEMO-2000487","required_inputs_list":["evidence_records","assessment_context","rubric_version","operational_definition"],"optional_inputs_list":["weights","thresholds","expert_adjudication"],"source_record":{"source_workbook":"Pure_Biomedical_Evidence_Research_Metrics_Part_2(3).xlsx","source_sheet":"Biomedical Evidence Metrics","source_row":"494","Category":"Study Design and Internal Validity","Metric":"Study Design Appropriateness","Scientific Definition":"The extent to which study design is sufficient and fit for the stated biomedical inference.","What It Measures":"Assesses study design appropriateness using evidence appropriate to study design and internal validity, distinguishing random uncertainty from systematic error.","Why It Matters":"Material weakness in study design appropriateness can change the direction, magnitude, certainty, or biological interpretation of the research conclusion.","Measurement Criteria":"Clear intended use or estimand; accepted reference or criterion; prespecified acceptance thresholds; independent validation; performance across relevant conditions.","Typical Methods of Assessment":"Structured critical appraisal; quantitative estimation with uncertainty; prespecified thresholds; independent replication; sensitivity and subgroup analyses.","Units or Scale (if applicable)":"Ordinal rubric, domain judgment, or normalized score","Evidence Level Where Used":"Result / experiment / study / body of evidence, as applicable","Applicable Study Types":"Randomized trials, nonrandomized intervention studies, cohort, case-control, cross-sectional studies","Related Frameworks":"CONSORT; STROBE; RoB 2; ROBINS-I; ROBINS-E; NIH Quality Tools; JBI","Closely Related Metrics":"Temporal Precedence; Protocol Fidelity; Outcome Ascertainment Validity","Common Misinterpretations":"Treating study design appropriateness as interchangeable with overall study quality, or interpreting a favorable value as proof that all other bias and validity domains are satisfactory.","Limitations":"Depends on context, operational definition, thresholds, data quality, and evaluator judgment; it should not be used as a stand-alone summary score without domain-level evidence.","Scientific Importance (1–10)":"8","Frequency of Use":"Common","Maturity of Metric":"Established","References or Origin":"https://www.consort-spirit.org/ | https://www.equator-network.org/reporting-guidelines/strobe/ | https://www.riskofbias.info/welcome/rob-2-0-tool | https://www.riskofbias.info/welcome/home/current-version-of-robins-i | https://www.riskofbias.info/welcome/robins-e-tool | https://www.nhlbi.nih.gov/health-topics/study-quality-assessment-tools | https://jbi.global/critical-appraisal-tools"},"computation_profile":{"metric_id":"BEMO:2000487","preferred_label":"Study Design Appropriateness","computation_mode":"RuleBasedRubricComputation","computation_readiness":"RubricRequired","formula_status":"ContextSpecificProtocolRequired","human_readable_formula":"Apply a versioned, prespecified domain rubric or validated normalized scoring model.","machine_formula_language":"BEMO-Expression-JSON","machine_formula_expression":{"language":"BEMO-Expression-JSON","operator":"external_protocol","protocolRef":"REQUIRED","allowed_outputs":["ordinal_category","normalized_score"]},"required_inputs":"evidence_records; assessment_context; rubric_version; operational_definition","optional_inputs":"weights; thresholds; expert_adjudication","output_concept":"BEMO:0000200 Metric Assessment","output_datatype":"xsd:string_or_decimal","scale_type":"OrdinalOrNormalizedScoreScale","unit_ontology_iri":"","unit_text":"Ordinal rubric, domain judgment, or normalized score","minimum_value":null,"maximum_value":null,"null_value":null,"directionality":"ContextDependentDirection","aggregation_rule":"Not specified in source; must be defined and versioned before composite use.","normalization_method":"None by default; any normalization must be justified, versioned, and validated.","missing_data_policy":"Must be declared before computation; report missingness; no silent imputation; perform sensitivity analysis when material.","uncertainty_required":true,"uncertainty_method":"Confidence interval, credible interval, bootstrap distribution, inter-rater reliability, or sensitivity analysis as scientifically applicable.","confidence_interval_required":false,"threshold_policy":"Prespecify, justify, version, and sensitivity-test thresholds; do not derive and evaluate on the same data without correction.","decision_thresholds":"","quality_control_requirements":"Source provenance; input validation; duplicate control; uncertainty reporting; independent or orthogonal validation where applicable.","calibration_requirement":"Required for normalized scores; inter-rater reliability required for human rubrics.","validation_status":"Source maturity: Established; BEMO computation profile requires independent validation.","benchmark_requirement":"A representative, versioned benchmark set is required before production use.","gold_standard_requirement":"Use an independent reference standard when one exists; document expert-adjudicated alternatives.","external_validation_required":true,"computational_method_family":"structured critical appraisal or rubric scoring","software_implementation_status":"Specification generated; metric-specific reference implementation pending.","reference_implementation":"scripts/compute_metric.py --metric BEMO:2000487","api_endpoint_template":"/v1/metrics/BEMO:2000487/compute","json_schema_ref":"schemas/bemo-assessment.schema.json","shacl_shape_iri":"https://w3id.org/bemo/shapes/BEMO_2000487_AssessmentShape","profile_version":"0.1.0","profile_iri":"https://w3id.org/bemo/profile/BEMO_2000487"},"relationships":[{"predicate":"closely_related_metric","source":"Source workbook: Closely Related Metrics","status":"SourceDeclared","directionality":"Directed as recorded; may be symmetric in the source","version":"0.1.0","metric_id":"BEMO:2000488","preferred_label":"Temporal Precedence"},{"predicate":"closely_related_metric","source":"Source workbook: Closely Related Metrics","status":"SourceDeclared","directionality":"Directed as recorded; may be symmetric in the source","version":"0.1.0","metric_id":"BEMO:2000484","preferred_label":"Protocol Fidelity"},{"predicate":"closely_related_metric","source":"Source workbook: Closely Related Metrics","status":"SourceDeclared","directionality":"Directed as recorded; may be symmetric in the source","version":"0.1.0","metric_id":"BEMO:2000480","preferred_label":"Outcome Ascertainment Validity"}],"incoming_relationships":[{"predicate":"closely_related_metric","source":"Source workbook: Closely Related Metrics","status":"SourceDeclared","directionality":"Directed as recorded; may be symmetric in the source","version":"0.1.0","metric_id":"BEMO:2000473","preferred_label":"Control Group Appropriateness"},{"predicate":"closely_related_metric","source":"Source workbook: Closely Related Metrics","status":"SourceDeclared","directionality":"Directed as recorded; may be symmetric in the source","version":"0.1.0","metric_id":"BEMO:2000484","preferred_label":"Protocol Fidelity"},{"predicate":"closely_related_metric","source":"Source workbook: Closely Related Metrics","status":"SourceDeclared","directionality":"Directed as recorded; may be symmetric in the source","version":"0.1.0","metric_id":"BEMO:2000488","preferred_label":"Temporal Precedence"}],"references":[{"order":1,"url":"https://www.consort-spirit.org/","type":"Source framework, standard, guideline, or origin","verification_status":"Source-preserved; URL reachability not asserted by this package"},{"order":2,"url":"https://www.equator-network.org/reporting-guidelines/strobe/","type":"Source framework, standard, guideline, or origin","verification_status":"Source-preserved; URL reachability not asserted by this package"},{"order":3,"url":"https://www.riskofbias.info/welcome/rob-2-0-tool","type":"Source framework, standard, guideline, or origin","verification_status":"Source-preserved; URL reachability not asserted by this package"},{"order":4,"url":"https://www.riskofbias.info/welcome/home/current-version-of-robins-i","type":"Source framework, standard, guideline, or origin","verification_status":"Source-preserved; URL reachability not asserted by this package"},{"order":5,"url":"https://www.riskofbias.info/welcome/robins-e-tool","type":"Source framework, standard, guideline, or origin","verification_status":"Source-preserved; URL reachability not asserted by this package"},{"order":6,"url":"https://www.nhlbi.nih.gov/health-topics/study-quality-assessment-tools","type":"Source framework, standard, guideline, or origin","verification_status":"Source-preserved; URL reachability not asserted by this package"},{"order":7,"url":"https://jbi.global/critical-appraisal-tools","type":"Source framework, standard, guideline, or origin","verification_status":"Source-preserved; URL reachability not asserted by this package"}]}