{
  "schema_version": 1,
  "updated_at": "2026-09-15",
  "name": "Brali Evidence Ledger",
  "description": "Reviewed source-by-source decisions showing what evidence supports, what it does not establish, its limitations, and how Brali changes guidance in response.",
  "methodology_url": "https://brali-lifeos.github.io/life-os/methodology/",
  "evidence_dataset_url": "https://brali-lifeos.github.io/life-os/datasets/evidence-decisions.json",
  "count": 63,
  "decision_counts": {
    "propose-protocol": 29,
    "support-existing": 20,
    "challenge-existing": 8,
    "watch": 6
  },
  "unsupported_or_overstated_claim_count": 381,
  "entries": [
    {
      "schema_version": 1,
      "id": "prequestions-targeted-learning-boundary-2025",
      "canonical_url": "https://brali-lifeos.github.io/evidence/prequestions-targeted-learning-boundary-2025/",
      "json_url": "https://brali-lifeos.github.io/evidence/prequestions-targeted-learning-boundary-2025/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-15",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The Effect of Prequestions on Learning: A Multilevel Meta-Analysis",
        "url": "https://doi.org/10.1007/s10648-025-10075-7",
        "type": "meta-analysis",
        "doi": "10.1007/s10648-025-10075-7",
        "citation_text": "King-Shepard QW, Walker J, Nokes-Malach TJ, Carpenter SK, Fraundorf SH. Educational Psychology Review. 2025;37:115.",
        "study_design": "Multilevel meta-analysis of the prequestion literature examining learning of prequestioned and non-prequestioned content across variations in learning event, participant sample, and assessment conditions.",
        "population": "Learners represented across the studies included in the meta-analysis; the practical protocol is intentionally population-neutral and does not convert subgroup findings into a universal dose.",
        "intervention_or_exposure": "Questions attempted before an educational activity, compared with conditions without those prequestions; some study conditions also provided feedback.",
        "outcomes": [
          "Learning of prequestioned information",
          "Learning of non-prequestioned information"
        ]
      },
      "supported_claim": "Prequestions can improve later learning of the specific information they target. The meta-analysis reported g = 0.66 for prequestioned information and g = 0.01 for non-prequestioned information; feedback alongside prequestions was associated with stronger targeted learning.",
      "unsupported_or_overstated_claims": [
        "Prequestions improve learning of the entire lesson.",
        "Any question asked before learning will be beneficial.",
        "There is one evidence-backed optimal number of prequestions.",
        "A wrong first answer is itself sufficient without later exposure or correction.",
        "Prequestions replace retrieval practice, practice problems, or feedback after learning."
      ],
      "limitations": [
        "The average benefit is specific to prequestioned information rather than a broad benefit to all content.",
        "The meta-analysis does not establish one universal question count, timing rule, or content format.",
        "Effect sizes summarize heterogeneous studies and do not guarantee the same result for an individual learning session."
      ],
      "target_hack_ids": [
        "prequestions-before-learning"
      ],
      "target_protocol_ids": [
        "brali:prequestions-before-learning"
      ],
      "risk_flags": [],
      "notes": "Editorial action: publish a new Skill Sprint protocol because no prequestion/pretesting surface was found in the canonical index or override layer. Make the target-versus-nontarget boundary prominent."
    },
    {
      "schema_version": 1,
      "id": "self-explanation-digital-learning-boundary-2025",
      "canonical_url": "https://brali-lifeos.github.io/evidence/self-explanation-digital-learning-boundary-2025/",
      "json_url": "https://brali-lifeos.github.io/evidence/self-explanation-digital-learning-boundary-2025/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-15",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Enhancing Academic Performance Through Self-Explanation in Digital Learning Environments (DLEs): A Three-Level Meta-Analysis",
        "url": "https://doi.org/10.1007/s10648-025-10001-x",
        "type": "meta-analysis",
        "doi": "10.1007/s10648-025-10001-x",
        "citation_text": "Tan LP, Gong SY, Wang YJ, Guo XR, Xu XZ, Wang YQ. Educational Psychology Review. 2025;37:20.",
        "study_design": "Three-level meta-analysis of 204 effect sizes extracted from 56 studies comparing self-explanation with no-self-explanation conditions in digital learning environments.",
        "population": "Learners in digital learning environments across the included studies. Effects varied with learning-environment and material characteristics.",
        "intervention_or_exposure": "Self-explanation prompts or activities during digital learning, compared with conditions without self-explanation.",
        "outcomes": [
          "Academic performance",
          "Retention",
          "Transfer"
        ]
      },
      "supported_claim": "Self-explanation in digital learning environments had a positive average effect on academic performance (reported g = 0.46), with positive average effects for retention and transfer. Effects were moderated by pacing and knowledge type, with larger effects reported for learner-centered pacing and conceptual or mixed knowledge than for procedural knowledge.",
      "unsupported_or_overstated_claims": [
        "Self-explanation is always better than other learning strategies.",
        "One explanation prompt or modality is universally optimal.",
        "Every paragraph should be self-explained.",
        "The average effect applies unchanged outside digital learning environments.",
        "A fluent explanation is accurate without checking it against the source."
      ],
      "limitations": [
        "The reviewed evidence is specifically about digital learning environments.",
        "The meta-analysis reports heterogeneity and moderation rather than a uniform effect across all settings.",
        "The authors identify the quality and standardization of self-explanations as areas needing further research."
      ],
      "target_hack_ids": [
        "self-explain-what-you-learn"
      ],
      "target_protocol_ids": [
        "brali:self-explain-what-you-learn"
      ],
      "risk_flags": [],
      "notes": "Editorial action: publish a new Skill Sprint protocol built around generate -> inspect gap -> source check -> correct. Preserve digital-learning scope and moderation rather than presenting self-explanation as a universal study rule."
    },
    {
      "schema_version": 1,
      "id": "green-route-stress-walk-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/green-route-stress-walk-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/green-route-stress-walk-boundary-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-13",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Effects of green exercise on mental health: a systematic review and meta-analysis",
        "url": "https://www.frontiersin.org/journals/psychology/articles/10.3389/fpsyg.2026.1802759/full",
        "type": "systematic-review-meta-analysis",
        "doi": "10.3389/fpsyg.2026.1802759",
        "citation_text": "Liu X, Sun Z, Wang X, Dong D, Samsudin SB. Front Psychol. 2026;17:1802759.",
        "study_design": "PROSPERO-registered systematic review and random-effects meta-analysis restricted to peer-reviewed randomized controlled trials and crossover RCTs. PubMed, Web of Science, EBSCOhost, PsycINFO and CENTRAL were searched through 7 November 2025. The review included 51 studies, used duplicate screening/data extraction, PEDro study-quality assessment, GRADE certainty assessment and sensitivity analyses.",
        "population": "3,092 healthy adults aged 18-85 across 51 studies. Ten studies included only men, three only women and 38 included both. Most interventions were walking; the review authors state that the findings primarily reflect low-intensity walking in natural environments.",
        "intervention_or_exposure": "Physical activity performed in environments containing natural green elements, compared with non-exercise, indoor exercise or outdoor exercise in built-up environments with minimal greenery. Intervention formats and settings varied; most studies were short-term and follow-up was generally under three months.",
        "outcomes": [
          "Overall wellbeing",
          "Positive and negative affect",
          "Stress, anxiety and depression symptom scores",
          "Calm, vigor, anger, fatigue and confusion"
        ]
      },
      "supported_claim": "For an existing Brali stress-walk protocol, a safe greener route can be offered as a low-friction optional modifier when it is as convenient as an ordinary route. In the reviewed randomized evidence, green exercise improved average wellbeing and affect relative to pooled non-exercise, indoor-exercise and built-up outdoor comparators. Stress-specific pooled effects favored green exercise in all three comparator categories. This supports 'prefer green when equally practical', not a stronger prescription.",
      "unsupported_or_overstated_claims": [
        "A park, forest or specific amount of vegetation is required for a useful reset walk.",
        "Ten minutes, or any other fixed duration, is the evidence-based optimal dose for stress relief.",
        "Green exercise guarantees an immediate reduction in stress for an individual.",
        "The pooled effects prove that greenery itself, rather than movement, context, expectations or their interaction, uniquely caused the benefit.",
        "Green walking reliably treats depression, anxiety disorders, burnout or another mental-health condition.",
        "The review establishes durable long-term mental-health benefits after the intervention ends.",
        "The evidence supports a specific pace, heart-rate target, breathing pattern, mindfulness routine or screen rule.",
        "Subgroup or setting differences should be converted into a ranked route prescription."
      ],
      "limitations": [
        "Several pooled outcomes showed substantial or high heterogeneity, although sensitivity analyses suggested some results were driven by a small number of studies.",
        "Most individual studies were small, and the authors noted incomplete reporting of randomization procedures and limited blinding in some trials.",
        "Most interventions were walking, so generalization to other forms or intensities of exercise is limited.",
        "Most interventions were short-term, with follow-up generally under three months, so long-term durability is uncertain.",
        "Settings, exercise formats and psychological measures varied across studies.",
        "Stress-specific comparisons contained relatively few studies: three versus non-exercise, six versus indoor exercise and two versus built-up exercise.",
        "The meta-analysis does not isolate a unique biological or psychological mechanism for any observed advantage of greener settings."
      ],
      "target_hack_ids": [
        "10-minute-stress-walk"
      ],
      "target_protocol_ids": [
        "brali:10-minute-stress-walk"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Editorial action: update the existing reviewed stress-walk page rather than create a second green-walk hack. Preserve the anti-magic-timer framing, make safety and convenience primary, add 'prefer a greener route when it costs essentially nothing' as an optional modifier, and add a simple observation loop comparing ordinary versus greener routes without presenting N-of-1 impressions as causal proof."
    },
    {
      "schema_version": 1,
      "id": "consider-opposite-social-judgment-boundary-1984",
      "canonical_url": "https://brali-lifeos.github.io/evidence/consider-opposite-social-judgment-boundary-1984/",
      "json_url": "https://brali-lifeos.github.io/evidence/consider-opposite-social-judgment-boundary-1984/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Considering the opposite: A corrective strategy for social judgment",
        "url": "https://doi.org/10.1037/0022-3514.47.6.1231",
        "type": "primary-study",
        "doi": "10.1037/0022-3514.47.6.1231",
        "citation_text": "Lord CG, Lepper MR, Preston E. J Pers Soc Psychol. 1984;47(6):1231-1243.",
        "study_design": "Two conceptually parallel experiments manipulating whether participants considered possibilities opposed to their current beliefs or impressions, compared with alternative instructions including generic fairness/unbiasedness prompts.",
        "population": "150 undergraduate participants across two experiments.",
        "intervention_or_exposure": "Explicit consider-the-opposite instructions or materials designed to make opposite possibilities salient during evaluation of ambiguous evidence and hypothesis testing about personality impressions.",
        "outcomes": [
          "Biased assimilation of evidence",
          "Attitude polarization",
          "Biased hypothesis testing in social judgment"
        ]
      },
      "supported_claim": "When a judgment is vulnerable to one-sided evidence processing, deliberately generating an opposed possibility can reduce bias on some tasks more effectively than simply telling oneself to be fair or unbiased. Brali can use this as a concrete pre-decision check while preserving the possibility that the original conclusion remains correct.",
      "unsupported_or_overstated_claims": [
        "Consider-the-opposite eliminates confirmation bias.",
        "The technique transfers automatically to every real-world decision.",
        "The opposite conclusion should be preferred once generated.",
        "More counterarguments are always better.",
        "The strategy should be used for every trivial or reversible choice."
      ],
      "limitations": [
        "Classic laboratory/social-judgment evidence from undergraduate samples.",
        "Only two focal task domains were tested in the original article.",
        "Long-term persistence and broad transfer were not established.",
        "The authors note that considering the opposite can in some circumstances overweight disconfirming evidence and create a different form of partiality.",
        "Demand characteristics and task-specific effects remain possible."
      ],
      "target_hack_ids": [
        "consider-the-opposite"
      ],
      "target_protocol_ids": [
        "brali:consider-the-opposite"
      ],
      "risk_flags": [],
      "notes": "Publish as a single structured alternative-generation step followed by a discriminating-evidence question. Do not tell users to reverse their decision by default."
    },
    {
      "schema_version": 1,
      "id": "deadline-evidence-reversal-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/deadline-evidence-reversal-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/deadline-evidence-reversal-2026/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Replication of Procrastination, Deadlines, and Performance: Self-Control by Precommitment",
        "url": "https://doi.org/10.1177/09567976261460772",
        "type": "replication-study",
        "doi": "10.1177/09567976261460772",
        "citation_text": null,
        "study_design": null,
        "population": null,
        "intervention_or_exposure": null,
        "outcomes": []
      },
      "supported_claim": "The 2026 replication did not reproduce the classic deadline effect, while the publisher now marks the 2002 source retracted.",
      "unsupported_or_overstated_claims": [
        "Artificial deadlines are a settled evidence-based productivity intervention.",
        "Retraction proves that deadlines never matter in any context."
      ],
      "limitations": [
        "The 2026 paper is a replication of a specific deadline experiment and does not test every form of deadline, planning, precommitment, or time-management intervention.",
        "Failure to reproduce the classic effect lowers confidence in that specific evidence claim; it does not establish that deadlines have no effect in every task, population, incentive structure, or real-world setting.",
        "The publisher's retraction of the 2002 paper means Brali should not use that paper as positive evidence, but retraction alone is not evidence for the opposite universal claim.",
        "This decision is a provenance and evidence-boundary correction, not a complete systematic review of the broader deadline literature."
      ],
      "target_hack_ids": [],
      "target_protocol_ids": [],
      "risk_flags": [],
      "notes": "Block any new Brali hack whose scientific rationale rests mainly on Ariely & Wertenbroch (2002), and audit any existing deadline guidance for provenance."
    },
    {
      "schema_version": 1,
      "id": "debiasing-education-transfer-boundary-2025",
      "canonical_url": "https://brali-lifeos.github.io/evidence/debiasing-education-transfer-boundary-2025/",
      "json_url": "https://brali-lifeos.github.io/evidence/debiasing-education-transfer-boundary-2025/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Systematic review and meta-analysis of educational approaches to reduce cognitive biases among students",
        "url": "https://www.nature.com/articles/s41562-025-02253-y",
        "type": "systematic-review-meta-analysis",
        "doi": "10.1038/s41562-025-02253-y",
        "citation_text": "Swaryandini G, Graham J, Griffith S, et al. Nat Hum Behav. 2025;9:2510-2538.",
        "study_design": "Systematic review of 54 randomized controlled trials with 383 effect sizes and meta-analysis of 160 effects from 41 studies evaluating educational interventions intended to reduce cognitive biases.",
        "population": "10,941 participants across the randomized studies, primarily students in educational settings and spanning multiple targeted cognitive biases.",
        "intervention_or_exposure": "Educational debiasing approaches including cognitive strategies, feedback, games, instructional materials and multi-component training; some included consider-the-opposite components.",
        "outcomes": [
          "Performance on targeted cognitive-bias tasks",
          "Bias reduction after educational interventions",
          "Evidence about retention and transfer where reported"
        ]
      },
      "supported_claim": "Debiasing education can produce a small average improvement on targeted bias tasks. This supports modest expectations for a consider-the-opposite check while making the boundary explicit: depth of learning and transfer to meaningful real-world decisions remain uncertain.",
      "unsupported_or_overstated_claims": [
        "A brief consider-the-opposite prompt has a meta-analytically proven effect size of g 0.26.",
        "Debiasing training reliably transfers to professional and personal decisions.",
        "All cognitive biases respond similarly to training.",
        "One educational strategy was established as universally best.",
        "A small targeted-task effect implies large practical decision improvements."
      ],
      "limitations": [
        "All included studies were rated unclear or high risk of bias.",
        "There was some evidence of publication bias.",
        "Interventions, bias targets and outcome measures were highly heterogeneous.",
        "Transfer beyond explicitly trained tasks was limited or uncertain in many studies.",
        "The pooled effect represents diverse educational interventions rather than the consider-the-opposite strategy alone."
      ],
      "target_hack_ids": [
        "consider-the-opposite"
      ],
      "target_protocol_ids": [
        "brali:consider-the-opposite"
      ],
      "risk_flags": [],
      "notes": "Use as a calibration source: the direction is encouraging, the average effect is small, and real-world transfer must remain an open question."
    },
    {
      "schema_version": 1,
      "id": "email-checking-frequency-stress-boundary-2015",
      "canonical_url": "https://brali-lifeos.github.io/evidence/email-checking-frequency-stress-boundary-2015/",
      "json_url": "https://brali-lifeos.github.io/evidence/email-checking-frequency-stress-boundary-2015/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Checking email less frequently reduces stress",
        "url": "https://doi.org/10.1016/j.chb.2014.11.005",
        "type": "primary-study",
        "doi": "10.1016/j.chb.2014.11.005",
        "citation_text": "Kushlev K, Dunn EW. Comput Hum Behav. 2015;43:220-228.",
        "study_design": "Two-week within-subject experimental field study with randomized condition order. Each participant completed a limited-email week and an unlimited-email week.",
        "population": "124 adult email users participating in a field experiment over two weeks.",
        "intervention_or_exposure": "During the limited condition participants were instructed to check email only three times per day and keep mail software closed otherwise; during the unlimited condition they were instructed to check as often as possible.",
        "outcomes": [
          "Daily perceived stress",
          "Tension during an important activity",
          "Self-reported email-checking frequency",
          "Exploratory wellbeing, mindfulness, productivity, sleep and social outcomes"
        ]
      },
      "supported_claim": "When role expectations allow it, checking email less frequently than one's normal pattern can reduce daily stress. A practical Brali implementation is to use deliberate email windows while keeping a separate route for genuinely urgent work. The study supports less-frequent checking, not three checks per day as an optimized rule.",
      "unsupported_or_overstated_claims": [
        "Three email checks per day is the optimal schedule.",
        "Less-frequent email checking directly causes higher productivity, better sleep, more meaning, or greater social connectedness.",
        "Every occupation can safely close email between windows.",
        "Inbox Zero or a particular email-processing method was tested.",
        "The exact self-reported number of daily checks can be treated as objectively measured behavior."
      ],
      "limitations": [
        "Checking frequency was self-reported rather than objectively logged.",
        "The study explored a broad set of outcomes, increasing the chance of isolated significant results.",
        "Direct causal evidence was strongest for stress; broader wellbeing links were indirect correlational analyses through stress.",
        "There was no passive measurement-only control condition.",
        "The experiment did not control the response expectations imposed by participants' workplaces and contacts."
      ],
      "target_hack_ids": [
        "deliberate-email-checking-windows"
      ],
      "target_protocol_ids": [
        "brali:deliberate-email-checking-windows"
      ],
      "risk_flags": [],
      "notes": "Publish as a work-system experiment with explicit operational exceptions. Measure stress and missed obligations, not inbox aesthetics or unsupported productivity claims."
    },
    {
      "schema_version": 1,
      "id": "grayscale-screen-time-friction-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/grayscale-screen-time-friction-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/grayscale-screen-time-friction-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "A cross-over feasibility trial of smartphone grayscale mode in medical students",
        "url": "https://www.frontiersin.org/journals/digital-health/articles/10.3389/fdgth.2026.1816095/full",
        "type": "primary-study",
        "doi": "10.3389/fdgth.2026.1816095",
        "citation_text": "Hagerty J, Saretha S, Phung B, et al. Front Digit Health. 2026;8:1816095.",
        "study_design": "Two-week cross-over feasibility trial with consecutive one-week grayscale and color conditions. Assignment used a balanced alternating sequence based on enrollment order; participants served as their own controls. Mixed-effects models analyzed repeated daily screen-time observations.",
        "population": "First-year medical students at one US allopathic medical school. Sixty-one enrolled and 51 completed both study phases.",
        "intervention_or_exposure": "Smartphone display set to grayscale for one week versus normal color for one week. Participants transcribed daily screen-time values from their device's built-in tracking feature and completed exploratory end-of-week surveys.",
        "outcomes": [
          "Daily smartphone screen time",
          "Exploratory self-reported productivity",
          "Exploratory self-reported sleep quality",
          "Feasibility and qualitative usability feedback"
        ]
      },
      "supported_claim": "For a person who already wants to reduce habitual smartphone use, grayscale is a reasonable low-cost, reversible experiment. In this small short cross-over trial, grayscale was associated with about 28 fewer minutes of daily phone screen time on average. The result supports a screen-time-friction protocol, not a productivity, sleep, mental-health, or durable behavior-change claim.",
      "unsupported_or_overstated_claims": [
        "Grayscale reliably improves productivity or sleep.",
        "Grayscale treats problematic smartphone use or mental-health conditions.",
        "Everyone will reduce phone use by about 28 minutes per day.",
        "One week is an optimal or necessary duration.",
        "Lower phone screen time necessarily means less total digital distraction rather than substitution to another device."
      ],
      "limitations": [
        "Small single-school sample of first-year medical students.",
        "Short two-week study with one week per condition and no washout period.",
        "Participants and investigators could not blind the visible display condition.",
        "Screen-time values were participant-reported from built-in device trackers without screenshot verification.",
        "Exploratory productivity and sleep measures were non-validated and not powered for formal inference.",
        "Qualitative comments described color-dependent usability problems and substitution to other devices."
      ],
      "target_hack_ids": [
        "grayscale-phone-friction"
      ],
      "target_protocol_ids": [
        "brali:grayscale-phone-friction"
      ],
      "risk_flags": [],
      "notes": "Publish as a reversible interface nudge. Track the actual target behavior and substitution, allow color for functional tasks, and do not use broad digital-detox language."
    },
    {
      "schema_version": 1,
      "id": "implementation-intentions-meta-boundary-2024",
      "canonical_url": "https://brali-lifeos.github.io/evidence/implementation-intentions-meta-boundary-2024/",
      "json_url": "https://brali-lifeos.github.io/evidence/implementation-intentions-meta-boundary-2024/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The when and how of planning: Meta-analysis of implementation intentions in 642 tests",
        "url": "https://doi.org/10.1080/10463283.2024.2334563",
        "type": "meta-analysis",
        "doi": "10.1080/10463283.2024.2334563",
        "citation_text": null,
        "study_design": null,
        "population": null,
        "intervention_or_exposure": null,
        "outcomes": []
      },
      "supported_claim": "Cue-response implementation intentions are a well-supported self-regulation mechanism; contingent if-then formats performed better on average in the reviewed literature.",
      "unsupported_or_overstated_claims": [
        "Every if-then sentence guarantees behavior change.",
        "Aggregate effect sizes apply to every productivity task."
      ],
      "limitations": [
        "The meta-analysis aggregates many populations, goals, outcomes and implementation-intention variants, so its average effect is not a guaranteed effect size for one productivity task or one person.",
        "Reported moderators such as motivation, contingent if-then form and rehearsal describe variation across the evidence base; they do not establish one universally optimal Brali recipe.",
        "The evidence supports the implementation-intention mechanism broadly, not fixed numbers of plans, fixed rehearsal counts, fixed time windows, or productivity estimates.",
        "The reviewed synthesis does not validate Brali's exact wording, interface or reminder implementation as a separately tested intervention."
      ],
      "target_hack_ids": [
        "if-then-trigger-plan"
      ],
      "target_protocol_ids": [
        "brali:if-then-trigger-plan"
      ],
      "risk_flags": [],
      "notes": "Merge into existing WOOP/task-initiation surfaces when overlapping rather than creating duplicate content. Keep the executable form explicit: If [observable cue], then I will [specific action]."
    },
    {
      "schema_version": 1,
      "id": "micro-conversation-belonging-affect-boundary-2014",
      "canonical_url": "https://brali-lifeos.github.io/evidence/micro-conversation-belonging-affect-boundary-2014/",
      "json_url": "https://brali-lifeos.github.io/evidence/micro-conversation-belonging-affect-boundary-2014/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Is Efficiency Overrated?: Minimal Social Interactions Lead to Belonging and Positive Affect",
        "url": "https://doi.org/10.1177/1948550613502990",
        "type": "primary-study",
        "doi": "10.1177/1948550613502990",
        "citation_text": "Sandstrom GM, Dunn EW. Soc Psychol Personal Sci. 2014;5(4):437-442.",
        "study_design": "Randomized real-world field experiment at a busy urban coffee shop. Participants were assigned to make an already-required cashier interaction either social or maximally efficient; a research assistant blind to condition collected immediate outcomes after the purchase.",
        "population": "60 adult coffee-shop customers recruited in person from a busy urban shopping district, spanning multiple adult age groups.",
        "intervention_or_exposure": "A brief social service interaction involving smiling, eye contact and conversation versus an efficient transaction with unnecessary conversation avoided.",
        "outcomes": [
          "Immediate positive affect",
          "Immediate negative affect",
          "Satisfaction with the service experience",
          "Sense of belonging"
        ]
      },
      "supported_claim": "In a low-pressure context where interaction is already occurring and the other person is able to engage, adding a brief genuine social exchange can improve immediate affect, satisfaction and belonging relative to treating the exchange as purely efficient. Brali can offer this as a context-sensitive social micro-practice, not a loneliness or mental-health intervention.",
      "unsupported_or_overstated_claims": [
        "Brief stranger conversations treat loneliness or social anxiety.",
        "Everyone benefits from talking to strangers.",
        "Service workers should be expected to provide conversation.",
        "The effect generalizes to rushed, unsafe, high-pressure, culturally inappropriate or unwanted interactions.",
        "One brief interaction produces durable relationship or wellbeing changes."
      ],
      "limitations": [
        "Small single-site sample of 60 participants.",
        "Immediate self-reported outcomes rather than long-term behavior or wellbeing.",
        "The context was a coffee shop where friendly customer interaction was normal.",
        "Participants could infer the social purpose of the manipulation, creating possible demand characteristics.",
        "The authors explicitly warn that other service contexts may not support the same kind of interaction."
      ],
      "target_hack_ids": [
        "genuine-micro-conversation"
      ],
      "target_protocol_ids": [
        "brali:genuine-micro-conversation"
      ],
      "risk_flags": [],
      "notes": "Publish with an explicit willingness/context gate. The action is an optional extension of an existing exchange, not a quota of strangers or a treatment claim."
    },
    {
      "schema_version": 1,
      "id": "notification-batching-boundary-2019",
      "canonical_url": "https://brali-lifeos.github.io/evidence/notification-batching-boundary-2019/",
      "json_url": "https://brali-lifeos.github.io/evidence/notification-batching-boundary-2019/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Batching smartphone notifications can improve well-being",
        "url": "https://doi.org/10.1016/j.chb.2019.07.016",
        "type": "primary-study",
        "doi": "10.1016/j.chb.2019.07.016",
        "citation_text": "Fitz N, Kushlev K, Jagannathan R, Lewis T, Paliwal D, Ariely D. Comput Hum Behav. 2019;101:84-94.",
        "study_design": "Two-week randomized field experiment using a custom smartphone application. Participants were assigned to usual notification delivery, hourly batching, three daily batches, or no notifications.",
        "population": "237 adults recruited through Amazon Mechanical Turk for a smartphone field experiment.",
        "intervention_or_exposure": "Notification delivery was changed from the phone's usual variable stream to predictable batches, either hourly or three times daily, or removed entirely in the no-notification condition.",
        "outcomes": [
          "Perceived notification interruptions",
          "Stress",
          "Attention",
          "Mood",
          "Perceived productivity",
          "Sense of control over the phone",
          "Anxiety and fear of missing out",
          "Phone-use behavior"
        ]
      },
      "supported_claim": "Predictable batching of non-urgent notifications is a defensible attention-environment experiment. In this field trial, three daily batches reduced perceived interruption and stress relative to usual delivery, while hourly batching changed little and complete notification removal increased anxiety/FoMO. Brali should preserve urgent exceptions and treat schedule frequency as a user-fit parameter rather than a fixed dose.",
      "unsupported_or_overstated_claims": [
        "Three notification batches per day is universally optimal.",
        "Hourly batching never helps anyone.",
        "Turning all notifications off is generally harmful.",
        "Notification batching reliably increases objective productivity or long-term wellbeing.",
        "The result automatically generalizes to every operating system, notification type, occupation, or family situation."
      ],
      "limitations": [
        "Single short field experiment rather than a replicated long-term evidence base.",
        "Participants were recruited through an online labor market and may not represent all smartphone users or work contexts.",
        "The intervention used a custom notification-management implementation.",
        "Many psychological outcomes were self-reported and the study intentionally used broad exploratory measurement.",
        "The study compared a small set of schedules and cannot identify an optimal individualized batching frequency."
      ],
      "target_hack_ids": [
        "batch-non-urgent-notifications"
      ],
      "target_protocol_ids": [
        "brali:batch-non-urgent-notifications"
      ],
      "risk_flags": [],
      "notes": "Publish the mechanism, not the number: separate urgent channels, batch the rest predictably, observe manual checking and missed information, and avoid total-silence superiority claims."
    },
    {
      "schema_version": 1,
      "id": "notification-blocking-workday-boundary-2023",
      "canonical_url": "https://brali-lifeos.github.io/evidence/notification-blocking-workday-boundary-2023/",
      "json_url": "https://brali-lifeos.github.io/evidence/notification-blocking-workday-boundary-2023/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Effects of task interruptions caused by notifications from communication applications on strain and performance",
        "url": "https://doi.org/10.1002/1348-9585.12408",
        "type": "primary-study",
        "doi": "10.1002/1348-9585.12408",
        "citation_text": null,
        "study_design": null,
        "population": null,
        "intervention_or_exposure": null,
        "outcomes": []
      },
      "supported_claim": "Disabling automatic notifications can reduce notification interruptions; in this field experiment, fewer interruptions mediated better perceived performance and lower irritation.",
      "unsupported_or_overstated_claims": [
        "All notifications should be off all day for everyone.",
        "The study proves an optimal focus-block duration.",
        "The effect is independent of FoMO and response expectations."
      ],
      "limitations": [
        "The study was a one-day field experiment, so it does not establish durable effects of notification blocking over longer work periods or adaptation over time.",
        "Performance and irritation outcomes were self-reported, and the direct intervention effect on performance was not significant; the reported performance pathway was indirect through fewer interruptions.",
        "Work roles differ in response-time obligations, telepressure, fear of missing out and safety or escalation requirements, so an always-off notification rule is not supported.",
        "The study tests communication-application notification interruptions, not every form of digital distraction or every focus-window design."
      ],
      "target_hack_ids": [
        "batch-non-urgent-notifications"
      ],
      "target_protocol_ids": [
        "brali:batch-non-urgent-notifications"
      ],
      "risk_flags": [],
      "notes": "Strengthen the existing notification-batching surface rather than create a redundant always-off hack. Protect non-urgent focus windows while preserving explicit urgent-channel and role exceptions."
    },
    {
      "schema_version": 1,
      "id": "pmr-subjective-sleep-quality-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/pmr-subjective-sleep-quality-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/pmr-subjective-sleep-quality-boundary-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The effects of progressive muscle relaxation on sleep quality in adults: a systematic review and meta-analysis of randomized controlled trials",
        "url": "https://www.frontiersin.org/journals/public-health/articles/10.3389/fpubh.2026.1906525/full",
        "type": "systematic-review-meta-analysis",
        "doi": "10.3389/fpubh.2026.1906525",
        "citation_text": "Li K, Chen S, Liu Y, Yu N. Front Public Health. 2026;14:1906525.",
        "study_design": "PROSPERO-registered systematic review and random-effects meta-analysis of randomized controlled trials. Five databases were searched through May 2026; 14 RCTs were included. The review used RoB 2, Hartung-Knapp-adjusted random-effects models, subgroup analyses, leave-one-out sensitivity analysis, meta-regression for session count and publication-bias analyses.",
        "population": "957 adults across 14 randomized trials from China, Turkey, Iran and Egypt. Participants had cancer or other clinical conditions including coronary artery bypass graft, lung resection, restless legs syndrome, hemodialysis, rheumatoid arthritis, fractures, epilepsy, COPD and obstructive sleep apnea. The review did not include a clearly healthy general-population trial.",
        "intervention_or_exposure": "Progressive muscle relaxation delivered with heterogeneous protocols. Intervention duration ranged from 3 days to 12 weeks, and the number, frequency and length of sessions varied substantially across studies. Controls included routine care, no intervention, light activity and standard pulmonary rehabilitation.",
        "outcomes": [
          "Subjective sleep quality measured with the Pittsburgh Sleep Quality Index (PSQI)",
          "Subgroup estimates by cancer versus non-cancer clinical populations",
          "Subgroup estimates by mean age below versus at least 55 years",
          "Association between number of PMR sessions and sleep-quality effect size"
        ]
      },
      "supported_claim": "Across the included randomized trials, PMR improved subjective PSQI sleep-quality scores on average in clinically heterogeneous adult populations. The direction of the pooled effect remained after sensitivity and trim-and-fill analyses. This supports adding sleep as a bounded use case to Brali's existing conservative PMR protocol. It does not establish an optimal timer, session count or frequency, and the evidence should be described as subjective sleep-quality evidence in the studied clinical populations rather than a universal sleep effect.",
      "unsupported_or_overstated_claims": [
        "PMR reliably improves sleep for every adult or for healthy adults in general.",
        "The meta-analysis proves that PMR shortens sleep latency, increases total sleep time or improves objective sleep architecture.",
        "A specific 5/30-second timer, 7-minute routine, 12-group sequence, 18-group sequence or weekly frequency is the evidence-based optimum.",
        "More PMR sessions produce larger sleep benefits.",
        "PMR improves sleep because it lowers cortisol, reduces sympathetic activity or activates parasympathetic activity; these mechanisms were discussed, not tested by this meta-analysis.",
        "PMR does not work in adults aged 55 years or older.",
        "The pooled effect size can be treated as an expected individual improvement.",
        "PMR should replace assessment or treatment for persistent or severe sleep problems."
      ],
      "limitations": [
        "Between-study heterogeneity was very high (I2 85.5%) and was not explained by the reported subgroup analyses.",
        "The included trials involved diverse clinical populations, limiting direct generalization to healthy adults or any one condition.",
        "Sleep quality was assessed with the subjective PSQI rather than objective sleep measures.",
        "Intervention duration, frequency, session count and protocol details varied substantially across studies.",
        "Egger's test suggested possible publication bias; trim-and-fill reduced the pooled estimate while retaining the direction.",
        "Only 14 studies were available, limiting subgroup and meta-regression precision.",
        "The age >=55 subgroup estimate was imprecise, and the formal between-subgroup difference was not significant.",
        "The review did not establish long-term persistence of benefit after PMR stopped."
      ],
      "target_hack_ids": [
        "guided-pmr-muscle-relaxation"
      ],
      "target_protocol_ids": [
        "brali:guided-pmr-muscle-relaxation"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Editorial action: strengthen the existing reviewed PMR page rather than create a duplicate sleep-relaxation hack. Preserve the NHS-derived technique and safety details, add a sleep-specific observation loop, and teach the evidence boundary: subjective sleep quality improved on average in heterogeneous clinical populations, while exact dose, healthy-population generalization, objective sleep effects and proposed autonomic/cortisol mechanisms remain unproven here."
    },
    {
      "schema_version": 1,
      "id": "post-meal-exercise-acute-glucose-boundary-2023",
      "canonical_url": "https://brali-lifeos.github.io/evidence/post-meal-exercise-acute-glucose-boundary-2023/",
      "json_url": "https://brali-lifeos.github.io/evidence/post-meal-exercise-acute-glucose-boundary-2023/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "After Dinner Rest a While, After Supper Walk a Mile? A Systematic Review with Meta-analysis on the Acute Postprandial Glycemic Response to Exercise Before and After Meal Ingestion in Healthy Subjects and Patients with Impaired Glucose Tolerance",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC10036272/",
        "type": "systematic-review-meta-analysis",
        "doi": "10.1007/s40279-022-01808-7",
        "citation_text": "Engeroff T, Groneberg DA, Wilke J. Sports Med. 2023;53(4):849-869.",
        "study_design": "Prospectively registered systematic review and random-effects meta-analysis restricted to three-arm randomized controlled designs comparing matched exercise before meals, exercise after meals, and no exercise. Eight crossover trials contributed 30 interventions.",
        "population": "116 participants across eight crossover trials: 47 with type 2 diabetes and 69 without type 2 diabetes.",
        "intervention_or_exposure": "Acute exercise, often walking, performed before or after a meal and compared with an inactive control. Protocols varied in exercise type, intensity, duration and delay after eating.",
        "outcomes": [
          "Acute postprandial blood or interstitial glucose excursions",
          "Timing of exercise relative to the meal as a moderator"
        ]
      },
      "supported_claim": "Replacing some post-meal sitting with safe walking or other suitable movement relatively soon after eating is a defensible practical option when the goal is ordinary movement and the acute post-meal glucose mechanism is relevant. The meta-analysis found lower acute postprandial glucose after post-meal exercise than after no exercise and better average results than matched pre-meal exercise. Brali should not prescribe one duration, pace, start minute, or medical target.",
      "unsupported_or_overstated_claims": [
        "Everyone should walk exactly 20 minutes after every meal.",
        "There is a proven exact number of minutes after eating when exercise must start.",
        "Post-meal walking treats diabetes or permits medication changes.",
        "Acute glucose reductions prove long-term prevention of diabetes, cardiovascular disease, or other conditions.",
        "Exercise before meals is useless for health in general.",
        "A casual personal glucose reading can validate the protocol's causal effect for an individual."
      ],
      "limitations": [
        "Only eight randomized crossover trials with 116 total participants met the strict inclusion criteria.",
        "Included trials were rated high risk of bias.",
        "Exercise protocols, meals, participant characteristics and glucose measurement methods varied.",
        "The type 2 diabetes subgroup was small and some subgroup comparisons were imprecise.",
        "The review concerned acute postprandial responses and cannot establish long-term clinical outcomes.",
        "Moderator evidence about delay after eating does not define an exact individual optimum."
      ],
      "target_hack_ids": [
        "post-meal-walk"
      ],
      "target_protocol_ids": [
        "brali:post-meal-walk"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Publish as a safe movement substitution, not glucose self-treatment. Include medication/exercise-plan boundaries and judge the routine primarily by whether it replaces sitting and remains practical."
    },
    {
      "schema_version": 1,
      "id": "progress-monitoring-goal-attainment-meta-2016",
      "canonical_url": "https://brali-lifeos.github.io/evidence/progress-monitoring-goal-attainment-meta-2016/",
      "json_url": "https://brali-lifeos.github.io/evidence/progress-monitoring-goal-attainment-meta-2016/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Does monitoring goal progress promote goal attainment?",
        "url": "https://doi.org/10.1037/bul0000025",
        "type": "meta-analysis",
        "doi": "10.1037/bul0000025",
        "citation_text": null,
        "study_design": null,
        "population": null,
        "intervention_or_exposure": null,
        "outcomes": []
      },
      "supported_claim": "Increasing progress monitoring improves goal attainment on average; physically recording or reporting progress was associated with larger effects.",
      "unsupported_or_overstated_claims": [
        "Daily tracking is uniquely optimal.",
        "Streaks or public leaderboards are universally beneficial.",
        "More monitoring is always better."
      ],
      "limitations": [
        "The meta-analysis combines heterogeneous goals, populations and monitoring interventions, so the pooled effect does not identify one optimal metric, cadence or tracking interface.",
        "Larger effects associated with physically recording or reporting progress do not establish that public leaderboards, streaks, social accountability or any specific app feature is universally beneficial.",
        "The evidence supports increasing useful progress monitoring on average, not maximal monitoring; excessive or poorly chosen measurement can still create burden or metric gaming.",
        "The synthesis does not validate Brali's exact checkpoint wording or determine which proxy measure is best for a particular user's goal."
      ],
      "target_hack_ids": [
        "visible-progress-monitoring"
      ],
      "target_protocol_ids": [
        "brali:visible-progress-monitoring"
      ],
      "risk_flags": [],
      "notes": "Evidence backbone for existing tracker/progress surfaces. Prefer one measurable progress variable at a task-appropriate checkpoint over indiscriminate tracking."
    },
    {
      "schema_version": 1,
      "id": "ready-to-resume-interruption-boundary-2018",
      "canonical_url": "https://brali-lifeos.github.io/evidence/ready-to-resume-interruption-boundary-2018/",
      "json_url": "https://brali-lifeos.github.io/evidence/ready-to-resume-interruption-boundary-2018/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Tasks Interrupted: Ready-to-Resume Plan",
        "url": "https://doi.org/10.1287/orsc.2017.1184",
        "type": "primary-study",
        "doi": "10.1287/orsc.2017.1184",
        "citation_text": null,
        "study_design": null,
        "population": null,
        "intervention_or_exposure": null,
        "outcomes": []
      },
      "supported_claim": "When unfinished work must be interrupted, a short resumption plan can reduce attention residue and protect performance in the studied interruption contexts.",
      "unsupported_or_overstated_claims": [
        "It eliminates all switching costs.",
        "There is a validated exact three-line template.",
        "It has a known universal productivity percentage."
      ],
      "limitations": [
        "The evidence comes from a small set of controlled interruption studies and does not represent every form of complex, collaborative or high-stakes real-world work.",
        "A ready-to-resume plan can mitigate attention residue in the studied contexts; it does not make task switching cost-free or imply that avoidable interruptions should be accepted.",
        "The research supports making a concrete resumption plan, not Brali's exact three-prompt note format, note length, writing medium or timing rule.",
        "The studies do not establish a universal productivity percentage, a guaranteed performance benefit, or the same effect for every individual and task."
      ],
      "target_hack_ids": [
        "ready-to-resume-plan"
      ],
      "target_protocol_ids": [
        "brali:ready-to-resume-plan"
      ],
      "risk_flags": [],
      "notes": "Highest-value genuinely new productivity protocol from this pass. Proposed action: before a forced or deliberate switch from unfinished work, capture where you stopped, the next concrete action, and one critical restart detail."
    },
    {
      "schema_version": 1,
      "id": "savoring-emotional-outcomes-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/savoring-emotional-outcomes-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/savoring-emotional-outcomes-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Effectiveness of savoring interventions: A systematic review and meta-analysis of randomized controlled trials",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC12968602/",
        "type": "systematic-review-meta-analysis",
        "doi": "10.1111/aphw.70134",
        "citation_text": "Chen PH, Tung HH, Wu YC, Sie J. Appl Psychol Health Well Being. 2026;18(2):e70134.",
        "study_design": "Systematic review and three-level random-effects meta-analysis of 20 independent randomized controlled trials, accounting for 45 dependent effect sizes. Risk of bias was assessed with Cochrane RoB 2 and multiple publication-bias analyses were reported.",
        "population": "4,805 adult participants across 20 trials in Western and Eastern cultural contexts. Samples and intervention formats varied; included participants were adults only.",
        "intervention_or_exposure": "Savoring interventions delivered as individual exercises, online programs, or smartphone-based interventions. Intervention content and duration varied; the median intervention duration was 14 days.",
        "outcomes": [
          "Positive psychological states",
          "Negative emotional states",
          "Negative emotional symptoms",
          "Overall emotional outcomes"
        ]
      },
      "supported_claim": "Savoring interventions improve emotional outcomes on average across randomized trials, but the evidence does not identify one optimal micro-practice. Brali can reasonably offer a minimal nonclinical implementation—deliberately attending to an already-positive experience—as a personal experiment while clearly labeling the concrete format as an implementation choice rather than the meta-analysis's tested universal recipe.",
      "unsupported_or_overstated_claims": [
        "A brief pause in one positive moment has itself been proven to produce the pooled meta-analytic effect.",
        "Savoring reliably treats depression, anxiety, or another clinical condition.",
        "There is an evidence-based optimal duration or daily frequency for savoring.",
        "Effects are guaranteed to persist long term.",
        "Savoring requires gratitude, positive reframing, denial of negative emotion, or forced optimism."
      ],
      "limitations": [
        "Total heterogeneity was high, with I2 86.61% across effects.",
        "Six of 20 studies were rated high risk of bias and ten had some concerns.",
        "Most outcomes were self-reported and therefore vulnerable to recall and social-desirability bias.",
        "Intervention content, delivery format, populations and duration varied substantially.",
        "The median intervention duration was about 14 days, limiting long-term conclusions.",
        "All included samples were adults, limiting generalization to children and adolescents.",
        "Active-control point estimates were smaller than passive-control estimates, although control type did not reach statistical significance as a moderator."
      ],
      "target_hack_ids": [
        "savor-positive-moment"
      ],
      "target_protocol_ids": [
        "brali:savor-positive-moment"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Publish as a small attention practice around genuine positive experiences. Explicitly reject forced positivity and clinical treatment framing, and state that the exact Brali micro-format is a conservative implementation of a heterogeneous evidence base."
    },
    {
      "schema_version": 1,
      "id": "temptation-bundling-exercise-boundary-2020",
      "canonical_url": "https://brali-lifeos.github.io/evidence/temptation-bundling-exercise-boundary-2020/",
      "json_url": "https://brali-lifeos.github.io/evidence/temptation-bundling-exercise-boundary-2020/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Teaching temptation bundling to boost exercise: A field experiment",
        "url": "https://doi.org/10.1016/j.obhdp.2020.09.003",
        "type": "primary-study",
        "doi": "10.1016/j.obhdp.2020.09.003",
        "citation_text": "Milkman KL, Duckworth AL, Choi J, Laibson D, Madrian B. Organ Behav Hum Decis Process. 2020.",
        "study_design": "Large real-world gym field program containing a preregistered randomized comparison of a free audiobook plus temptation-bundling instruction versus a free audiobook alone, plus a broader comparison with a no-audiobook control. Objective gym check-in data were used before, during and after the intervention.",
        "population": "Adult members of 24 Hour Fitness who enrolled in a four-week exercise-boosting program. The preregistered eligible comparison included 2,334 participants; the broader three-condition analysis included 6,792 participants.",
        "intervention_or_exposure": "A free audiobook paired with explicit instruction and reminders to reserve desired media for gym exercise, compared with audiobook-only and no-audiobook program conditions. All StepUp participants also encountered planning prompts, reminders and small visit-linked rewards.",
        "outcomes": [
          "Probability of at least one weekly gym workout",
          "Average weekly gym visits during the program",
          "Gym attendance after the intervention"
        ]
      },
      "supported_claim": "Pairing a delayed-benefit behavior with a compatible immediate reward can modestly increase exercise participation in some field settings. Brali can offer temptation bundling as a task-initiation experiment, while keeping its strongest empirical anchor in exercise and requiring that the reward not impair the useful activity.",
      "unsupported_or_overstated_claims": [
        "Temptation bundling works equally well for every habit.",
        "The audiobook itself proves the mechanism independent of planning, reminders and other program components.",
        "A reward should always be restricted exclusively to the target behavior.",
        "The technique reliably creates permanent habits.",
        "Adding entertainment is appropriate for driving, complex learning, technical work or exercise where divided attention creates risk."
      ],
      "limitations": [
        "The strongest evidence is concentrated in exercise/gym behavior rather than arbitrary habits.",
        "The large StepUp program included multiple behavior-change components, complicating attribution in broader control comparisons.",
        "The incremental effect of explicit temptation-bundling teaching over receiving the audiobook alone was modest.",
        "Participants self-selected into an exercise-boosting program and may have been more motivated than typical gym members.",
        "The original field experiment showed that effects can decay and be disrupted by context changes such as holidays."
      ],
      "target_hack_ids": [
        "temptation-bundling"
      ],
      "target_protocol_ids": [
        "brali:temptation-bundling"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Publish with a compatibility/safety gate and behavior tracking. Keep generalization from exercise to other low-complexity habits explicitly tentative."
    },
    {
      "schema_version": 1,
      "id": "time-management-specificity-boundary-2021",
      "canonical_url": "https://brali-lifeos.github.io/evidence/time-management-specificity-boundary-2021/",
      "json_url": "https://brali-lifeos.github.io/evidence/time-management-specificity-boundary-2021/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-09-11",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Does time management work? A meta-analysis",
        "url": "https://doi.org/10.1371/journal.pone.0245066",
        "type": "meta-analysis",
        "doi": "10.1371/journal.pone.0245066",
        "citation_text": null,
        "study_design": null,
        "population": null,
        "intervention_or_exposure": null,
        "outcomes": []
      },
      "supported_claim": "Time management has moderate positive relationships with performance and wellbeing at the construct level, with substantial heterogeneity.",
      "unsupported_or_overstated_claims": [
        "This validates a specific planner, calendar ratio, time-block length or branded method.",
        "All observed relationships are causal."
      ],
      "limitations": [
        "The meta-analysis pools studies using different definitions and measures of time management, job or academic performance and wellbeing, which limits claims about any one technique.",
        "A substantial part of the underlying evidence is correlational, so associations between time-management behavior and outcomes should not be read as causal effects of a specific intervention.",
        "The synthesis does not validate an exact planner layout, calendar ratio, time-block duration, prioritization rule or scheduling cadence.",
        "Observed relationships vary across contexts and outcomes, so broad time-management evidence should not be converted into precise universal productivity claims."
      ],
      "target_hack_ids": [],
      "target_protocol_ids": [],
      "risk_flags": [],
      "notes": "Wording guardrail: broad time-management evidence cannot be used to invent scientifically optimal implementation details."
    },
    {
      "schema_version": 1,
      "id": "exercise-snacks-cardiorespiratory-fitness-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/exercise-snacks-cardiorespiratory-fitness-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/exercise-snacks-cardiorespiratory-fitness-boundary-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-09",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Effects of exercise snacks on cardiorespiratory fitness, body composition, and blood lipids in adults across different age groups: a systematic review and meta-analysis",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC13486332/",
        "type": "systematic-review-meta-analysis",
        "doi": "10.1016/j.jnha.2026.100940",
        "citation_text": "Meng L, Li Y, Chen J, Zheng H. J Nutr Health Aging. 2026;30(10):100940.",
        "study_design": "Prospectively registered PRISMA systematic review and random-effects meta-analysis of randomized controlled trials. Six databases were searched from inception through April 15, 2026; 21 studies were included. Risk of bias, subgroup analyses, meta-regression, sensitivity analyses, publication bias and GRADE certainty were reported.",
        "population": "921 adult participants across 21 randomized studies conducted in multiple countries. Study mean ages ranged from 18.9 to 75.7 years, with substantial variation in body mass index, health status and training context.",
        "intervention_or_exposure": "Exercise snacks or similar fragmented low-volume exercise accumulated through multiple brief bouts. Included modalities mainly involved stair climbing, resistance exercise, cycling and sprint-interval exercise, with heterogeneous intensity, bout duration, daily distribution, weekly frequency, intervention duration and supervision.",
        "outcomes": [
          "Maximal oxygen uptake (VO2max)",
          "Peak power output",
          "Body fat percentage",
          "Total cholesterol",
          "HDL cholesterol",
          "LDL cholesterol",
          "Triglycerides"
        ]
      },
      "supported_claim": "Across the included randomized studies, exercise snacks improved VO2max and peak power output on average with moderate-certainty evidence. This supports offering distributed brief exercise bouts as a practical cardiorespiratory-fitness option when a longer continuous workout is difficult to fit. The evidence does not establish one modality, timer, number of bouts, weekly frequency or intervention duration as the optimal prescription.",
      "unsupported_or_overstated_claims": [
        "A 20-seconds-on/10-seconds-off Tabata timer is the evidence-based exercise-snack dose.",
        "Three or four one-minute exercise snacks per day are universally optimal.",
        "The subgroup with more than five daily bouts proves that more bouts are better for each individual.",
        "Four to five days per week or seven to nine weeks is an optimized prescription.",
        "Exercise snacks reliably reduce body fat or improve cholesterol and triglycerides.",
        "Exercise snacks are equivalent to or should replace conventional exercise.",
        "The nonsignificant VO2max result in studies with mean age above 50 proves exercise snacks do not work for people older than 50.",
        "A short exercise bout automatically produces a meaningful training stimulus regardless of intensity, modality or execution."
      ],
      "limitations": [
        "Cardiorespiratory-fitness effects showed substantial between-study heterogeneity (I2 74% for VO2max and 63% for peak power output), which contributed to GRADE downgrading.",
        "Certainty was low for body-fat and blood-lipid outcomes because of heterogeneity, imprecision and possible publication bias for some outcomes.",
        "Interventions differed markedly in modality, intensity, bout duration, within-day distribution, weekly frequency, intervention duration and supervision.",
        "Dietary control varied across trials and could contribute to inconsistent body-composition and lipid results.",
        "Age analyses used study-level mean age rather than individual participant data, creating ecological-bias risk and making individual age cutoffs inappropriate.",
        "Potential publication bias was detected for VO2max and body-fat percentage.",
        "Subgroup analyses compared relatively small numbers of studies and cannot identify a causal optimal dose.",
        "Long-term scalability remains uncertain."
      ],
      "target_hack_ids": [
        "4-minute-hiit-tabata-workout"
      ],
      "target_protocol_ids": [
        "brali:4-minute-hiit-tabata-workout"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Editorial action: update the existing reviewed HIIT/Tabata protocol rather than add a duplicate exercise-snack hack. Explain two distinct implementations—one compact interval session versus brief bouts distributed across the day—and position exercise snacks primarily as a cardiorespiratory-fitness access strategy. Keep timer choice flexible and make the evidence boundary explicit for body composition, lipids, age subgroups and supposed optimal cadence."
    },
    {
      "schema_version": 1,
      "id": "exercise-snacks-real-world-t2d-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/exercise-snacks-real-world-t2d-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/exercise-snacks-real-world-t2d-boundary-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-09",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Exercise Snacks Are Feasible to Perform in the Real World and Improve Physical Capacity for Adults Living With Non-Insulin Treated Type 2 Diabetes: A Randomised Trial",
        "url": "https://pubmed.ncbi.nlm.nih.gov/42502209/",
        "type": "primary-study",
        "doi": "10.1111/dom.71052",
        "citation_text": "Babir FJ, Marcotte-Chénard A, Sandilands RE, et al. Diabetes Obes Metab. 2026;28(10):9294-9303.",
        "study_design": "12-week randomized real-world trial with an active comparator. Participants were assigned to four 1-minute vigorous exercise-snack bouts or four 1-minute low-intensity mobility/stretching bouts on at least five days per week. Feasibility based on adherence was the primary outcome; multiple fitness, glycemic, biomarker and anthropometric measures were secondary outcomes.",
        "population": "69 insufficiently active adults living with non-insulin-treated type 2 diabetes; 46 were female and mean age was 58 ± 11 years.",
        "intervention_or_exposure": "Four one-minute vigorous exercise bouts per day on at least five days per week for 12 weeks, delivered remotely and compared with the same frequency and duration of low-intensity mobility/stretching.",
        "outcomes": [
          "Weekly adherence and exercise enjoyment",
          "Perceived exertion and peak heart rate",
          "HbA1c",
          "Cardiometabolic blood biomarkers",
          "30-second sit-to-stand capacity",
          "Grip strength",
          "Estimated maximal oxygen uptake",
          "Anthropometrics"
        ]
      },
      "supported_claim": "In this specific real-world type 2 diabetes sample, the four-by-one-minute vigorous exercise-snack program was feasible and improved 30-second sit-to-stand performance more than an active mobility/stretching comparator. The trial does not show a broad cardiometabolic or aerobic-fitness advantage for this exact prescription.",
      "unsupported_or_overstated_claims": [
        "Four one-minute vigorous bouts on five days per week are an optimal exercise-snack prescription.",
        "High adherence to exercise snacks guarantees improvement in VO2max, HbA1c or cardiometabolic biomarkers.",
        "The trial proves exercise snacks improve diabetes control.",
        "The findings can be generalized unchanged to healthy adults, younger adults, insulin-treated diabetes or other clinical populations.",
        "A null result for most secondary outcomes proves exercise snacks are ineffective in general.",
        "The observed sit-to-stand benefit proves reduced fall risk or improved long-term clinical outcomes."
      ],
      "limitations": [
        "The sample was limited to insufficiently active adults with non-insulin-treated type 2 diabetes and well-controlled glycemia at baseline.",
        "The primary outcome was feasibility; efficacy outcomes were secondary and the study was not a definitive comparative-effectiveness trial for all health outcomes.",
        "The mobility/stretching control was active rather than no-exercise, so the contrast estimates the added effect of the vigorous condition over a matched low-intensity routine.",
        "The intervention was remotely delivered and real-world exercise intensity was lower than some prior supervised laboratory exercise-snack protocols.",
        "Estimated VO2max relied on a submaximal predictive test that the investigators reported performed poorly in this population, limiting interpretation of that fitness outcome.",
        "The trial lasted 12 weeks and does not establish longer-term adherence or clinical outcomes."
      ],
      "target_hack_ids": [
        "4-minute-hiit-tabata-workout"
      ],
      "target_protocol_ids": [
        "brali:4-minute-hiit-tabata-workout"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Use as a calibration source, not as the main efficacy anchor. It demonstrates that a plausible exercise-snack schedule can be feasible yet fail to produce broad secondary-outcome superiority in free-living conditions. Teach readers to notice whether the work is genuinely challenging and whether function or fitness changes over time instead of assuming timer completion equals physiological adaptation."
    },
    {
      "schema_version": 1,
      "id": "woop-mcii-academic-procrastination-rct-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/woop-mcii-academic-procrastination-rct-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/woop-mcii-academic-procrastination-rct-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-09",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Beyond positive thinking: A randomized trial of mental contrasting with implementation intentions to curb academic procrastination",
        "url": "https://pubmed.ncbi.nlm.nih.gov/41601124/",
        "type": "randomized-controlled-trial",
        "doi": "10.1016/j.actpsy.2025.106168",
        "citation_text": "Zhou X, Wider W, Wu H, Xu Y, Qin M, Borromeo AS. Acta Psychologica. 2026;262:106168.",
        "study_design": "Randomized controlled comparison of MCII with a positive-thinking strategy control. The study randomized 101 eligible students; 20 were excluded during the seven-day reporting period for missing reports or intervention nonadherence, leaving 81 in the final sample. Outcomes were measured through pre-test, post-test, seven-day diaries and a one-week follow-up.",
        "population": "101 eligible first-year undergraduate students at one private Chinese university were randomized; 81 remained in the final analyzed sample after 19.8% attrition.",
        "intervention_or_exposure": "MCII applied to academic procrastination compared with a positive-thinking strategy control.",
        "outcomes": [
          "Task aversiveness",
          "Outcome utility",
          "Willingness to initiate academic tasks",
          "Diary-based task initiation"
        ]
      },
      "supported_claim": "In this specific undergraduate academic-procrastination study, MCII reduced task aversiveness and improved willingness to initiate academic tasks relative to a positive-thinking control, with effects reported through one-week follow-up. Diary reports also provided a behavioral task-initiation signal. It provides narrow corroborating evidence for the obstacle-plus-plan mechanism in an academic setting.",
      "unsupported_or_overstated_claims": [
        "This trial proves WOOP works for all goals or populations.",
        "The trial establishes long-term goal achievement effects.",
        "The results generalize to workplaces, clinical populations or high-stakes decisions.",
        "WOOP eliminates procrastination.",
        "A one-week follow-up establishes durable behavior change."
      ],
      "limitations": [
        "The study was single-site and restricted to first-year undergraduates at one private Chinese university.",
        "Twenty of 101 randomized participants were excluded for missing reports or intervention nonadherence, leaving 81 in the final sample (19.8% attrition).",
        "The main psychological outcomes were self-reported; task initiation was derived from participant diary reports rather than an objective academic-performance endpoint.",
        "The context was academic procrastination rather than general goal pursuit.",
        "The follow-up was one week, so long-term persistence is unknown.",
        "The study does not justify generalizing to broad life outcomes, workplaces, clinical treatment, or safety-critical decisions."
      ],
      "target_hack_ids": [
        "woop-goal-planner"
      ],
      "target_protocol_ids": [
        "brali:woop-goal-planner"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Use as corroborating boundary evidence only. Reuse the existing research-candidate identity; do not turn this narrow RCT into the primary public claim or imply cross-domain efficacy."
    },
    {
      "schema_version": 1,
      "id": "woop-mcii-goal-attainment-meta-analysis-2021",
      "canonical_url": "https://brali-lifeos.github.io/evidence/woop-mcii-goal-attainment-meta-analysis-2021/",
      "json_url": "https://brali-lifeos.github.io/evidence/woop-mcii-goal-attainment-meta-analysis-2021/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-09",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "A Meta-Analysis of the Effects of Mental Contrasting With Implementation Intentions on Goal Attainment",
        "url": "https://www.frontiersin.org/journals/psychology/articles/10.3389/fpsyg.2021.565202/full",
        "type": "systematic-review-meta-analysis",
        "doi": "10.3389/fpsyg.2021.565202",
        "citation_text": "Wang G, Wang Y, Gai X. Frontiers in Psychology. 2021;12:565202.",
        "study_design": "Random-effects meta-analysis of 21 empirical articles contributing 24 independent MCII effect sizes across field interventions and goal domains, with 15,907 participants in total.",
        "population": "Children, college students and adults across academic, health, relationship and personal goal contexts represented in the included intervention studies. The corpus was heterogeneous and was dominated numerically by two very large MOOC samples.",
        "intervention_or_exposure": "Mental contrasting with implementation intentions: identify a desired future and an obstacle in reality, then form an if-obstacle-then-behavior plan. Delivery varied between experimenter-led and document-based interventions.",
        "outcomes": [
          "Goal attainment across included studies"
        ]
      },
      "supported_claim": "Across the included field-intervention studies, mental contrasting with implementation intentions was associated with a small-to-medium average improvement in goal attainment. The pooled random-effects estimate was Hedges' g=0.336, while publication-bias adjustment produced a smaller estimate. This supports WOOP/MCII as a bounded self-regulation option for feasible goals, not as a guarantee of individual goal success.",
      "unsupported_or_overstated_claims": [
        "WOOP guarantees that a person will achieve a goal.",
        "The pooled effect size applies to every person, goal domain or delivery format.",
        "WOOP is always better than mental contrasting or implementation intentions used separately.",
        "A document or self-guided WOOP exercise is as effective as an experimenter-led intervention.",
        "The meta-analysis establishes a clinically meaningful treatment effect.",
        "The study proves that one fixed WOOP script is optimal."
      ],
      "limitations": [
        "The included studies were heterogeneous in population, goal domain, intervention delivery and outcome measurement.",
        "The meta-analysis reported medium heterogeneity and mixed publication-bias diagnostics; trim-and-fill yielded a smaller adjusted estimate.",
        "Two very large MOOC studies accounted for most participants and were excluded from moderator analyses because of their numerical dominance.",
        "The number of studies was limited for moderator inference, and the authors explicitly called for additional studies.",
        "An average standardized effect cannot predict whether the protocol will help one individual with one goal."
      ],
      "target_hack_ids": [
        "woop-goal-planner"
      ],
      "target_protocol_ids": [
        "brali:woop-goal-planner"
      ],
      "risk_flags": [],
      "notes": "Editorial outcome: retain WOOP as reviewed practical self-regulation guidance for meaningful, reasonably feasible goals. Keep the user-visible wording focused on the obstacle-to-if-then planning process and an observable attempt, not a promise of goal attainment or an advertised effect size."
    },
    {
      "schema_version": 1,
      "id": "movement-breaks-cardiometabolic-vanherle-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/movement-breaks-cardiometabolic-vanherle-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/movement-breaks-cardiometabolic-vanherle-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-08",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Optimizing physical activity bouts to interrupt sedentary behaviour for cardiometabolic health: a systematic review and meta-analyses of randomized controlled trials",
        "url": "https://academic.oup.com/eurjpc/advance-article/doi/10.1093/eurjpc/zwag079/8571395",
        "type": "meta-analysis",
        "doi": "10.1093/eurjpc/zwag079",
        "citation_text": "Vanherle J, Franssen GHLM, Ivanova A, Eijnde BO, Franssen WMA. European Journal of Preventive Cardiology. Published online 1 April 2026. doi:10.1093/eurjpc/zwag079.",
        "study_design": "Systematic review and meta-analyses of randomized controlled trials comparing acute laboratory-based aerobic physical-activity interruptions with sedentary control conditions; risk of bias assessed with Cochrane RoB 2 and certainty with GRADE.",
        "population": "Adults aged 18-65 years, including healthy participants and people with cardiometabolic conditions, across 144 studies, 247 intervention arms, and 2,216 participants.",
        "intervention_or_exposure": "Physical-activity bouts interrupting prolonged sedentary behaviour, varying in frequency, duration, and intensity, compared with uninterrupted sedentary control; standing-only conditions were examined separately.",
        "outcomes": [
          "Acute glucose concentrations",
          "Acute insulin concentrations",
          "Blood pressure",
          "Triglycerides",
          "Endothelial function"
        ]
      },
      "supported_claim": "In acute controlled settings, interrupting prolonged sitting with actual physical activity can improve glucose and insulin responses on average. More frequent light-to-moderate activity appears useful for acute glucose control, while standing alone did not show the same cardiometabolic signal. For Brali this supports replacing an avoidable sitting block with safe movement, including walking when a meeting can travel, without prescribing one universal interval.",
      "unsupported_or_overstated_claims": [
        "A walking meeting has a special metabolic, productivity, creativity, or cognitive effect beyond the movement it contains.",
        "Everyone should stand or walk every 20 or 30 minutes.",
        "Standing still is metabolically equivalent to a walking or physical-activity break.",
        "Acute laboratory improvements prove long-term prevention or treatment of diabetes, hypertension, cardiovascular disease, or other conditions.",
        "A walking meeting should target a fixed step count, heart-rate increase, duration, or physiological dose.",
        "The review supports inherited claims about wind noise, phone battery drain, meeting quality, mental clarity, or screen fatigue."
      ],
      "limitations": [
        "The evidence is dominated by acute controlled experimental settings and cannot establish long-term clinical benefit or sustainable free-living adherence.",
        "Participant characteristics, meals, timing, activity protocols, and outcome measurement varied substantially across studies.",
        "Many individual studies had small samples; only a minority of intervention comparisons were rated low overall risk of bias, with most having some concerns.",
        "The review focused on adults aged 18-65 years, so the findings should not be generalized automatically to children or older adults.",
        "The evidence concerns physical-activity interruptions, not walking meetings as a distinct behavioral or workplace intervention."
      ],
      "target_hack_ids": [
        "walking-meeting-assistant"
      ],
      "target_protocol_ids": [
        "brali:walking-meeting-assistant"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Editorial action: replace the inherited walking-meeting page with a mechanism-first implementation protocol. The meeting is only a delivery surface for exchanging some sitting for movement. Remove unsupported numerical targets and health/productivity halo claims; teach users to observe whether sitting was actually replaced, the conversation still worked, and the route remained safe and practical. The Gale 2026 meta-analysis is already represented by Evidence Decision activity-breaks-postprandial-metabolism-boundary-2026 and is reused as converging evidence rather than duplicated."
    },
    {
      "schema_version": 1,
      "id": "caffeine-dose-timing-sleep-boundary-2025",
      "canonical_url": "https://brali-lifeos.github.io/evidence/caffeine-dose-timing-sleep-boundary-2025/",
      "json_url": "https://brali-lifeos.github.io/evidence/caffeine-dose-timing-sleep-boundary-2025/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-06",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Dose and timing effects of caffeine on subsequent sleep: a randomized clinical crossover trial",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC11985402/",
        "type": "primary-study",
        "doi": "10.1093/sleep/zsae230",
        "citation_text": "Gardiner CL, Weakley J, Burke LM, et al. Sleep. 2025;48(4):zsae230.",
        "study_design": "Placebo-controlled, double-blind randomized crossover trial with seven conditions separated by 48-hour washouts. Participants received placebo or a single 100 mg or 400 mg caffeine dose 12, 8, or 4 hours before habitual bedtime; sleep was assessed with in-home partial polysomnography and sleep diaries.",
        "population": "23 healthy males aged 18–40 years (mean age 25.3 years) with moderate habitual caffeine intake below 300 mg per day.",
        "intervention_or_exposure": "Acute caffeine anhydrous capsules containing 100 mg or 400 mg, administered once at 12, 8, or 4 hours before habitual bedtime, compared with placebo.",
        "outcomes": [
          "Objective total sleep time",
          "Sleep efficiency",
          "Sleep onset latency and latency to persistent sleep",
          "Wake after sleep onset and awakenings",
          "Sleep-stage distribution including N3 sleep",
          "Subjective total sleep time and sleep quality"
        ]
      },
      "supported_claim": "Caffeine timing should be interpreted together with dose rather than as one universal after-lunch rule. In this small randomized crossover trial, a single 400 mg dose disrupted multiple objective sleep measures when taken within 12 hours of bedtime, with larger disruption closer to bedtime; no statistically significant sleep effect was detected for 100 mg at the tested 12-, 8-, and 4-hour timings. For Brali, this supports moving large late-day caffeine doses earlier when sleep matters and using bedtime-relative timing rather than a fixed clock cutoff.",
      "unsupported_or_overstated_claims": [
        "Everyone should stop caffeine at 1 p.m. or immediately after lunch.",
        "A 100 mg dose four hours before bedtime is safe or sleep-neutral for every person.",
        "Twelve hours is a scientifically established cutoff for every caffeine dose or pattern of use.",
        "People can infer whether they metabolize caffeine quickly from whether they fall asleep after an afternoon coffee.",
        "The study establishes an ideal daily caffeine dose, a body-weight dosing range, or a morning dosing schedule.",
        "The trial tested ordinary coffee, tea, energy drinks, or repeated real-world servings rather than standardized caffeine capsules.",
        "A personal sleep log can establish a causal individual caffeine dose-response curve."
      ],
      "limitations": [
        "The sample was small (23 participants) and included only healthy men aged 18–40, limiting generalizability to women, older adults, adolescents, people with sleep disorders, and other populations.",
        "Participants were moderate habitual caffeine users; results may differ in people with very low, very high, or irregular caffeine exposure.",
        "The intervention used acute standardized caffeine capsules rather than the variable doses, ingredients, absorption patterns, and repeated servings found in real-world caffeinated products.",
        "The study tested only two caffeine doses and three pre-bed timing points, so it cannot define a continuous or personalized cutoff.",
        "Sleep was measured with in-home partial polysomnography rather than full laboratory polysomnography.",
        "A nonsignificant result for 100 mg in this sample is not proof of no effect for every individual, especially given known variability in caffeine pharmacokinetics and sensitivity."
      ],
      "target_hack_ids": [
        "stop-caffeine-after-lunch"
      ],
      "target_protocol_ids": [
        "brali:stop-caffeine-after-lunch"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Editorial action: replace the inherited fixed after-lunch/1 p.m. rule with a dose-aware, bedtime-relative protocol. Start with the highest-leverage low-risk change: move large late-day doses earlier or reduce them. Let users track last caffeine timing/rough amount and broad sleep/daytime-function trends without deliberately challenging themselves with late caffeine. Keep the 100 mg null finding as an evidence boundary, not a clearance rule."
    },
    {
      "schema_version": 1,
      "id": "genai-learning-augmentation-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/genai-learning-augmentation-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/genai-learning-augmentation-boundary-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-05",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Evidence of impact and interpretational limits of generative AI in STEM education: a systematic review and meta-analysis on cognitive learning outcomes",
        "url": "https://link.springer.com/article/10.1007/s10462-026-11665-9",
        "type": "systematic-review-and-meta-analysis",
        "doi": "10.1007/s10462-026-11665-9",
        "citation_text": "Boolzen C, Kuhn J, Flegr S, et al. Artificial Intelligence Review. Published August 25, 2026.",
        "study_design": "Preregistered systematic review and meta-analysis of peer-reviewed quantitative STEM-education studies involving learner interaction with generative AI and a comparison or control group. Searches of ERIC, PsycINFO and Web of Science were updated May 7, 2026 and supplemented by citation tracking. Eighty-five studies met review criteria; 49 studies contributed 59 effect sizes to the quantitative meta-analysis. Random-effects meta-analysis, risk-of-bias assessment, moderator analyses, coincidence analysis and Robust Bayesian Meta-Analysis were used to examine heterogeneity and publication bias.",
        "population": "Learners in STEM education. Most reviewed work used text-based generative-AI systems and much of the literature was in higher education; among the 49 meta-analyzed articles, the detailed review classified 25 as university studies, 23 as K-12 and one with educational level unavailable.",
        "intervention_or_exposure": "Learners used generative-AI systems during STEM learning activities. The review classified studies partly by whether AI substituted, augmented or redefined learners' cognitive activity relative to the comparison condition and by whether outcomes primarily assessed knowledge or skills.",
        "outcomes": [
          "Externally assessed cognitive learning outcomes",
          "Knowledge versus skill outcomes",
          "Heterogeneity across studies",
          "Publication-bias-adjusted overall effect",
          "Moderation by augmentation/substitution of cognitive activity",
          "Reported learner challenges and instructional supports"
        ]
      },
      "supported_claim": "Brali should not treat generative AI as a learning intervention by itself. In this highly heterogeneous literature, the apparent positive pooled effect did not survive robust publication-bias correction, while more informative moderator evidence suggested that outcomes depend partly on what cognitive work the learner still performs. This supports the existing human-first AI collaboration principle in learning contexts: preserve a meaningful learner attempt, explanation, retrieval, reasoning or verification step and use AI to extend or refine that activity rather than silently replacing the activity being learned.",
      "unsupported_or_overstated_claims": [
        "Generative AI has been shown to improve STEM learning overall.",
        "Human-first AI workflows are proven superior for every learning task; the review compared heterogeneous instructional designs rather than one canonical Brali workflow.",
        "AI substitution is always harmful or augmentation is always beneficial.",
        "The meta-analysis establishes an optimal prompt, model, tutoring script, feedback style or amount of AI assistance.",
        "Knowledge outcomes are universally improved while skill outcomes are not; substantial residual heterogeneity remained within categories.",
        "A large effect reported in an individual AI-learning study should be assumed causal when intervention and control groups performed different levels of cognitive activity.",
        "The findings establish long-term skill retention, transfer to real work, or independent performance after AI is removed.",
        "The results generalize unchanged beyond STEM education."
      ],
      "limitations": [
        "Between-study heterogeneity was extreme (I²=96.32%), and the prediction interval included negative, null and strongly positive true effects.",
        "Funnel-plot asymmetry and the Robust Bayesian Meta-Analysis indicated substantial publication bias; after correction the evidence favored no stable overall positive or negative main effect.",
        "The review's risk-of-bias assessment found high risk across included studies, with no study meeting the low-risk criteria.",
        "Many studies used cognitively incomparable intervention and control conditions; the authors identified numerous apparently large effects where this comparison problem was present.",
        "Most studies used text-based systems and many were conducted in higher education, limiting generalization to other learners, modalities and tasks.",
        "AI literacy, metacognitive skill, delegation behavior, prompt quality and verification behavior were often underreported, leaving important mechanisms unresolved.",
        "Evidence about learner challenges and instructional supports was sparse and inconsistently reported, so those qualitative findings should not be turned into general effect estimates.",
        "Moderator patterns explained only part of the heterogeneity and did not produce a sufficient if-then configuration that guaranteed large effects."
      ],
      "target_hack_ids": [
        "human-first-ai-collaboration"
      ],
      "target_protocol_ids": [
        "brali:human-first-ai-collaboration"
      ],
      "risk_flags": [],
      "notes": "Evidence role: strengthen and teach the existing boundary rather than create another AI hack. For a learning task, ask what cognitive operation the user is trying to acquire. Require a meaningful unaided attempt or explicit reasoning step where that operation matters; then use AI for examples, critique, comparison, feedback or extension; finally verify or reproduce the target skill without leaning on the generated answer. The sarcastic version is accurate enough for editorial use: if the AI performs the exact thinking you are trying to learn, excellent news for the AI. The learner still needs a turn."
    },
    {
      "schema_version": 1,
      "id": "protect-off-work-time-from-work-apps-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/protect-off-work-time-from-work-apps-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/protect-off-work-time-from-work-apps-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-05",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Reducing Work-Related Screen-Time in Healthcare Workers During Leisure Time (REDUCE SCREEN) – A Randomized Controlled Trial",
        "url": "https://link.springer.com/article/10.1007/s10916-026-02338-9",
        "type": "primary-study",
        "doi": "10.1007/s10916-026-02338-9",
        "citation_text": "Bartels K, Shah K, Sanchez Rodriguez E, et al. Journal of Medical Systems. 2026;50:11.",
        "study_design": "Preregistered pragmatic parallel randomized controlled trial conducted from November 2021 to November 2023. Eight hundred fifteen healthcare workers were randomized 1:1 before a selected work-free weekend to usual behavior or a three-component educational intervention. The trial was registered before enrollment (NCT05106647). Analysts were masked to allocation until the primary analysis; participants could not be masked. Because 295 randomized participants did not provide post-weekend outcomes, the main analysis was a modified intention-to-treat analysis among 520 responders.",
        "population": "815 adult US healthcare workers who routinely used a smartphone and had a work email application installed. The post-weekend analysis included 520 respondents; the modal age group was 25-34 years, 57% were women, and participants included nurses, physicians, advanced practice providers, trainees and other healthcare professionals.",
        "intervention_or_exposure": "Before an off-work weekend, participants received education encouraging three strategies: activate an automatic reply to incoming work email, reduce leisure-time screen use, and uninstall work applications from personal devices. Participants could choose to implement all, some, or none of the components. Control participants received no intervention.",
        "outcomes": [
          "Perceived Stress Scale-10 after the weekend adjusted for baseline",
          "Change in self-reported device screen time based on operating-system screen-time summaries",
          "Exploratory associations between chosen components and stress"
        ]
      },
      "supported_claim": "In this healthcare-worker sample, assigning a practical package designed to reduce work-related smartphone intrusion during a genuinely off-duty weekend produced a modest additional reduction in perceived stress versus an untreated off weekend and reduced reported screen time. Brali can therefore test a bounded off-work digital-boundary protocol: when a period is truly off duty, make non-availability explicit, remove or disable the easiest work-app entry points that are not required, and restore them when the off-duty period ends.",
      "unsupported_or_overstated_claims": [
        "Uninstalling work applications by itself has been proven to reduce stress; component use was self-selected within the intervention group.",
        "Automatic out-of-office replies by themselves reduce stress.",
        "All screen time should be reduced during leisure, regardless of content or purpose.",
        "The intervention prevents burnout, depression, anxiety, medical errors, or workforce attrition.",
        "The effect generalizes unchanged to non-healthcare workers, shift workers, caregivers, or people who are on call.",
        "One weekend is an optimal dose or proves durable long-term benefit.",
        "People should disable channels needed for emergencies, on-call duties, authentication, safety, or required care responsibilities.",
        "A one-hour reduction in screen time is a universal target or causal threshold."
      ],
      "limitations": [
        "About 36% of randomized participants did not provide the post-weekend outcome, so attrition could bias the modified intention-to-treat result despite similar measured baseline characteristics among completers and non-completers.",
        "The randomized intervention was a three-component educational package, so the causal contribution of automatic replies, general screen reduction, or work-app removal cannot be isolated.",
        "Which components participants actually adopted was self-selected; the apparently larger stress reduction among those who removed work apps is exploratory and vulnerable to motivation and selection bias.",
        "The primary outcome was self-reported perceived stress measured over a short off-work weekend; burnout and long-term recovery were not established.",
        "Screen time came from participant-reported operating-system summaries and did not distinguish work-related from leisure-related use.",
        "The study was conducted in US healthcare workers, where work urgency and after-hours obligations may differ materially from other occupations.",
        "The protocol must not be applied to periods when the user is formally on call or must remain reachable for safety-critical responsibilities."
      ],
      "target_hack_ids": [
        "protect-off-work-time-from-work-apps"
      ],
      "target_protocol_ids": [
        "brali:protect-off-work-time-from-work-apps"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Protocol-design direction: for a user-chosen period that is genuinely off duty, first communicate the boundary (for example, an approved away status or handoff), then disable/remove optional work-app access from the personal device or sign out of it, while preserving any explicitly required emergency/on-call channel. Restore access after the protected period. Measure whether work checking actually falls and whether subjective recovery/stress improves. Treat this as a boundary experiment, not a treatment for burnout or a blanket anti-screen rule."
    },
    {
      "schema_version": 1,
      "id": "sitting-breaks-cognition-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/sitting-breaks-cognition-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/sitting-breaks-cognition-boundary-2026/index.json",
      "decision": "watch",
      "decision_label": "Watch, do not prescribe",
      "reviewed_at": "2026-09-05",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Acute cognitive effects of interruptions to prolonged sitting with brief standing or physical activity breaks: a systematic review and three-level meta-analysis",
        "url": "https://link.springer.com/article/10.1186/s12966-026-01953-6",
        "type": "systematic-review-meta-analysis",
        "doi": "10.1186/s12966-026-01953-6",
        "citation_text": "Zhuang M, Yin M, Liu Z, et al. International Journal of Behavioral Nutrition and Physical Activity. 2026;23:79.",
        "study_design": "Prospectively registered systematic review (PROSPERO CRD420251273873) following PRISMA. It searched five major databases through January 2026 for randomized crossover trials comparing prolonged sitting with brief standing or physical-activity interruptions and cognitive outcomes. Twenty-one trials involving 433 participants were included. A three-level random-effects meta-analysis and cluster-robust variance estimation addressed dependent effect sizes; risk of bias used RoB 2 and certainty used GRADE.",
        "population": "433 participants across 21 randomized crossover trials. Samples and protocols varied; exploratory subgroup analyses suggested differences by age and weight status, but the review explicitly judged these insufficient for prescriptive recommendations.",
        "intervention_or_exposure": "Brief standing or physical-activity interruptions during experimentally prolonged sitting, compared with uninterrupted sitting. Interruption frequency, duration, intensity and modality varied across studies.",
        "outcomes": [
          "Executive function",
          "Memory function",
          "Global cognition",
          "Information processing",
          "Attention",
          "GRADE certainty and exploratory protocol moderators"
        ]
      },
      "supported_claim": "The existing Brali sitting-break protocol has limited additional cognitive support: across randomized crossover studies, interrupting prolonged sitting was associated with small acute improvements in executive function and memory. Because certainty was low and attention/information-processing results were not clear, cognition should remain a secondary possible benefit rather than the reason or promise for the protocol.",
      "unsupported_or_overstated_claims": [
        "Movement breaks reliably improve attention or concentration; the pooled attention estimate was not beneficial.",
        "Movement breaks reliably improve overall cognition or information-processing speed.",
        "A specific break frequency is proven optimal for executive function or memory.",
        "Low-intensity walking is definitively superior to other feasible movement modes for cognition.",
        "The small acute cognitive effects translate into better workplace productivity, fewer errors, academic achievement, or long-term cognitive protection.",
        "The exploratory subgroup results establish who should use a more frequent or less frequent break schedule."
      ],
      "limitations": [
        "Only 21 crossover trials with 433 total participants were included, leaving many domain and subgroup estimates based on small evidence bases.",
        "GRADE certainty was low for executive function, memory, information processing and attention and very low for global cognition.",
        "The evidence concerns acute experimental responses, not long-term cognitive outcomes, productivity or real-world work performance.",
        "Intervention protocols varied substantially in frequency, intensity, duration and mode.",
        "Subgroup and meta-regression findings were exploratory and should not be converted into individualized dosing rules.",
        "Some apparently favorable domain estimates may be sensitive to the small number of studies and heterogeneous cognitive tasks."
      ],
      "target_hack_ids": [
        "break-up-prolonged-sitting"
      ],
      "target_protocol_ids": [
        "brali:break-up-prolonged-sitting"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Evidence role: watch-level boundary for the existing protocol. Keep the action simple—interrupt long sitting with feasible brief movement—but do not sell the protocol as an attention or productivity hack and do not infer an exact timer from exploratory moderator analyses."
    },
    {
      "schema_version": 1,
      "id": "wellbeing-multiple-pathways-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/wellbeing-multiple-pathways-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/wellbeing-multiple-pathways-boundary-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-09-05",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "A systematic review and network meta-analysis of randomized controlled trials of well-being-focused interventions",
        "url": "https://www.nature.com/articles/s41562-025-02369-1",
        "type": "systematic-review-network-meta-analysis",
        "doi": "10.1038/s41562-025-02369-1",
        "citation_text": "Wilkie L, Fisher Z, Geidel A, et al. Nature Human Behaviour. 2026;10:715-726.",
        "study_design": "Preregistered systematic review and random-effects network meta-analysis (PROSPERO CRD42023403480) of randomized controlled trials. Searches of MEDLINE, PsycINFO, CENTRAL and Scopus through March 2023 identified 183 trials used in the final network meta-analysis. Risk of bias was assessed with RoB 2, network confidence with CINeMA, and sensitivity analyses excluded high-risk and small studies, restricted outcomes, and added grey literature.",
        "population": "22,811 adults without diagnosed conditions. Mean participant age was 38.3 years with a reported range of 18-82. Studies were conducted in universities, workplaces, communities and online settings; 79% were conducted in Western countries. Demographic reporting was incomplete across many trials.",
        "intervention_or_exposure": "Structured well-being interventions including mindfulness-based, compassion-based, positive-psychology, exercise, yoga, ACT, educational, nature-based and some combined psychological-plus-exercise programmes, compared directly or indirectly through a connected randomized-trial network.",
        "outcomes": [
          "Validated well-being outcomes",
          "Relative intervention effects versus control",
          "Network treatment rankings",
          "Risk of bias",
          "Certainty of network comparisons",
          "Sensitivity analyses"
        ]
      },
      "supported_claim": "For general adult well-being, Brali can reasonably offer more than one small practice rather than pretending there is a single best technique. In this network meta-analysis, several established approaches—including mindfulness, exercise, yoga, compassion-based and positive-psychology interventions—showed moderate benefits versus inactive controls, while many head-to-head differences were not statistically clear. This supports the existing short-meditation protocol as one optional practice and supports a broader Brali rule: choose a feasible evidence-informed route, test it, and keep the one that is useful in the user's context.",
      "unsupported_or_overstated_claims": [
        "Mindfulness is the single best well-being intervention for everyone.",
        "A five-minute meditation is an evidence-established dose.",
        "Exercise-plus-psychological interventions are definitively superior; that node contained only three studies and two were high risk of bias.",
        "These interventions treat depression, anxiety or another diagnosed mental-health condition; clinical samples were excluded.",
        "Nature-based practices do not work; the nature node was small, heterogeneous and weakly connected to the network.",
        "Treatment rankings prove one intervention should be preferred clinically or personally over another.",
        "Immediate post-intervention gains prove durable long-term improvement.",
        "The pooled effect sizes predict how much one individual will benefit."
      ],
      "limitations": [
        "The literature search ended in March 2023 even though the synthesis was published in 2026, so it is a current synthesis of an older evidence window rather than a review of trials through 2025-2026.",
        "Only 12 of 183 studies were rated low risk of bias overall; 110 (60%) were rated high risk.",
        "Only 33% of studies used intention-to-treat analysis.",
        "Funnel-plot asymmetry and Egger tests suggested publication bias, although sensitivity analyses did not materially change the overall pattern.",
        "Many trials reported only immediate post-intervention outcomes, leaving long-term durability poorly characterized.",
        "Many network comparisons relied on indirect evidence rather than direct head-to-head trials.",
        "Seventy-nine percent of studies were conducted in Western countries and demographic reporting was incomplete.",
        "Broad intervention nodes combined programmes that can differ materially in content, dose and delivery, so the review cannot identify the active ingredient or an optimal protocol duration."
      ],
      "target_hack_ids": [
        "5-minute-meditation-habit-tracker"
      ],
      "target_protocol_ids": [
        "brali:5-minute-meditation-habit-tracker"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Editorial role: strengthen the short-meditation page without turning it into a sales pitch for meditation. The useful lesson is pluralism: structured mindfulness is a defensible option, but exercise, yoga, compassion and positive-psychology approaches also have evidence. Brali should explain the mechanism/action clearly, invite a small experiment, and make switching practices normal when one approach is unhelpful."
    },
    {
      "schema_version": 1,
      "id": "records-maintenance-precommitment-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/records-maintenance-precommitment-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/records-maintenance-precommitment-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-04",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Using behavioral science to improve compliance in government records management: Evidence from a large field experiment",
        "url": "https://doi.org/10.1016/j.giq.2026.102106",
        "type": "primary-study",
        "doi": "10.1016/j.giq.2026.102106",
        "citation_text": "Hopkins VC, Kormos C. Government Information Quarterly. 2026;43(2):102106.",
        "study_design": "Pre-registered randomized field experiment using repeated administrative observations of employee network-drive storage from January through December 2021. Employees were randomized in June to an injunctive-norm intervention, a personal-incentive intervention, or status quo control; the control group received the injunctive-norm intervention later in October, providing an additional within-study replication pattern. The treatment was a bundle rather than a factorial test of individual components.",
        "population": "22,774 employees across 35 agencies in one Canadian provincial government. Employees used centrally managed network drives, allowing monthly storage measurement. Data were deidentified, so treatment-effect heterogeneity by age, gender, tenure or occupation could not be examined.",
        "intervention_or_exposure": "An initial behaviorally framed email and calendar invitation. Accepting the invitation created a recurring 10-minute weekly records-management appointment for two months. Each reminder included one expert-approved practical tip, such as emptying the recycling bin, finding large duplicate files or moving sensitive files to secure government archives. Initial messaging emphasized either personal benefits or institutional/injunctive norms.",
        "outcomes": [
          "Monthly network-drive storage volume per employee",
          "Persistence of storage reduction after intervention",
          "Relative performance of personal-benefit versus injunctive-norm framing"
        ]
      },
      "supported_claim": "In the studied Canadian government setting, a multi-component package combining a recurring calendar commitment, timely records-management guidance and behavioral framing reduced network-drive storage relative to status quo control, with effects persisting over subsequent months. This supports testing a bounded Brali protocol that reserves a recurring maintenance block and uses one organization-approved records action per session when retention and archival rules are already known.",
      "unsupported_or_overstated_claims": [
        "Ten minutes every week is the universally optimal records-maintenance dose.",
        "A calendar reminder alone produced the observed effect.",
        "Personal-benefit framing is superior to institutional or injunctive-norm framing.",
        "The intervention proves employees deleted only records that should have been deleted.",
        "Deleting old, duplicated or large files indiscriminately is safe.",
        "The intervention has been shown to improve retrieval speed, reduce security incidents, increase productivity or improve information-access compliance directly.",
        "The reported monetary savings generalize to other organizations.",
        "The intervention is already validated for private-sector organizations or personal knowledge-management systems.",
        "Generative AI or automation should automatically delete or archive records on the basis of this study."
      ],
      "limitations": [
        "The primary outcome was storage volume, an objective administrative measure but only a proxy for records-policy compliance; individual file decisions were not observed.",
        "The intervention bundled pre-commitment, recurring reminders, weekly practical guidance and motivational framing, so the causal contribution of each component cannot be isolated.",
        "The study took place in one Canadian provincial government, whose accountability structures, retention requirements, infrastructure and employment context may differ from private, nonprofit or personal settings.",
        "Deidentified data prevented subgroup analysis by employee age, gender, tenure or occupation.",
        "The unbalanced panel experienced attrition over time, although reported robustness checks did not materially change the findings.",
        "The study did not test alternative maintenance frequencies or durations, so the 10-minute weekly schedule should not be treated as an optimized dose.",
        "The study did not directly measure retrieval quality, data-breach risk, employee productivity or downstream information-access outcomes.",
        "The treatment worked in an environment with mandatory records-management training and expert-approved retention guidance; applying cleanup actions without equivalent policy knowledge could create harm."
      ],
      "target_hack_ids": [
        "records-maintenance-block"
      ],
      "target_protocol_ids": [
        "brali:records-maintenance-block"
      ],
      "risk_flags": [],
      "notes": "Protocol-design direction: reserve a short recurring records-maintenance appointment; at each session execute one approved, low-ambiguity action such as clearing disposable material, resolving duplicates or moving sensitive records to the correct archive. Treat 10 minutes weekly for two months as the tested implementation, not a universal prescription. Where retention rules are unclear, stop and resolve policy rather than deleting. Any future AI assistance should surface candidates and preserve human accountability for retention or deletion decisions."
    },
    {
      "schema_version": 1,
      "id": "leisure-crafting-work-spillover-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/leisure-crafting-work-spillover-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/leisure-crafting-work-spillover-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-03",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The leisure crafting intervention: Effects on work and non-work outcomes and the moderating role of age",
        "url": "https://journals.sagepub.com/doi/10.1177/00187267251407641",
        "type": "primary-study",
        "doi": "10.1177/00187267251407641",
        "citation_text": "Petrou P, Den Dulk L, Michaelides G. Human Relations. First published online 8 January 2026.",
        "study_design": "Five-week randomized online intervention with repeated weekly measurements and a passive control group. Participants were randomized after baseline. The intervention group created a leisure-based personal development plan and reflected on it weekly; growth trajectories were estimated with Bayesian multilevel models.",
        "population": "462 employed adults from multiple occupational sectors in the Netherlands recruited through an internet research panel. Participants worked at least three days per week and either had a leisure activity or were willing to find one. The analyzed intervention group included 196 participants and the passive control 266. Retention criteria required sufficient survey participation and leisure activity; intervention participants also had to report more than minimal effort.",
        "intervention_or_exposure": "A brief educational video introduced leisure crafting. Participants chose a leisure activity and wrote a four-week personal development plan covering three elements: an autonomously chosen goal, something to learn or develop, and human connection. At each weekly survey they reflected on progress, what happened during the activity and what they wanted to adjust for the following week.",
        "outcomes": [
          "Leisure crafting",
          "Meaning at work",
          "Employee creativity",
          "Work engagement",
          "Meaning in life",
          "Need satisfaction",
          "Affective well-being",
          "Sense of community"
        ]
      },
      "supported_claim": "In this Dutch employee sample, deliberately structuring a leisure activity around autonomous goal-setting, learning and human connection, with weekly reflection, produced greater five-week growth than passive control in self-reported meaning at work and employee creativity. Brali can therefore test a bounded 'craft your leisure' protocol as a personal-development experiment whose outcome is reviewed rather than assumed.",
      "unsupported_or_overstated_claims": [
        "Leisure crafting has been shown to improve objective job performance or productivity.",
        "The intervention reliably improves work engagement, meaning in life, need satisfaction or sense of community.",
        "The intervention generally improves affective well-being for working adults of all ages; the observed well-being effect was age-dependent.",
        "Sports, collective activities or any particular hobby are superior choices.",
        "Five weeks, one weekly reflection, or any specific activity frequency is an optimal dose.",
        "Autonomous goal-setting, learning or human connection has been isolated as the causal ingredient; the intervention tested them as a bundle.",
        "Leisure should be optimized primarily for work output rather than enjoyment, recovery or personal values."
      ],
      "limitations": [
        "The comparator was passive, so expectancy, attention and structured reflection were not controlled.",
        "All reported outcomes were self-report and measured at the same weekly occasions; objective performance was not measured.",
        "The study lasted five weeks and modeled linear growth, so persistence and longer-term trajectories are unknown.",
        "Participants were Dutch employees recruited from one internet panel, limiting cultural and occupational generalization.",
        "Post-randomization retention criteria included leisure participation and, for intervention participants, self-reported intervention effort; this can select for participants who engaged with the program.",
        "Activity type, frequency and intensity were not experimentally varied, and most intervention participants chose sports activities.",
        "The authors note that goal-setting may have been overrepresented relative to learning and human connection in how participants engaged with the intervention."
      ],
      "target_hack_ids": [
        "craft-your-leisure"
      ],
      "target_protocol_ids": [
        "brali:craft-your-leisure"
      ],
      "risk_flags": [],
      "notes": "Protocol-design direction: choose one leisure activity because it matters outside work; add one self-chosen developmental target, one learning edge and one connection element where natural; review progress weekly for several weeks and separately record whether any useful work spillover appears. Do not turn leisure into a productivity obligation, and stop or simplify if the structure removes enjoyment or recovery value."
    },
    {
      "schema_version": 1,
      "id": "supervisory-feedback-characteristics-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/supervisory-feedback-characteristics-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/supervisory-feedback-characteristics-boundary-2026/index.json",
      "decision": "watch",
      "decision_label": "Watch, do not prescribe",
      "reviewed_at": "2026-09-03",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Examining the relationship between supervisory feedback characteristics and employee feedback processing: a systematic review and meta-analysis",
        "url": "https://link.springer.com/article/10.1007/s12144-026-09323-y",
        "type": "systematic-review-meta-analysis",
        "doi": "10.1007/s12144-026-09323-y",
        "citation_text": "Zyberaj J. Current Psychology. 2026;45:1221.",
        "study_design": "Systematic review and random-effects multilevel meta-analysis of 24 peer-reviewed empirical studies contributing 75 effect sizes across 26 supervisory-feedback characteristics. Correlations were synthesized separately for feedback perception, acceptance, desire to respond and intended response, while accounting for dependent effect sizes.",
        "population": "Employees represented across the 24 included workplace studies; the pooled participant count was 595,950, heavily influenced by at least one very large sample. Context, culture, channel and feedback-frequency information were not reported consistently enough for robust moderator analysis.",
        "intervention_or_exposure": "Observed characteristics of supervisory feedback and feedback sources, including credibility, feedback quality, valence, support and other source/message features. This was primarily an observational evidence base rather than a synthesis of randomized feedback interventions.",
        "outcomes": [
          "Perceived feedback",
          "Feedback acceptance",
          "Desire to respond",
          "Intended response"
        ]
      },
      "supported_claim": "Across the reviewed workplace literature, supervisor credibility and feedback quality were consistently positively associated with several stages of employee feedback processing. These are useful design hypotheses for Brali feedback protocols, but the synthesis does not establish that changing credibility, specificity or constructiveness will causally improve acceptance, behavior, performance or learning.",
      "unsupported_or_overstated_claims": [
        "Training supervisors to appear more credible has been shown by this meta-analysis to cause better employee performance.",
        "Specific and constructive feedback is proven here to causally increase learning or behavior change.",
        "One delivery channel, feedback cadence or communication style is optimal.",
        "Positive-valence feedback is generally better than corrective or negative feedback.",
        "The pooled correlations can be interpreted as intervention effect sizes.",
        "The findings generalize uniformly across cultures, roles, remote work and digital feedback systems."
      ],
      "limitations": [
        "Most included studies were correlational and relied heavily on self-report, limiting causal inference and raising common-method concerns.",
        "Between-study heterogeneity was very high (I²=95%).",
        "Nine studies had potential quality or risk-of-bias limitations under the review's appraisal.",
        "Only peer-reviewed studies were included, so unpublished and grey literature were excluded.",
        "Context variables such as delivery channel, feedback frequency and organizational climate were too inconsistently reported for robust meta-regression.",
        "Some characteristics and processing facets were represented by only a small number of studies, constraining fine-grained conclusions."
      ],
      "target_hack_ids": [],
      "target_protocol_ids": [],
      "risk_flags": [],
      "notes": "Keep this as a research boundary and design input for the existing work-feedback search lens. Before Brali publishes a causal feedback-delivery protocol, prioritize randomized or strong longitudinal studies that manipulate message quality/source behavior and measure downstream behavior or performance rather than only acceptance intentions."
    },
    {
      "schema_version": 1,
      "id": "vocabulary-pretesting-guess-feedback-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/vocabulary-pretesting-guess-feedback-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/vocabulary-pretesting-guess-feedback-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-02",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Duolingo-inspired pretesting with words and pictures improves vocabulary learning",
        "url": "https://link.springer.com/article/10.1186/s41235-026-00708-y",
        "type": "primary-study",
        "doi": "10.1186/s41235-026-00708-y",
        "citation_text": "Chua TJE, Pan SC. Cognitive Research: Principles and Implications. 2026;11:20.",
        "study_design": "Four preregistered controlled online experiments comparing within-person pretesting versus reading for Spanish vocabulary. Experiments 1 and 3 used intermixed criterial tests; Experiments 2 and 4 additionally randomized participants to blocked versus intermixed test arrangements. Items and condition assignments were randomized and counterbalanced.",
        "population": "Adult Prolific participants aged 21–45, fluent in English and resident in English-speaking countries. Final samples were 58 in Experiment 1, 123 in Experiment 2, 47 in Experiment 3 and 113 in Experiment 4 after exclusions for prior Spanish knowledge, comprehension failures, non-completion or off-task behavior.",
        "intervention_or_exposure": "For word-image learning, learners saw a Spanish word and chose its meaning from four images; for image-word learning, they saw an image and chose the Spanish word. They had up to 8 seconds to guess and then 5 seconds of correct-answer feedback. Reading trials showed the correct pair directly for 5 seconds. Memory was tested after five one-minute distractor tasks using cued recall and multiple choice.",
        "outcomes": [
          "Short-delay cued recall of word meanings or Spanish words",
          "Short-delay multiple-choice recognition",
          "Learner preference and metacognitive postdictions"
        ]
      },
      "supported_claim": "For adults learning new, concrete second-language vocabulary paired with images, adding a forced multiple-choice guess before immediately revealing the correct pairing can produce a modest short-delay memory advantage over simply reading the pair. A bounded Brali protocol may therefore use a cue → guess → immediate correct feedback → later recall sequence for vocabulary practice, while keeping the claim limited to this learning format and time horizon.",
      "unsupported_or_overstated_claims": [
        "Making errors before learning is generally better than studying correct information.",
        "Pretesting improves grammar, reading comprehension, conversation skill, pronunciation or complex skill learning.",
        "The observed advantage has been shown to persist over days, weeks or months.",
        "The exact 8-second guess window, 5-second feedback window, four-option format or image format is optimal.",
        "Multiple-choice pretesting always improves later multiple-choice performance; Experiment 3 did not show a significant multiple-choice advantage.",
        "The effect proves that error correction, associative strengthening or any single proposed mechanism caused the benefit.",
        "The findings generalize to children, advanced language learners, non-concrete vocabulary or learners outside the sampled online populations."
      ],
      "limitations": [
        "All criterial tests followed a short same-session distractor period rather than a delayed retention interval, so long-term retention was not established.",
        "The studies used concrete Spanish nouns paired with images; vocabulary type and language-learning context were narrow.",
        "Participants were adult Prolific users from English-speaking countries and exclusions were substantial in several experiments.",
        "Pretesting trials included up to 8 seconds of cue-only guessing before the same 5 seconds of correct pair exposure used in reading, so total cue exposure was longer in the pretesting condition.",
        "Initial guess accuracy was about 35–38%, raising the possibility of some prior familiarity despite exclusions; the authors note that benefits also appeared for incorrectly guessed items.",
        "Learning condition was primarily within-subjects, and the authors identify between-subject replication as a future research need.",
        "Multiple-choice benefits were not significant in Experiment 3, so recognition effects were less uniform than cued-recall effects."
      ],
      "target_hack_ids": [
        "guess-before-reveal-vocabulary"
      ],
      "target_protocol_ids": [
        "brali:guess-before-reveal-vocabulary"
      ],
      "risk_flags": [],
      "notes": "Protocol-design direction: use this only when introducing unfamiliar concrete L2 vocabulary. Show a cue, require a real guess before the answer is visible, reveal the correct answer immediately, then test unaided recall later in the session. Do not reward confident guessing as correctness, and do not replace spaced later review with this one-pass encoding technique."
    },
    {
      "schema_version": 1,
      "id": "custom-social-media-brake-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/custom-social-media-brake-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/custom-social-media-brake-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-01",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Promoting Self-Regulated Social Media Use on Smartphones With a Mobile Intervention App (Wellspent): Randomized Controlled Trial",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC13062480/",
        "type": "primary-study",
        "doi": "10.2196/56824",
        "citation_text": "Mertens L, Brockmeier LC, Roitzheim C, Radtke T, Dingler T, Keller J. JMIR Mhealth Uhealth. 2026;14:e56824.",
        "study_design": "Preregistered 3-week randomized controlled trial. Seventy iPhone users were randomized 1:1 to a customizable digital self-control app or a no-treatment control. The intervention was active in week 2 and optional in week 3. Linear mixed models were used across baseline and two follow-ups; formal blinding was not possible.",
        "population": "70 English-proficient adult iPhone users who regularly used at least one social-media app, recruited through social channels and Freie Universität Berlin study-opportunity mailing lists. Mean age was 26.2 years, 67% were women, and most were students or young professionals. Fifty-two participants completed week 3; the primary problematic-social-media-use model contained 46 participants.",
        "intervention_or_exposure": "Users selected social-media apps they considered problematic, set daily usage goals and continuous-use nudge intervals, customized reminder frequency and tone, chose an alternative activity and could define vulnerable time windows. After a self-defined interval, a full-screen reminder required an active dismissal while preserving the choice to continue. The app also supplied feedback and self-monitoring.",
        "outcomes": [
          "Objectively recorded daily minutes on the self-identified most problematic social-media app, entered by participants from iOS screen-time records",
          "Problematic smartphone use",
          "Problematic social-media use",
          "Self-efficacy to self-regulate social-media use",
          "Reminder acceptance and voluntary continued app use"
        ]
      },
      "supported_claim": "For an adult who already wants to spend less time in a specific social-media app, a short trial of a user-configured digital brake is reasonable: choose the target app, choose a personally meaningful session threshold, make continued use require an explicit quit-or-continue decision, and identify an alternative activity while retaining the ability to override the prompt. In this small randomized trial, the bundled intervention reduced time on the target app over the short study period. The source supports the package as a behavior-change experiment; it does not identify the full-screen checkpoint, customization, self-monitoring or any other component as the unique cause.",
      "unsupported_or_overstated_claims": [
        "A full-screen quit-or-continue reminder by itself has been proven to reduce social-media use.",
        "Ten minutes, 45 minutes, three-minute reminder repeats, or any other example threshold in the app is an evidence-based optimal setting.",
        "Reducing app time in this study improved wellbeing, attention, productivity, sleep or mental health.",
        "The intervention treated social-media addiction or another clinical condition.",
        "The intervention reliably increased self-control or self-efficacy.",
        "The result generalizes to Android users, older populations, all apps or people who do not want to reduce their use.",
        "Short-term reduced use will persist after the intervention is removed.",
        "Blocking access more aggressively would necessarily work better."
      ],
      "limitations": [
        "The randomized sample was small at 70 participants, and the study did not reach its original recruitment target.",
        "Twenty-six percent did not complete week 3; the final model for the primary problematic-social-media-use outcome included 46 participants.",
        "Participants were iPhone users, mostly students or young professionals, and self-selected into a study about social-media self-regulation.",
        "The control group received no active comparator and participants could not be blinded.",
        "The intervention bundled several behavior-change components, so the study cannot isolate the causal contribution of the quit-or-continue checkpoint, customization, alternative activities, goals, feedback or self-monitoring.",
        "Some psychological outcomes were self-reported; problematic social-media use and self-efficacy did not show robust improvement.",
        "The intervention period was short and there was no long-term follow-up establishing durable habit change.",
        "One author was employed by Wellspent GmbH during the intervention period and another was a company cofounder."
      ],
      "target_hack_ids": [
        "custom-social-media-brake"
      ],
      "target_protocol_ids": [
        "brali:custom-social-media-brake"
      ],
      "risk_flags": [
        "health",
        "mental-health"
      ],
      "notes": "Protocol-design direction: use this as a lower-friction alternative to Brali's stronger phone-offline experiment, not as a replacement. Ask the user to identify one app they already want to use less, choose a personally relevant continuous-use checkpoint, select a realistic alternative action, and force a deliberate continue-or-exit choice without hard blocking. Review actual target-app minutes and whether prompts were useful after a bounded trial. Do not prescribe a universal session length or infer mental-health benefit."
    },
    {
      "schema_version": 1,
      "id": "if-then-fruit-vegetable-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/if-then-fruit-vegetable-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/if-then-fruit-vegetable-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-09-01",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "A systematic review and meta-analysis on the effectiveness of if-then plans – in a strict sense – to facilitate fruit and vegetable consumption in adults",
        "url": "https://link.springer.com/article/10.1186/s12966-026-01915-y",
        "type": "meta-analysis",
        "doi": "10.1186/s12966-026-01915-y",
        "citation_text": "Melum SK, Martiny-Huenger T. Int J Behav Nutr Phys Act. 2026;23:51.",
        "study_design": "Systematic review and random-effects meta-analysis of randomized controlled trials comparing strict if-then planning with active control conditions. Searches of MEDLINE, Embase and PsycInfo were last run April 3, 2025. Ten articles contributed 12 comparisons and 2,399 participants, with outcomes measured from 1 week to 24 months.",
        "population": "Adult samples from studies conducted mainly in western Europe, plus one Canadian study. Four studies recruited university students and five recruited broader adult samples such as employees or members of a health-insurance population. Across studies, participants were predominantly young-to-middle-aged, female and Caucasian.",
        "intervention_or_exposure": "Strict if-then planning that explicitly linked a perceivable situational cue with a goal-directed response relevant to increasing fruit or vegetable intake. Examples across the reviewed interventions included planning when and where to buy, prepare, pack or eat fruit and vegetables. Active controls received similar health information and encouragement without the if-then linkage.",
        "outcomes": [
          "Self-reported daily fruit and vegetable intake",
          "Pooled mean difference between if-then planning and active control at the latest reported outcome"
        ]
      },
      "supported_claim": "When an adult already intends to eat more fruit or vegetables, adding one or more explicit cue-response plans is a defensible low-resource experiment: identify a recurring situation that can be noticed and link it to a concrete action under the person's control, such as preparing, packing, buying or adding the intended food. Across 12 randomized active-control comparisons, strict if-then planning produced a small average increase in self-reported fruit-and-vegetable intake. Brali should present this as a planning aid with a modest expected effect, not as a diet prescription.",
      "unsupported_or_overstated_claims": [
        "If-then planning produces a large change in diet.",
        "The observed increase in self-reported fruit and vegetable intake has been shown to prevent disease or improve clinical health outcomes.",
        "A particular cue, meal, food, portion count, reminder schedule or number of if-then plans is optimal.",
        "If-then planning is superior to all other nutrition interventions.",
        "The evidence establishes effectiveness for weight loss, calorie reduction, protein intake, supplements or other dietary targets.",
        "The effect generalizes equally across cultures, age groups, sexes or populations not represented in the included trials.",
        "The review validates medical nutrition advice for people with allergies, eating disorders, metabolic disease or other conditions requiring individualized care.",
        "Self-reported intake in the trials is equivalent to objectively verified dietary change."
      ],
      "limitations": [
        "All included fruit-and-vegetable intake outcomes relied on retrospective self-report, often one or two questions estimating portions per day.",
        "The included samples were predominantly western European, female, Caucasian and young-to-middle-aged, limiting generalizability.",
        "The pooled benefit was small: about 0.29 self-reported portions per day relative to active controls.",
        "The review's protocol was not prospectively registered; an unpublished protocol existed and the authors narrowed the planned nutrition outcome to fruit and vegetables before analysis.",
        "The review's strict intervention definition required judgment about whether study instructions truly linked a perceivable cue to a response.",
        "The evidence concerns fruit and vegetable intake specifically and cannot be extended to other nutrition outcomes.",
        "The meta-analysis does not establish the clinical relevance of the modest reported intake increase."
      ],
      "target_hack_ids": [
        "if-then-produce-plan"
      ],
      "target_protocol_ids": [
        "brali:if-then-produce-plan"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Protocol-design direction: first confirm that the user already wants to increase fruit or vegetables; this is a volitional aid, not persuasion. Pick one repeatable cue and one visible response, for example: 'If I pack lunch tonight, then I will add one fruit' or 'If I serve dinner, then I will add a vegetable.' Track whether the cue and response occurred. Avoid medical or weight-loss claims and do not prescribe a fixed portion target from this source."
    },
    {
      "schema_version": 1,
      "id": "chatbot-difficult-conversation-rehearsal-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/chatbot-difficult-conversation-rehearsal-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/chatbot-difficult-conversation-rehearsal-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-31",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Harnessing Artificial Intelligence to Facilitate Difficult Conversations",
        "url": "https://journals.sagepub.com/doi/10.1177/19485506261444068",
        "type": "preregistered-randomized-behavioral-experiment",
        "doi": "10.1177/19485506261444068",
        "citation_text": "Maheshka C, Folk D, Dunn E. Social Psychological and Personality Science. First published online May 5, 2026.",
        "study_design": "Preregistered two-part randomized behavioral study. Participants first identified a real difficult conversation they were considering, then were randomized to spend 6-12 minutes preparing it with a chatbot or to ask the chatbot three neutral factual questions. Of 1,465 participants completing Part 1, 1,371 returned for a follow-up survey beginning five days later. The primary behavioral outcome was whether the participant reported having the difficult conversation in the intervening period.",
        "population": "Adult Prolific workers in the United States and Canada who could identify a difficult conversation they were considering having. Most target conversations involved romantic partners, family or friends; about one-fifth were professional conversations.",
        "intervention_or_exposure": "A short GPT-4o-mini conversation used to prepare a real interpersonal conversation. Participants commonly planned what to say, role-played the other person, sought phrasing or contingency advice, and anticipated possible reactions. Control participants interacted with the chatbot on unrelated factual questions after identifying and describing their difficult conversation.",
        "outcomes": [
          "Self-reported initiation of the difficult conversation during follow-up",
          "Anticipated conversation quality",
          "Perceived importance, difficulty and anxiety",
          "Post-intervention affect",
          "Among participants who conversed: satisfaction and perceived quality, difficulty and thoroughness"
        ]
      },
      "supported_claim": "For adults already considering an ordinary difficult conversation, a brief chatbot preparation session modestly increased the probability that they reported actually initiating it over the following week: 65% in the preparation condition versus 59% in control. Brali can therefore test a bounded rehearsal protocol whose goal is to reduce the initiation barrier: clarify the goal, draft an opening, role-play a likely response and prepare one or two contingency moves, then return to the real human conversation.",
      "unsupported_or_overstated_claims": [
        "Chatbot rehearsal makes difficult conversations go better once they occur.",
        "AI resolves interpersonal conflict or improves relationship quality.",
        "AI preparation is better than preparing with a friend, coach, therapist, manager or journal.",
        "The 6-12 minute duration is an established optimal dose.",
        "The result generalizes unchanged outside US and Canadian adults familiar enough with chatbots to complete the intervention.",
        "The protocol is appropriate for abuse, coercion, threats, harassment, legal disputes, HR investigations or other situations requiring safety planning, professional support or a durable written record.",
        "A specific chatbot model or prompt is necessary to obtain the effect."
      ],
      "limitations": [
        "The primary real-world outcome was participant self-report rather than independently observed conversation behavior.",
        "The absolute difference in conversation initiation was modest: six percentage points.",
        "The experimental and control chatbot interactions differed in duration, engagement and prompting, so the study does not isolate which preparation component caused the effect.",
        "The sample was limited to US and Canadian Prolific participants, restricting cultural generalizability.",
        "Among participants who did have the difficult conversation, the study found no significant differences in satisfaction, perceived quality, difficulty or thoroughness.",
        "The study followed behavior for roughly one week and does not establish repeated-use or long-term effects.",
        "Preparing the conversation increased negative affect immediately after the exercise, and repeated reliance on a chatbot was not studied.",
        "Safety-sensitive, coercive, legal and high-power-imbalance conversations were not established as appropriate use cases."
      ],
      "target_hack_ids": [
        "rehearse-difficult-conversation"
      ],
      "target_protocol_ids": [
        "brali:rehearse-difficult-conversation"
      ],
      "risk_flags": [],
      "notes": "Protocol direction: use AI as a private rehearsal surface, not as the relationship partner or final authority. Ask the user to name the conversation goal, draft one opening sentence, role-play a plausible response, identify the main failure mode and prepare one recovery phrase. The protocol should stop or redirect when safety, coercion, harassment, legal/HR sensitivity or professional mental-health support is relevant. Success should initially be measured as whether the intended safe conversation was initiated, not whether AI supposedly improved the relationship."
    },
    {
      "schema_version": 1,
      "id": "structured-peer-feedback-provision-2025",
      "canonical_url": "https://brali-lifeos.github.io/evidence/structured-peer-feedback-provision-2025/",
      "json_url": "https://brali-lifeos.github.io/evidence/structured-peer-feedback-provision-2025/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-08-31",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Enhancing the Peer-Feedback Process Through Instructional Support: A Meta-Analysis",
        "url": "https://link.springer.com/article/10.1007/s10648-025-10017-3",
        "type": "meta-analysis",
        "doi": "10.1007/s10648-025-10017-3",
        "citation_text": "Hornstein J, Keller MV, Greisel M, Dresel M, Kollar I. Educational Psychology Review. 2025;37:42.",
        "study_design": "Meta-analysis of 32 peer-reviewed journal studies with 3,806 learners, comparing peer-feedback with instructional support against peer-feedback without instructional support. Robust variance estimation was used to handle dependent effect sizes, and analyses distinguished support during feedback provision versus reception and content-specific versus generic support.",
        "population": "Learners in the experimental and quasi-experimental educational peer-feedback studies included in the review. Tasks and educational contexts varied across studies.",
        "intervention_or_exposure": "Instructional support around peer feedback, including preparatory activities, rubrics or rating schemes, sentence starters, guiding questions, integrated support and broader content-specific or generic scaffolds, compared with unsupported peer-feedback processes.",
        "outcomes": [
          "Quality of feedback provision",
          "Quality of feedback reception or revision",
          "Subject-matter-related knowledge"
        ]
      },
      "supported_claim": "In the reviewed educational literature, adding instructional support to peer feedback improved the peer-feedback process overall, and support aimed at the person providing feedback was associated with better feedback-provision quality. This supports Brali's existing educational writing protocol in making the review target explicit and giving the peer reviewer a small amount of structure instead of requesting vague general feedback.",
      "unsupported_or_overstated_claims": [
        "One specific rubric, sentence starter, checklist or guiding-question set is proven best.",
        "Instructional support for receiving feedback has been shown ineffective; that evidence base was too small for a stable conclusion.",
        "Structured peer feedback reliably increases subject-matter knowledge in every setting.",
        "The meta-analysis establishes the same effects for professional, technical, marketing or creative writing outside educational contexts.",
        "More detailed reviewer instructions are always better.",
        "Peer feedback is universally superior to instructor, expert, self or automated feedback."
      ],
      "limitations": [
        "Only 32 journal studies met the inclusion criteria, limiting fine-grained moderator analysis.",
        "Only three studies with 14 effect sizes examined feedback-reception support, making conclusions about that phase underpowered and unstable.",
        "The authors had to collapse specific supports such as rubrics, sentence starters and guiding questions into broader categories because each was represented by too few studies.",
        "Study heterogeneity was high, and publication-bias tests were mixed: Egger's test was significant whereas Begg's test was not and the funnel plot was relatively symmetrical.",
        "The meta-analysis concerns educational peer-feedback processes and should not be treated as direct evidence for all professional writing or workplace review contexts.",
        "Coarse outcome categories do not reveal which exact support mechanism is best for a specific writing task."
      ],
      "target_hack_ids": [
        "ask-a-peer-to-proofread"
      ],
      "target_protocol_ids": [
        "brali:ask-a-peer-to-proofread"
      ],
      "risk_flags": [],
      "notes": "Evidence role: support the existing 'match feedback to the revision target' protocol, especially its instruction to tell the reviewer what to inspect. Do not yet add a canonical universal rubric. A future content revision may offer optional lightweight structures such as two or three guiding questions, while keeping the educational-context boundary visible."
    },
    {
      "schema_version": 1,
      "id": "task-choice-response-latency-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/task-choice-response-latency-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/task-choice-response-latency-boundary-2026/index.json",
      "decision": "watch",
      "decision_label": "Watch, do not prescribe",
      "reviewed_at": "2026-08-31",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Give Me a Choice! A Field Experiment on Task Choice Enabled by Wearables",
        "url": "https://journals.sagepub.com/doi/10.1177/10591478251400469",
        "type": "plant-level-field-experiment-difference-in-differences",
        "doi": "10.1177/10591478251400469",
        "citation_text": "Kwasnitschka D, Franke H, von Maydell R, Netland T. Production and Operations Management. First published online November 12, 2025.",
        "study_design": "Field study in two comparable manufacturing plants. The plant receiving the choice interface was selected by a plant-level random allocation, and pre/post outcomes were estimated using Difference-in-Differences with machine and time controls and robustness checks. The dataset contained 31,429 work tasks and 66,233 machine status reports; the post-treatment observation window covered three weeks.",
        "population": "Production workers in two plants of one global automotive-industry supplier in Germany and Italy, working in highly automated molding operations and responding to short machine-interruption troubleshooting tasks through wearable devices.",
        "intervention_or_exposure": "The digital assignment interface changed from presenting a single delegated task to presenting a list of all currently available tasks matching the worker's codified skills. Other assignment-system features were intended to remain constant.",
        "outcomes": [
          "Task response time",
          "Accepted-task completion time",
          "Machine-level productivity",
          "Exploratory task-choice behavior",
          "Auxiliary single-item satisfaction with task allocation"
        ]
      },
      "supported_claim": "In this specific automated manufacturing setting, giving workers task choice changed where time was spent rather than improving aggregate productivity: response time rose by about 80 seconds, accepted-task completion time fell by about 26 seconds, and machine productivity did not significantly change. This makes a hybrid choice-versus-delegation design worth further study: offer bounded choice where response latency is not critical, while preserving direct ownership for urgent tasks.",
      "unsupported_or_overstated_claims": [
        "Task autonomy increases productivity in general.",
        "Workers should always choose their own next task.",
        "Direct assignment is generally superior to choice.",
        "The observed response-time and completion-time effects generalize to knowledge work, software teams or office workflows.",
        "The study proves that limiting task choices or notifying fewer workers will improve performance; those are design implications that require testing.",
        "The reported improvement in satisfaction establishes a robust well-being effect.",
        "The results establish an optimal number of task options or an optimal urgency threshold."
      ],
      "limitations": [
        "Only two plants from one company were studied, so plant-level random allocation still leaves only two clusters and limits causal generalization.",
        "The post-treatment period was only three weeks.",
        "The assignment system could not perfectly encode worker skills, so improved matching and motivational autonomy effects cannot be fully disentangled.",
        "The tasks were short troubleshooting responses to machine interruptions, where response latency is unusually important.",
        "Productivity was measured at the machine level rather than directly for individual tasks or workers.",
        "The treatment group started from relatively high productivity, which may have limited detectable gains.",
        "The satisfaction result came from an auxiliary single-item survey and was acknowledged by the authors as anecdotal.",
        "Findings around task choice remain mixed in the broader literature, according to the authors' own discussion."
      ],
      "target_hack_ids": [
        "bounded-task-choice"
      ],
      "target_protocol_ids": [
        "brali:bounded-task-choice"
      ],
      "risk_flags": [],
      "notes": "Keep on watch rather than publishing a general autonomy protocol. Research direction: look for replications in service work, knowledge work and longer-duration tasks, and specifically test hybrid routing where urgent/latency-sensitive work is assigned while non-urgent work offers a constrained choice set. A new work-task-autonomy search lens has been added so Research Scout can pursue this boundary systematically."
    },
    {
      "schema_version": 1,
      "id": "activity-breaks-postprandial-metabolism-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/activity-breaks-postprandial-metabolism-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/activity-breaks-postprandial-metabolism-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-30",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The Acute Effects of Interrupting Prolonged Sitting With Regular Activity Breaks on Postprandial Glucose and Insulin in Adults: A Systematic Review and Meta-Analysis",
        "url": "https://onlinelibrary.wiley.com/doi/10.1111/obr.70152",
        "type": "meta-analysis",
        "doi": "10.1111/obr.70152",
        "citation_text": "Gale JT, Martin H, Haszard JJ, Peddie MC. Obesity Reviews. 2026:e70152.",
        "study_design": "Systematic review and random-effects meta-analysis of acute laboratory randomized crossover trials. Fifty-three studies met eligibility criteria and 39 were included in meta-analysis. Eligible interventions lasted at least three hours but less than 24 hours, compared prolonged sitting with at least three brief activity breaks, and measured serial postprandial glucose and/or insulin responses.",
        "population": "Adults aged at least 18 years across healthy, overweight/obesity and some clinical/metabolic-risk groups. Studies varied in age, sex, nationality, weight status, health status and laboratory protocol.",
        "intervention_or_exposure": "Interrupting prolonged sitting with regular activity breaks lasting no more than ten minutes, using standing, walking, resistance or other activity modes, compared with uninterrupted prolonged sitting of the same duration.",
        "outcomes": [
          "Postprandial glucose incremental area under the curve",
          "Postprandial insulin incremental area under the curve",
          "Subgroup estimates by activity mode",
          "Subgroup estimates by break frequency",
          "Risk-of-bias assessment"
        ]
      },
      "supported_claim": "For acute post-meal metabolism, replacing small portions of a prolonged sitting period with brief activity is better supported than remaining continuously seated. Across randomized crossover evidence, regular activity breaks reduced postprandial glucose and insulin responses, with walking producing the strongest overall mode estimates. Brali can therefore propose a bounded protocol to interrupt long sitting bouts with brief movement, especially walking when feasible, while treating the exact timing and dose as adaptable rather than universally fixed.",
      "unsupported_or_overstated_claims": [
        "A break every 20 minutes is the universal optimal schedule for all people and settings.",
        "Brief movement breaks prevent diabetes, cardiovascular disease, obesity, or mortality.",
        "The protocol improves concentration, productivity, mood, pain, or cognitive performance.",
        "Standing alone produces the same metabolic effect as walking or active movement.",
        "Brief sitting interruptions can replace recommended weekly physical activity or structured exercise.",
        "The subgroup estimates establish which BMI or clinical group will benefit most.",
        "One exact walking speed, intensity, duration, or exercise routine is required."
      ],
      "limitations": [
        "All included studies were acute laboratory randomized crossover trials lasting less than 24 hours, so the evidence does not establish long-term health effects or sustainability.",
        "Prolonged-sitting laboratory protocols, especially the longer ones, may not resemble habitual real-world sitting.",
        "Many pooled and subgroup estimates had substantial heterogeneity that remained after subgroup analysis.",
        "Most studies were rated fair or good rather than excellent; participant blinding was impossible and reporting of attrition and assessor blinding was inconsistent.",
        "The review was not prospectively registered in PROSPERO or another registry and no review protocol was prepared.",
        "The search was limited to peer-reviewed English-language studies, so relevant evidence may have been missed.",
        "Frequency and mode subgroup comparisons do not by themselves prove that the largest subgroup estimate is the optimal prescription for an individual."
      ],
      "target_hack_ids": [
        "break-up-prolonged-sitting"
      ],
      "target_protocol_ids": [
        "brali:break-up-prolonged-sitting"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Protocol direction: when a work or leisure block has become prolonged sitting, use a low-friction cue to stand up and add brief movement such as a short walk, then return to the task. Do not market a universal 20-minute timer. The protocol should coexist with, not replace, Brali's weekly movement-baseline guidance. Track whether the pattern is feasible and whether long uninterrupted sitting actually decreases."
    },
    {
      "schema_version": 1,
      "id": "hourly-microexercise-workplace-watch-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/hourly-microexercise-workplace-watch-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/hourly-microexercise-workplace-watch-2026/index.json",
      "decision": "watch",
      "decision_label": "Watch, do not prescribe",
      "reviewed_at": "2026-08-30",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Micro-exercise breaks every hour: a feasible strategy to improve metabolic health in sedentary office workers",
        "url": "https://link.springer.com/article/10.1186/s12889-026-26484-4",
        "type": "primary-study",
        "doi": "10.1186/s12889-026-26484-4",
        "citation_text": "Fang Y, Li H, Dong P, Wan F. BMC Public Health. 2026;26:763.",
        "study_design": "Twelve-week individually randomized controlled workplace trial. Eighty-six sedentary office workers were allocated 1:1 to an hourly micro-exercise routine or usual behavior. Allocation used computer-generated permuted blocks with sex/workplace stratification and concealed envelopes; outcome assessors and laboratory staff were blinded. Seventy-nine participants completed 12-week follow-up, and analyses were reported as intention-to-treat.",
        "population": "Sedentary office workers aged 25-55 from three institutions in Nanchang, China, sitting more than six hours per workday and reporting less than 150 minutes of moderate-to-vigorous activity per week. Mean BMI was about 28.5 kg/m² and 40.5% met the study's fasting-glucose criterion for prediabetes.",
        "intervention_or_exposure": "Seven 3-minute equipment-free micro-exercise breaks across an eight-hour workday, using marching, desk/wall push-ups, squats, heel raises, shoulder/arm movements and torso twists, compared with usual behavior.",
        "outcomes": [
          "Fasting blood glucose",
          "Two-hour postprandial glucose",
          "HOMA-IR",
          "Anthropometric and blood-pressure measures",
          "Lipid profile",
          "Accelerometer-measured activity and sedentary time",
          "Adherence and adverse events"
        ]
      },
      "supported_claim": "This trial makes the general movement-break idea more plausible in an ordinary workplace over twelve weeks: a simple hourly bodyweight routine was feasible for many participants, increased light activity, reduced sedentary time and improved several measured metabolic markers. It is useful ecological support for continuing to develop a sitting-interruption protocol, but it should remain watch-level evidence for the exact hourly three-minute dose until stronger preregistered and replicated field evidence is available.",
      "unsupported_or_overstated_claims": [
        "Three minutes every hour is an established optimal dose.",
        "The intervention prevents or delays diabetes or cardiovascular disease.",
        "The results generalize unchanged to normal-weight, highly active, older, younger, shift-working, disabled, or non-office populations.",
        "The routine improves work productivity or cognitive performance; those outcomes were not established with the same objective causal standard.",
        "All seven prescribed breaks are necessary to obtain benefit.",
        "The specific six-exercise sequence is superior to walking or other feasible movement.",
        "The study proves long-term adherence or safety beyond twelve weeks."
      ],
      "limitations": [
        "The trial was not registered in a clinical-trial registry, which weakens confidence that outcomes and analyses were fixed prospectively.",
        "It was a single-city study across three workplaces with only 86 randomized participants and 79 completing follow-up.",
        "The sample was predominantly overweight, limiting generalizability to metabolically different populations.",
        "Participants and intervention facilitators could not be blinded and there was no attention-control condition.",
        "Adherence was monitored primarily by self-report, although workplace observations and accelerometry provided partial objective corroboration.",
        "Dietary intake was not measured, leaving a potential unmeasured contributor to metabolic change.",
        "Twelve weeks is too short to establish durable adherence or clinical endpoints such as diabetes incidence.",
        "The sample was not powered to establish reliable subgroup effects."
      ],
      "target_hack_ids": [
        "break-up-prolonged-sitting"
      ],
      "target_protocol_ids": [
        "brali:break-up-prolonged-sitting"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Evidence role: ecological support and dose-boundary watch. Do not promote the exact hourly three-minute routine to canonical Brali guidance from this trial alone. If later preregistered field trials converge, revisit frequency and implementation details."
    },
    {
      "schema_version": 1,
      "id": "ai-review-correction-friction-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/ai-review-correction-friction-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/ai-review-correction-friction-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-29",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Bias in the Loop: How Humans Evaluate AI-Generated Suggestions",
        "url": "https://hdsr.mitpress.mit.edu/pub/nrcn4h7d/release/2",
        "type": "primary-study",
        "doi": "10.1162/99608f92.0e98898d",
        "citation_text": "Beck J, Eckman S, Kern C, Kreuter F. Harvard Data Science Review. 2026;8(2).",
        "study_design": "Factorial randomized experiment with 2,784 U.S.-based Prolific participants reviewing ten corporate greenhouse-gas tables pre-annotated by an AI system. The experiment manipulated the correctness of the first three AI suggestions, whether rejecting an AI suggestion required entering a corrected value, and whether high accuracy received a performance bonus. A subset had completed an AI-attitudes survey one week earlier.",
        "population": "Adult U.S.-based crowdworkers on Prolific. Participants were generally experienced online task workers but were not selected for greenhouse-gas accounting expertise.",
        "intervention_or_exposure": "Human reviewers judged whether AI-extracted values were correct. In one randomized condition, flagging an AI error also required typing the corrected value, making rejection more effortful than acceptance.",
        "outcomes": [
          "Annotation accuracy",
          "Correction rate",
          "Undercorrection",
          "Overcorrection",
          "Annotation time"
        ]
      },
      "supported_claim": "When humans review AI-generated suggestions, the review interface itself can bias behavior. In this experiment, adding repair work to the act of rejecting an AI suggestion reduced correction activity and increased undercorrection. Brali can therefore justify a bounded workflow rule: make it cheap to flag or reject an AI output, and separate validation from repair when the repair burden would otherwise make acceptance the path of least resistance.",
      "unsupported_or_overstated_claims": [
        "Making correction easier will always increase overall accuracy.",
        "People who distrust AI are universally better reviewers.",
        "Performance bonuses cannot improve AI review in other settings.",
        "The same effect size applies to expert, medical, legal, financial, or safety-critical review.",
        "A human-in-the-loop label by itself guarantees reliable oversight.",
        "AI suggestions should be hidden from reviewers."
      ],
      "limitations": [
        "The experiment used crowdworkers rather than domain experts.",
        "The task was limited to ten greenhouse-gas reporting tables and one pre-annotation workflow.",
        "Several difficult items required domain knowledge that many annotators lacked.",
        "The randomized correction-burden manipulation changed correction behavior, but the regression analysis did not show a clear overall accuracy loss because overcorrections also decreased.",
        "AI attitudes predicted behavior observationally rather than through randomized manipulation.",
        "The study did not include a no-AI baseline and could not analyze item-order effects."
      ],
      "target_hack_ids": [
        "make-ai-rejection-cheap"
      ],
      "target_protocol_ids": [
        "brali:make-ai-rejection-cheap"
      ],
      "risk_flags": [],
      "notes": "Protocol direction: design AI review so Accept and Flag/Reject have comparable friction. Let the reviewer flag an output first; route correction, rewriting, or remediation into a separate step or queue when doing both at once would create asymmetric effort. Track undercorrection and overcorrection separately rather than treating 'human reviewed' as a quality guarantee."
    },
    {
      "schema_version": 1,
      "id": "ai-structured-intake-human-judgment-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/ai-structured-intake-human-judgment-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/ai-structured-intake-human-judgment-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-29",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Voice AI in Firms: A Natural Field Experiment on Automated Job Interviews",
        "url": "https://arxiv.org/html/2607.28222",
        "type": "primary-study",
        "doi": "10.2139/ssrn.5395709",
        "citation_text": "Jabarian B, Henkel L. Voice AI in Firms: A Natural Field Experiment on Automated Job Interviews. SSRN Working Paper, posted 2025; revised 2026.",
        "study_design": "Preregistered natural field experiment at PSG Global Solutions. Of 70,884 applications received during the experiment, 67,056 eligible applications were randomized to an AI interviewer, a human interviewer, or a choice condition. The direct causal comparison changed who conducted the information-collection interview while human recruiters evaluated applications and made every final hiring decision.",
        "population": "Applicants for 48 entry-level customer-service job postings across 41 client accounts, processed at 26 sites in 19 cities in the Philippines. Most applicants were aged 20-30 and had prior customer-service experience.",
        "intervention_or_exposure": "A voice AI agent conducted a structured but adaptive initial interview instead of a human recruiter. Human recruiters later evaluated interview information and standardized test results and retained final hiring authority.",
        "outcomes": [
          "Job offer rate",
          "Job starts",
          "Worker retention",
          "Measured worker productivity",
          "Interview structure and consistency",
          "Applicant experience",
          "Technical and refusal failures"
        ]
      },
      "supported_claim": "A defensible way to divide some repetitive high-volume workflows is to automate structured information collection while keeping consequential evaluation with a human. In this field experiment, that division improved several downstream hiring outcomes without a measured decline in worker productivity. Transcript evidence is consistent with greater standardization and comparability as a mechanism, but does not prove that mechanism independently. Brali should treat this as a task-allocation pattern to test, not as evidence that AI should make final hiring or other high-stakes decisions.",
      "unsupported_or_overstated_claims": [
        "AI interviewers are generally better than human interviewers.",
        "AI should make final hiring decisions.",
        "The result generalizes to specialized, relationship-heavy, tacit-knowledge, executive, clinical, legal, or other high-stakes work.",
        "Automating information collection removes discrimination or guarantees fairness.",
        "The same voice-AI system will produce the same results in other firms, languages, cultures, or labor markets.",
        "Human oversight automatically prevents automation bias.",
        "The controlled-variance mechanism is causally proven by the experiment."
      ],
      "limitations": [
        "The source is a working paper rather than a peer-reviewed journal article.",
        "The experiment was conducted with one recruitment-process outsourcing firm and entry-level customer-service hiring in the Philippines.",
        "Five percent of AI interviews ended because applicants were unwilling to continue with AI and seven percent experienced technical failure.",
        "The transcript-based mechanism analysis is associative even though interviewer assignment was randomized.",
        "The candidate-experience survey had a low response rate and may not represent all applicants.",
        "Applicants who were allowed to choose showed negative sorting into AI, limiting interpretation of the choice condition.",
        "Employment selection has legal, fairness, accessibility, and accountability requirements that this study does not resolve."
      ],
      "target_hack_ids": [
        "automate-intake-keep-judgment"
      ],
      "target_protocol_ids": [
        "brali:automate-intake-keep-judgment"
      ],
      "risk_flags": [],
      "notes": "Protocol direction: decompose a workflow into information collection and consequential evaluation. Consider AI for the repetitive collection stage only when inputs can be structured, auditable, and failure-handled; preserve an explicit human judgment stage, expose source material and uncertainty, and measure both process variance and downstream outcomes. Do not infer that the human stage is safe merely because it exists."
    },
    {
      "schema_version": 1,
      "id": "wakeful-rest-complex-learning-challenge-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/wakeful-rest-complex-learning-challenge-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/wakeful-rest-complex-learning-challenge-2026/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-08-28",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Wakeful rest and memory consolidation in an ecologically valid educational setting: no benefit over distractor tasks",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC12945945/",
        "type": "primary-study",
        "doi": "10.1007/s00426-026-02265-x",
        "citation_text": "Seban P, Šikl R, Prošek T, Urban K. Psychological Research. 2026;90:42.",
        "study_design": "Randomized four-group classroom study. University students read an expository text for eight minutes and were then assigned to eight minutes of wakeful rest, social-media use, math problems or a content-similar interference text. They completed an immediate test and a different but conceptually aligned test one week later.",
        "population": "161 university students tested in a supervised classroom setting.",
        "intervention_or_exposure": "Eight minutes of guided wakeful rest based on an autogenic-training procedure, compared with eight minutes of social-media use, math problems or reading a content-similar interference text after learning an expository passage.",
        "outcomes": [
          "Immediate factual recall",
          "Immediate conceptual understanding",
          "One-week factual retention",
          "One-week conceptual understanding"
        ]
      },
      "supported_claim": "The existing Brali wakeful-rest protocol must not promise better comprehension or long-term retention for complex educational reading. In this more ecologically valid randomized study, eight minutes of post-reading wakeful rest did not produce a consistent advantage over social media, math or similar-text reading on immediate or one-week factual and conceptual tests. Keep the protocol narrow: it is an optional declarative-memory tactic, not a general learning upgrade.",
      "unsupported_or_overstated_claims": [
        "This single study proves that wakeful rest never helps learning.",
        "Social-media use immediately after learning is harmless in every setting.",
        "All post-learning distractor tasks are equivalent.",
        "Autogenic training or guided rest is ineffective for every outcome.",
        "The null classroom result cancels the broader declarative-memory meta-analysis.",
        "The findings generalize unchanged beyond university students and the specific expository material tested."
      ],
      "limitations": [
        "The study used one expository learning task in a university-student population.",
        "The wakeful-rest condition used an autogenic-training implementation and therefore does not represent every possible quiet-rest procedure.",
        "Only one eight-minute post-learning duration was tested.",
        "The delayed test used different but conceptually aligned questions rather than identical items from the immediate test.",
        "A single null study cannot establish absence of a wakeful-rest effect across all complex learning tasks or populations."
      ],
      "target_hack_ids": [
        "avoid-list-interference-memory-retention"
      ],
      "target_protocol_ids": [
        "brali:avoid-list-interference-memory-retention"
      ],
      "risk_flags": [],
      "notes": "Evidence role: explicit boundary for the existing reviewed protocol and public copy. State that laboratory/meta-analytic declarative-memory effects do not yet translate cleanly to comprehension of complex educational material. Prefer recall-sensitive use cases and keep active learning methods primary."
    },
    {
      "schema_version": 1,
      "id": "mobile-internet-block-attention-boundary-2025",
      "canonical_url": "https://brali-lifeos.github.io/evidence/mobile-internet-block-attention-boundary-2025/",
      "json_url": "https://brali-lifeos.github.io/evidence/mobile-internet-block-attention-boundary-2025/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-27",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Blocking mobile internet on smartphones improves sustained attention, mental health, and subjective well-being",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC11834938/",
        "type": "primary-study",
        "doi": "10.1093/pnasnexus/pgaf017",
        "citation_text": "Castelo N, Kushlev K, Ward AF, Esterman M, Reiner PB. PNAS Nexus. 2025;4(2):pgaf017.",
        "study_design": "Preregistered month-long randomized controlled delayed-intervention trial. Participants were randomized to block all mobile internet on their smartphone during the first or second two-week phase. The Freedom app objectively tracked whether the block was active. Main analyses were intention-to-treat.",
        "population": "467 unique adult participants in the United States and Canada who agreed to attempt a two-week mobile-internet block. The sample was strongly selected toward people motivated and confident about reducing smartphone use; retention to all three surveys was 67 percent.",
        "intervention_or_exposure": "Block all Wi-Fi and mobile-data internet access on the smartphone for two weeks while preserving calls and text messages and allowing internet access from other devices such as computers or tablets.",
        "outcomes": [
          "Objectively measured sustained attention on the gradCPT",
          "Subjective wellbeing",
          "Self-reported mental-health composite",
          "Smartphone screen time",
          "Time-use mediators"
        ]
      },
      "supported_claim": "For adults who actively want to reduce constant phone connectivity, temporarily removing mobile-internet access from the phone is a defensible attention experiment. In this randomized trial, the intervention produced a small intention-to-treat improvement in objectively measured sustained attention while calls, texts and non-phone internet remained available. Brali can therefore test a bounded phone-offline window or period when the user's goal is fewer digital temptations and more sustained attention.",
      "unsupported_or_overstated_claims": [
        "Blocking mobile internet treats depression, anxiety, ADHD or any mental-health condition.",
        "Two weeks is an optimal or necessary duration.",
        "Everyone should use a dumb phone or eliminate smartphone internet.",
        "A fixed daily screen-time threshold is scientifically established by this study.",
        "The intervention will improve productivity, job performance or learning outcomes.",
        "The self-reported wellbeing effects are free from expectancy or demand effects.",
        "The study proves that every notification, app or form of internet access is harmful."
      ],
      "limitations": [
        "Only 119 of 467 participants who agreed to the intervention met the preregistered compliance threshold, so the full two-week block was difficult to sustain.",
        "The sample was strongly motivated to reduce smartphone use, which limits generalizability to people who do not share that goal.",
        "Only 67 percent completed all three surveys; completers had somewhat better baseline mental health and sustained attention than non-completers.",
        "The delayed-intervention design did not provide an active placebo control, so expectancy and demand effects may have affected self-report outcomes.",
        "The sustained-attention task provides a cognitive outcome, not direct evidence of improved work, study, safety or life performance.",
        "The study tested a broad mobile-internet block, not shorter work windows, app-specific limits, notification batching or other lower-friction variants."
      ],
      "target_hack_ids": [
        "phone-offline-window"
      ],
      "target_protocol_ids": [
        "brali:phone-offline-window"
      ],
      "risk_flags": [
        "health",
        "mental-health"
      ],
      "notes": "Protocol direction: when a user identifies phone connectivity as a recurring attention conflict, preserve essential communication but remove the internet channel from the phone for one bounded, user-chosen period. Measure actual phone use and completion/focus on the intended task. Start with a lower-friction window rather than prescribing the study's two-week block; escalate only if the user wants to test a stronger version. Stop or redesign the experiment if access is needed for safety, care, work authentication, navigation or other essential functions."
    },
    {
      "schema_version": 1,
      "id": "situational-digital-disconnection-procrastination-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/situational-digital-disconnection-procrastination-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/situational-digital-disconnection-procrastination-boundary-2026/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-08-27",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Digital disconnection as a self-regulatory strategy against procrastination",
        "url": "https://www.nature.com/articles/s41598-026-46218-1",
        "type": "primary-study",
        "doi": "10.1038/s41598-026-46218-1",
        "citation_text": "Klingelhoefer J, Gilbert A, Meier A. Scientific Reports. 2026;16:17133.",
        "study_design": "Preregistered two-week naturalistic experience-sampling study with up to five surveys per day. Multilevel models separated within-person from between-person associations, with exploratory lagged and disconnection-level analyses.",
        "population": "237 young adults aged 18 to 35 recruited from a German online panel and two university mailing lists; all used Android smartphones. The dataset contained 12,408 valid situational observations.",
        "intervention_or_exposure": "Self-initiated, temporary digital-disconnection behaviors in everyday life, spanning device, application, feature, interaction and message levels. The study observed behavior rather than assigning an intervention.",
        "outcomes": [
          "Momentary self-reported procrastination",
          "Goal conflict",
          "Digital-disconnection behavior",
          "Exploratory lagged procrastination"
        ]
      },
      "supported_claim": "Situational digital disconnection is a plausible low-friction self-regulation tactic to include in Brali's phone-offline protocol design. Within individuals, moments with more deliberate disconnection were associated with slightly less procrastination, and channel-focused strategies at the device, app or feature level showed the most promising pattern. This source adds ecological support for targeting access to a temptation before it is opened, but it does not establish a causal effect.",
      "unsupported_or_overstated_claims": [
        "Digital disconnection has been proven to reduce procrastination causally.",
        "Putting the phone away will improve productivity or task performance for every user.",
        "Device-level disconnection is always superior to app- or feature-level controls.",
        "A specific duration, frequency or disconnection tool is optimal.",
        "Digital media breaks are inherently procrastination or should always be suppressed.",
        "The findings generalize unchanged beyond young adult Android users."
      ],
      "limitations": [
        "The design was observational, so reverse causality and time-varying situational confounding remain possible.",
        "Digital disconnection and procrastination were both self-reported.",
        "Surveys asked participants to recall the previous two hours, creating possible recall delay and measurement error.",
        "The sample was limited to young adults and Android users recruited largely through German university-related channels.",
        "The lagged analyses and comparisons among disconnection levels were exploratory.",
        "The study did not measure downstream work performance, academic achievement or objective task completion."
      ],
      "target_hack_ids": [
        "phone-offline-window"
      ],
      "target_protocol_ids": [
        "brali:phone-offline-window"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Evidence role: support the protocol's situational and preventive design, not its causal effectiveness. Prefer removing access at the channel level (phone, app or feature) before a known temptation appears, and let the user choose the narrowest intervention that protects the intended task without unnecessarily blocking useful communication."
    },
    {
      "schema_version": 1,
      "id": "couple-conflict-channel-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/couple-conflict-channel-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/couple-conflict-channel-boundary-2026/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-08-26",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "A systematic review of technology-mediated conflict management in couples",
        "url": "https://doi.org/10.1177/02654075261420989",
        "type": "systematic-review",
        "doi": "10.1177/02654075261420989",
        "citation_text": "Daspe ME, Métellus S, Emond M, et al. Journal of Social and Personal Relationships. 2026;43(6):1836-1861.",
        "study_design": "Systematic review of 15 quantitative studies that directly compared technology-mediated and face-to-face conflict management in couples; eight were experimental studies and seven were surveys.",
        "population": "Couples in studies published from 2007 through 2024, mostly in the United States. Samples were generally convenience samples of relationally satisfied participants with low levels of negative conflict outcomes.",
        "intervention_or_exposure": "Technology-mediated conflict communication, including text, phone and video channels, compared with face-to-face conflict communication.",
        "outcomes": [
          "Conflict behaviors",
          "Emotional responses",
          "Perceived conflict resolution",
          "Perceived communication quality",
          "Relationship satisfaction"
        ]
      },
      "supported_claim": "For romantic-couple conflict, current evidence does not justify a universal channel rule. Across the 15 reviewed studies, technology-mediated and face-to-face conflict often produced similar outcomes, with mixed results and preliminary evidence that person and relationship characteristics may moderate effects. A Brali voice-switch protocol therefore needs an explicit close-relationship boundary instead of extending the broader spoken-disagreement finding as a universal rule.",
      "unsupported_or_overstated_claims": [
        "Texting is generally worse than face-to-face conflict discussion for couples.",
        "Face-to-face or voice conflict is universally better for couples.",
        "One communication channel reliably improves relationship satisfaction.",
        "Attachment style or self-esteem can currently be used as validated rules for selecting a conflict channel.",
        "The review establishes guidance for distressed, violent or otherwise high-risk couples.",
        "Older findings transfer unchanged to current messaging platforms and norms."
      ],
      "limitations": [
        "Only 15 quantitative studies met inclusion criteria.",
        "Seven of the 15 studies were grey literature.",
        "Several experimental studies had very small samples, including samples around two dozen participants.",
        "The studies were heterogeneous in design, communication channel and outcome measures.",
        "Convenience samples tended to be relationally satisfied and low-conflict.",
        "Laboratory technology-mediated interactions may be more synchronous and controlled than everyday messaging.",
        "The evidence spans technologies and norms from 2007 to 2024, limiting temporal uniformity."
      ],
      "target_hack_ids": [
        "move-disagreement-to-voice"
      ],
      "target_protocol_ids": [
        "brali:move-disagreement-to-voice"
      ],
      "risk_flags": [],
      "notes": "Editorial boundary: if Brali implements a spoken-over-written disagreement protocol, label it as a tactic for manageable disagreements where mutual understanding is the goal, not a default rule for couples. Relationship conflict should preserve choice of asynchronous text when it supports reflection, safety, accessibility or documentation."
    },
    {
      "schema_version": 1,
      "id": "human-first-ai-collaboration-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/human-first-ai-collaboration-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/human-first-ai-collaboration-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-26",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Relying on AI at work reduces self-efficacy, ownership, and meaning while active collaboration mitigates the effects",
        "url": "https://doi.org/10.1038/s41598-026-42312-6",
        "type": "preregistered-randomized-experiment-plus-correlational-follow-up",
        "doi": "10.1038/s41598-026-42312-6",
        "citation_text": "Lee EH, Yin Y, Jia N, Wakslak CJ. Scientific Reports. 2026;16:13583.",
        "study_design": "Preregistered randomized experiment with 269 analyzed participants completing occupation-specific writing tasks under no-AI, copy-and-paste AI, or human-first-then-AI-refinement conditions, followed by a correlational survey of 270 workers about broader real-world AI-use patterns.",
        "population": "Working adults across occupations including consulting, data analysis, HR, management and marketing; the follow-up survey sampled workers in the United States and United Kingdom.",
        "intervention_or_exposure": "No AI; passive use defined as directly using AI-generated content without modification; or active collaboration defined as producing an initial human draft and then using AI to review or refine it.",
        "outcomes": [
          "AI-independent self-efficacy",
          "Psychological ownership",
          "Work meaningfulness",
          "Task enjoyment",
          "Outcome satisfaction"
        ]
      },
      "supported_claim": "When preserving a worker's sense of authorship, competence and connection to the task matters, a meaningful human first pass before AI refinement is a reasonable workflow to test. In the randomized writing task, direct copy-and-paste use produced lower psychological ownership and meaningfulness than both no-AI and human-first collaboration, while the human-first condition was comparable to no-AI on those outcomes.",
      "unsupported_or_overstated_claims": [
        "Human-first AI use is universally more productive or more accurate.",
        "AI-first workflows are harmful in every task.",
        "Using AI causes a durable loss of skill or intelligence.",
        "The experiment establishes the best workflow for coding, analysis, research, design or decision making.",
        "The correlational follow-up proves causal effects outside writing tasks.",
        "Any amount of editing after AI generation is equivalent to the studied human-first condition.",
        "Avoiding AI is psychologically superior to active collaboration."
      ],
      "limitations": [
        "The causal experiment used short occupation-specific writing tasks rather than the full range of knowledge work.",
        "The passive condition was intentionally extreme: participants used AI-generated content directly without modification.",
        "The active condition tested one workflow only: human draft first, then AI refinement.",
        "The study did not instrument the external AI tool deeply enough to analyze prompt content and detailed interaction sequences.",
        "Baseline AI skill and confidence were not comprehensively modeled.",
        "The broader follow-up survey was correlational and cannot establish direction of causality."
      ],
      "target_hack_ids": [
        "human-first-ai-collaboration"
      ],
      "target_protocol_ids": [
        "brali:human-first-ai-collaboration"
      ],
      "risk_flags": [],
      "notes": "Protocol direction: for work where ownership and independent capability matter, write the task goal, constraints, judgment criteria and a meaningful first-pass artifact before asking AI to improve it. Keep the final accept/reject judgment human. Do not market this as a performance-optimization result."
    },
    {
      "schema_version": 1,
      "id": "spoken-vs-written-disagreement-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/spoken-vs-written-disagreement-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/spoken-vs-written-disagreement-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-26",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Spoken disagreement is more constructive than written disagreement",
        "url": "https://doi.org/10.1038/s41467-026-71669-5",
        "type": "series-of-randomized-experiments",
        "doi": "10.1038/s41467-026-71669-5",
        "citation_text": "Bevis B, Schroeder J, Yeomans M. Nature Communications. 2026;17:5792.",
        "study_design": "Series of randomized experiments comparing spoken and written disagreement across laboratory and field-like settings, with 1,576 conversation partners, 1,842 conversations, and 1,432 observers reported across the paper. Several studies separately manipulated conversation structure and examined conversational receptiveness.",
        "population": "Participants were primarily college students at American universities and often strangers. Conversation topics were controversial opinion disagreements; overall conflict levels were relatively low.",
        "intervention_or_exposure": "Disagreeing partners were assigned to spoken or written conversation, with some studies also varying interactivity, duration or synchronicity.",
        "outcomes": [
          "Perceived understanding",
          "Perceived conflict",
          "Social impressions",
          "Attitude alignment",
          "Conversational receptiveness"
        ]
      },
      "supported_claim": "For manageable disagreement settings similar to those studied, moving a discussion from writing to speaking is a reasonable tactic to test when the primary goal is mutual understanding: spoken conversations tended to produce greater perceived understanding and lower perceived conflict than written conversations.",
      "unsupported_or_overstated_claims": [
        "Speaking is always better than writing.",
        "Voice reliably resolves high-stakes workplace or family conflict.",
        "Speaking is safer in situations involving coercion, harassment, abuse or power imbalance.",
        "Written records should be avoided in legal, compliance, employment or accountability-sensitive contexts.",
        "Speaking guarantees persuasion, agreement or durable relationship improvement.",
        "The findings generalize unchanged to close couples, older populations or all cultures.",
        "One call duration, script or video format is optimal."
      ],
      "limitations": [
        "Overall conflict was relatively low.",
        "Participants were mostly American university students and often strangers.",
        "Topics were controversial opinions rather than typical workplace or close-relationship disputes.",
        "Some moderation analyses were underpowered.",
        "The research does not cover situations where written documentation, asynchronous reflection or distance is needed for safety or accountability."
      ],
      "target_hack_ids": [
        "move-disagreement-to-voice"
      ],
      "target_protocol_ids": [
        "brali:move-disagreement-to-voice"
      ],
      "risk_flags": [],
      "notes": "Protocol direction: detect repetitive or escalating text disagreement, separate the need for a durable record from the need for mutual understanding, move one bounded issue to voice only when safe and appropriate, then document decisions and unresolved points."
    },
    {
      "schema_version": 1,
      "id": "walking-divergent-thinking-boundary-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/walking-divergent-thinking-boundary-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/walking-divergent-thinking-boundary-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-26",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The impact of walking on creative thinking: A systematic review and meta-analysis",
        "url": "https://doi.org/10.1371/journal.pone.0347878",
        "type": "systematic-review-and-meta-analysis",
        "doi": "10.1371/journal.pone.0347878",
        "citation_text": "Thabane A, Arora V, Boilard J, et al. PLOS One. 2026;21(5):e0347878.",
        "study_design": "Systematic review and meta-analysis of 23 studies from 16 articles including 1,036 participants: 12 randomized experiments, nine non-randomized experiments and two observational studies. Randomized-only sensitivity analyses were reported.",
        "population": "Adults, with most participants drawn from post-secondary student samples.",
        "intervention_or_exposure": "Walking compared with non-walking control conditions, with divergent and convergent creative-thinking outcomes analyzed separately.",
        "outcomes": [
          "Divergent thinking",
          "Convergent thinking"
        ]
      },
      "supported_claim": "Walking is a defensible context change for the idea-generation phase of creative work. The meta-analysis found moderate-certainty evidence of a large positive effect on divergent thinking, including in randomized-only sensitivity analysis. The same evidence does not establish a benefit for convergent evaluation or selection.",
      "unsupported_or_overstated_claims": [
        "Walking improves every form of creativity.",
        "Walking improves reasoning, intelligence or decision quality.",
        "Walking helps select the best idea.",
        "One walking duration, intensity, speed, route or indoor/outdoor setting is optimal.",
        "The pooled effect size predicts the benefit for an individual user.",
        "Walking is required for good creative work."
      ],
      "limitations": [
        "Substantial heterogeneity remained unexplained.",
        "Most participants were post-secondary students.",
        "Divergent-thinking measures dominated the literature.",
        "Evidence for convergent thinking was rated very low certainty and was based on few studies.",
        "Several positive studies originated from a limited set of research groups and settings.",
        "The review does not identify an optimal duration or intensity."
      ],
      "target_hack_ids": [
        "walk-for-options"
      ],
      "target_protocol_ids": [
        "brali:walk-for-options"
      ],
      "risk_flags": [],
      "notes": "Protocol direction: separate creative work into generation and selection. Walk while producing alternatives; stop treating walking as the intervention once the task becomes comparison, verification and final selection."
    },
    {
      "schema_version": 1,
      "id": "communication-coaching-reflective-statements-boundary-2023",
      "canonical_url": "https://brali-lifeos.github.io/evidence/communication-coaching-reflective-statements-boundary-2023/",
      "json_url": "https://brali-lifeos.github.io/evidence/communication-coaching-reflective-statements-boundary-2023/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-08-22",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Effect of a Coaching Intervention to Improve Cardiologist Communication: A Randomized Clinical Trial",
        "url": "https://doi.org/10.1001/jamainternmed.2023.0629",
        "type": "randomized-clinician-coaching-trial",
        "doi": "10.1001/jamainternmed.2023.0629",
        "citation_text": "Pollak KI, Olsen MK, Yang H, et al. JAMA Internal Medicine. 2023;183(6):544-553.",
        "study_design": "Two-arm randomized clinical trial of forty cardiologists. The intervention combined three individual coaching sessions, feedback on recorded encounters and five communication behaviors. Different patient samples were recorded before and after the intervention; 230 post-intervention encounters were included in communication-behavior analyses.",
        "population": "Forty cardiologists, mostly White men in an academic medical center or affiliated clinics, and 240 post-intervention adult cardiology patients. Nearly one-third of cardiologists reported previous communication training.",
        "intervention_or_exposure": "A multicomponent WISER coaching package: sit and make eye contact, ask open-ended questions, make reflective statements or paraphrases, use empathic statements, and invite patient questions. Coaching included didactic instruction and tailored feedback on recorded encounters.",
        "outcomes": [
          "Coded use of reflective statements and open-ended questions",
          "Coded empathic statements and responses to empathic opportunities",
          "Use of the question-eliciting phrase 'What questions do you have?'",
          "Patient-reported communication, trust and empathy measures",
          "Cardiologist burnout and post-training process ratings"
        ]
      },
      "supported_claim": "A multicomponent individual coaching program changed some observed cardiologist communication behaviors, particularly empathic responses and eliciting patient questions. Reflective statements and open-ended questions were explicit trained components, but the intervention did not significantly increase either behavior relative to control. This source is therefore a boundary against claiming that training the Brali sequence by itself has demonstrated communication or patient-outcome effects.",
      "unsupported_or_overstated_claims": [
        "The coaching trial shows that paraphrasing or reflective statements alone improved.",
        "The Brali sequence of paraphrase, correction and one open question is equivalent to the five-component WISER intervention.",
        "Active listening improved patient understanding, trust, satisfaction or clinical outcomes in this trial.",
        "Open-ended questions increased because of the coaching intervention.",
        "A single conversational attempt is equivalent to three individual coaching sessions with tailored feedback.",
        "The findings generalize directly from cardiology encounters to ordinary work, family, conflict or digital conversations.",
        "Patient self-reports confirmed an effect; ceiling effects prevented the planned arm comparison.",
        "Positive clinician opinions about the coaching demonstrate objective effectiveness of each component."
      ],
      "limitations": [
        "Five communication skills were trained together, so any observed effect cannot be attributed to paraphrasing or open questions independently.",
        "There was no significant intervention-versus-control difference in reflective statements or open-ended questions.",
        "Patient-perceived communication and trust effects could not be assessed because scores had substantial ceiling effects.",
        "The sample included mostly White male cardiologists in an academic setting and many had prior communication training, limiting generalizability.",
        "Pre- and post-intervention patient samples were different rather than repeated measures on the same patients.",
        "The intervention required repeated one-to-one coaching and feedback, so it does not test a self-guided protocol used once.",
        "Clinical and longer-term outcomes were not established."
      ],
      "target_hack_ids": [
        "active-listening-exercises"
      ],
      "target_protocol_ids": [
        "brali:active-listening-exercises"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Editorial outcome: use this trial as an explicit challenge boundary. It prevents converting a broad communication-coaching result into evidence for Brali's specific sequence. Combined with the small paraphrasing experiment, the defensible status remains practical with visible source boundaries, not reviewed effectiveness."
    },
    {
      "schema_version": 1,
      "id": "distributed-practice-verbal-recall-boundary-2006",
      "canonical_url": "https://brali-lifeos.github.io/evidence/distributed-practice-verbal-recall-boundary-2006/",
      "json_url": "https://brali-lifeos.github.io/evidence/distributed-practice-verbal-recall-boundary-2006/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-08-22",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Distributed practice in verbal recall tasks: A review and quantitative synthesis",
        "url": "https://doi.org/10.1037/0033-2909.132.3.354",
        "type": "meta-analysis",
        "doi": "10.1037/0033-2909.132.3.354",
        "citation_text": "Cepeda NJ, Pashler H, Vul E, Wixted JT, Rohrer D. Psychological Bulletin. 2006;132(3):354-380.",
        "study_design": "Quantitative synthesis of distributed-practice experiments in verbal recall. From 427 reviewed articles, 317 experiments in 184 articles met the inclusion criteria, providing 958 accuracy values, 839 distributed-practice assessments and 169 effect sizes. Analyses compared massed with spaced study and examined how inter-study interval, retention interval, relearning procedures and fixed or expanding schedules affected final recall.",
        "population": "Participants ranged from children to older adults, although the evidence base was dominated by young adults. Clinical populations were excluded. Included tasks involved verbal material assessed by recall, such as paired associates, lists, spelling, foreign-language items and paragraph-like verbal information; recognition and frequency-judgment outcomes were excluded.",
        "intervention_or_exposure": "Repeated study episodes of the same verbal material separated by an inter-study interval, compared with massed presentations or with shorter versus longer spacing. Some studies used two episodes and others used fixed, expanding or contracting schedules across more episodes. Retention intervals ranged from very short delays to months or longer.",
        "outcomes": [
          "Final-test verbal recall after massed versus spaced study",
          "Final-test recall under shorter versus longer inter-study intervals",
          "Joint relationship between inter-study interval and intended retention interval",
          "Recall under fixed versus expanding study intervals",
          "Differences associated with task, age, experimental design and relearning procedure"
        ]
      },
      "supported_claim": "For verbal material measured by later recall, separating repeated study episodes by a meaningful interval generally supports better retention than concentrating the same material into massed study. The spacing associated with the best later recall tended to increase as the intended retention interval increased. For material that must be retained over months or years, the reviewed evidence supports distributing study across days or longer rather than completing all review in one sitting or one day. These findings justify the rewritten Brali action only within a verbal-recall boundary and without one universal schedule.",
      "unsupported_or_overstated_claims": [
        "One fixed daily, weekly or software-generated interval is optimal for all material and retention targets.",
        "Twenty minutes every day is superior to two hours in one sitting by a specific ratio.",
        "An 85 to 90 percent retrieval-success threshold identifies the correct next interval.",
        "Expanding intervals are reliably superior to fixed intervals; the review described the available evidence as limited, variable and insufficient for that conclusion.",
        "Longer spacing is always better; the review found that intervals can become too long and that the useful interval depends on the retention target.",
        "The source identifies testing, feedback or deliberate recall attempts as the causal mechanism of the spacing effect.",
        "The source establishes consolidation windows, neuroplasticity, reconsolidation or another biological mechanism for the public protocol.",
        "The result extends beyond later verbal recall to complex skills, problem solving, creative work, workplace performance or broader transfer outcomes.",
        "A spaced schedule guarantees durable memory for an individual learner.",
        "Children's long-term retention over months or years is established with the same confidence as the young-adult evidence.",
        "Daily streaks, adherence scores, AI coaching or app reminders are evidence-supported components of distributed practice."
      ],
      "limitations": [
        "The synthesis was restricted to verbal memory tasks measured by recall and deliberately excluded recognition, frequency judgments and the heterogeneous skill-learning literature.",
        "The evidence base was dominated by young adults; the authors reported very little middle-aged and older-adult evidence and insufficient long-term child data for confident generalization.",
        "Many studies did not report the variance data needed for effect-size calculation, so several analyses relied on accuracy differences and included fewer effect-size estimates.",
        "Published null findings may be underrepresented because of the file-drawer problem.",
        "Study materials, presentation schedules, retention intervals and experimental procedures varied substantially.",
        "Some historical studies confounded longer spacing with more relearning trials; the review examined this problem but could not remove every design limitation from the literature.",
        "Binning inter-study and retention intervals supported broad patterns but reduced the ability to recommend exact intervals.",
        "The useful interval depends jointly on spacing and the later retention target, and the authors stated that exact long-term optimization could not be specified with certainty.",
        "Evidence comparing expanding and fixed schedules was sparse and inconsistent, with large between-study variability.",
        "The synthesis addresses later recall, not broader comprehension, transfer, motivation, study adherence or real-world performance.",
        "The review was published in 2006 and should be rechecked against newer syntheses before adding more precise scheduling claims."
      ],
      "target_hack_ids": [
        "spaced-recall-coach"
      ],
      "target_protocol_ids": [
        "brali:spaced-recall-coach"
      ],
      "risk_flags": [],
      "notes": "Editorial outcome: promote only the fully rewritten protocol to reviewed. The public action now asks the user to choose verbal material, state a retention target, separate repeated study sessions, use longer gaps for longer retention targets, check final recall and adjust. It explicitly rejects one universal interval, an expanding-schedule guarantee and transfer to complex skills. Any future numerical schedule, success-threshold or mechanism claim requires a separate direct review."
    },
    {
      "schema_version": 1,
      "id": "empathic-paraphrasing-immediate-emotion-boundary-2012",
      "canonical_url": "https://brali-lifeos.github.io/evidence/empathic-paraphrasing-immediate-emotion-boundary-2012/",
      "json_url": "https://brali-lifeos.github.io/evidence/empathic-paraphrasing-immediate-emotion-boundary-2012/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-08-22",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Effects of Empathic Paraphrasing – Extrinsic Emotion Regulation in Social Conflict",
        "url": "https://doi.org/10.3389/fpsyg.2012.00482",
        "type": "within-participant-experiment",
        "doi": "10.3389/fpsyg.2012.00482",
        "citation_text": "Seehausen M, Kazzer P, Bajbouj M, Prehn K. Frontiers in Psychology. 2012;3:482.",
        "study_design": "Within-participant experiment in which one trained interviewer asked ten standardized open questions about a real-life social conflict and alternated empathic paraphrasing with silent note-taking. Immediate self-reported emotional valence, voice measures and psychophysiological measures were compared across conditions.",
        "population": "Twenty healthy native German-speaking adults, ten women, mean age 27 years, discussing an ongoing or recent conflict that did not involve physical or psychological violence. Some physiological and voice records were lost, leaving 16 or 17 participants for those analyses.",
        "intervention_or_exposure": "The interviewer briefly summarized the facts, perceived feelings and what seemed important to the participant, then asked whether that understanding was correct. The comparison condition was silent note-taking. The same highly trained interviewer delivered all sessions.",
        "outcomes": [
          "Immediate self-reported emotional valence after paraphrasing versus note-taking",
          "Voice intensity and speech measures following each condition",
          "Heart rate, skin conductance and blood-volume-pulse measures during each condition"
        ]
      },
      "supported_claim": "In this small nonclinical social-conflict experiment, a trained interviewer used a behavior closely matching Brali's core correction loop: summarize the speaker's facts, feelings and priorities, then ask whether the understanding is accurate. Participants reported less negative immediate emotion after paraphrasing than after silent note-taking. This supports retaining paraphrase plus explicit correction as a bounded practice behavior, with the studied outcome and setting stated precisely.",
      "unsupported_or_overstated_claims": [
        "Paraphrasing improves understanding accuracy or prevents misunderstandings; the study asked for accuracy confirmation but did not measure comprehension accuracy as an outcome.",
        "The protocol resolves conflict, improves relationships, increases trust or produces better decisions.",
        "One open question after a paraphrase has an independent effect; open questions were part of the interview procedure and were not isolated.",
        "The emotional effect is durable beyond the immediate interview period.",
        "Any untrained listener will reproduce the effect in ordinary work, family or online conversations.",
        "Paraphrasing is calming in a simple physiological sense; several autonomic measures showed higher activation during paraphrasing.",
        "A wrong paraphrase is always beneficial or harmless; the study used one experienced mediator and excluded violent conflicts.",
        "The study validates counting summaries, eye contact, nods, fixed timing or a universal conversational script."
      ],
      "limitations": [
        "Only twenty participants contributed self-report data, making results vulnerable to individual variation and unsuitable for broad population estimates.",
        "One female interviewer with approximately 190 hours of conflict-resolution training delivered every paraphrase, so listener skill and delivery cannot be separated from the technique.",
        "The control condition was silent note-taking rather than another spoken response; differences may partly reflect receiving a verbal response rather than paraphrasing specifically.",
        "Participants may have interpreted note-taking as judgment despite the study explanation, potentially biasing the comparison.",
        "Only immediate reactions were measured; the authors explicitly described longer-term emotion resolution as speculative.",
        "Conflicts involving physical or psychological violence were excluded, so the result must not be applied as ordinary advice in unsafe or abusive situations.",
        "The study assessed emotional valence and arousal, not whether the listener understood more accurately or whether the conflict outcome improved."
      ],
      "target_hack_ids": [
        "active-listening-exercises"
      ],
      "target_protocol_ids": [
        "brali:active-listening-exercises"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Editorial outcome: keep the public protocol practical rather than promote it to reviewed on this study alone. The source supports the exact paraphrase-plus-accuracy-check behavior in a narrow setting, but the public protocol targets understanding and correction, whereas the study's measured outcome was immediate emotional valence. Preserve the correction invitation, avoid mind-reading and emotion-as-fact wording, and keep unsafe conflict contexts outside routine use."
    },
    {
      "schema_version": 1,
      "id": "microbreak-wellbeing-performance-boundary-2022",
      "canonical_url": "https://brali-lifeos.github.io/evidence/microbreak-wellbeing-performance-boundary-2022/",
      "json_url": "https://brali-lifeos.github.io/evidence/microbreak-wellbeing-performance-boundary-2022/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-08-22",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Give me a break! A systematic review and meta-analysis on the efficacy of micro-breaks for increasing well-being and performance",
        "url": "https://pmc.ncbi.nlm.nih.gov/articles/PMC9432722/",
        "type": "systematic-review-meta-analysis",
        "doi": "10.1371/journal.pone.0272460",
        "citation_text": "Albulescu P, Macsinga I, Rusu A, Sulea C, Bodnaru A, Tulbure BT. PLOS ONE. 2022;17(8):e0272460.",
        "study_design": "Systematic review and random-effects meta-analysis of experimental and quasi-experimental studies comparing work-task micro-breaks of 10 minutes or less with control conditions. Nineteen publications contributed 22 independent study samples with 2,335 participants; most studies involved students or employees and most were conducted in laboratory or workplace settings.",
        "population": "Healthy students and employees engaged in depleting work or work-like tasks. The included samples varied in setting, task demands and break activity; the review did not establish one universal workplace or individual profile.",
        "intervention_or_exposure": "Short task interruptions of no more than 10 minutes, with varied break content and duration, compared with continued work or another break/control condition. This is related to the break step in Brali's Pomodoro-style protocol but does not test the full 25-minute-work/5-minute-break sequence.",
        "outcomes": [
          "Vigor after a micro-break",
          "Fatigue after a micro-break",
          "Task performance after a micro-break"
        ]
      },
      "supported_claim": "Across the included studies, micro-breaks of up to 10 minutes were associated on average with small improvements in vigor and reduced fatigue. The pooled overall performance effect was not statistically significant, and performance effects varied by task demands and break duration. This supports retaining a short, adjustable break as a bounded recovery practice without presenting it as a reliable performance booster.",
      "unsupported_or_overstated_claims": [
        "A 5-minute break is an evidence-based optimal duration.",
        "Twenty-five minutes of work followed by five minutes of rest is superior to other focus schedules.",
        "Pomodoro reliably improves productivity or task performance.",
        "Four 25-minute blocks followed by a longer break is validated by this meta-analysis.",
        "Every short break reduces fatigue or increases vigor for every individual.",
        "Scrolling, walking, standing, water, stretching or any other specific break activity is independently proven superior by this pooled result.",
        "The review establishes a universal break schedule for cognitively demanding work.",
        "The pooled effects justify treatment, burnout-prevention or clinical claims."
      ],
      "limitations": [
        "The review combined heterogeneous break activities, durations, tasks and settings rather than testing a single Pomodoro-style intervention.",
        "Only four of twenty-two included study samples were rated low risk of bias across the assessed domains.",
        "The overall pooled performance effect was not statistically significant and performance heterogeneity was substantial.",
        "Performance effects differed by task type; the authors reported significant effects only for less cognitively demanding tasks in subgroup analyses.",
        "Longer break duration was associated with better performance in meta-regression, which argues against treating one short duration as universally optimal.",
        "Most outcomes were immediate post-break measures; the review does not establish long-term productivity or health effects.",
        "The source defines micro-breaks as no more than ten minutes and does not validate the work interval preceding the break."
      ],
      "target_hack_ids": [
        "25-minute-pomodoro-focus-sprints"
      ],
      "target_protocol_ids": [
        "brali:25-minute-pomodoro-focus-sprints"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Editorial outcome: keep `25-minute-pomodoro-focus-sprints` practical. Its current wording already treats 25 minutes as a conventional starting setting and the break as an adjustable task boundary. Do not add a general performance claim or present five minutes as an optimal recovery dose."
    },
    {
      "schema_version": 1,
      "id": "procrastination-treatment-exact-protocol-boundary-2018",
      "canonical_url": "https://brali-lifeos.github.io/evidence/procrastination-treatment-exact-protocol-boundary-2018/",
      "json_url": "https://brali-lifeos.github.io/evidence/procrastination-treatment-exact-protocol-boundary-2018/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-08-22",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Targeting Procrastination Using Psychological Treatments: A Systematic Review and Meta-Analysis",
        "url": "https://doi.org/10.3389/fpsyg.2018.01588",
        "type": "systematic-review-meta-analysis",
        "doi": "10.3389/fpsyg.2018.01588",
        "citation_text": "Rozental A, Bennett S, Forsström D, Ebert DD, Shafran R, Andersson G, Carlbring P. Frontiers in Psychology. 2018;9:1588.",
        "study_design": "Pre-registered systematic review and random-effects meta-analysis of randomized controlled trials comparing psychological interventions specifically targeting procrastination with inactive comparators. Twelve studies, 21 comparisons and 718 participants were included in the quantitative synthesis; the primary endpoint was self-reported procrastination at post-treatment.",
        "population": "Self-referred participants drawn from general-population, high-school and undergraduate samples in the included studies. Individual studies were generally small, with sample sizes reported by the review as ranging from 6 to 50 participants.",
        "intervention_or_exposure": "A heterogeneous set of psychological treatments specifically targeting procrastination, including cognitive-behavioral, emotion-focused, self-control, goal-setting and other approaches delivered individually, in groups, by class, online or through self-help over durations ranging from brief sessions to ten weeks.",
        "outcomes": [
          "Self-reported procrastination at post-treatment compared with an inactive control",
          "Secondary self-report outcomes where available, which were too heterogeneous for a pooled estimate"
        ]
      },
      "supported_claim": "Across the reviewed randomized comparisons, psychological treatments targeting procrastination had a small average post-treatment benefit on self-reported procrastination relative to inactive controls, but effects varied substantially between studies and the evidence base had important quality limitations. This supports treating task initiation as a legitimate intervention target, not claiming that one brief Brali routine or one component has been validated.",
      "unsupported_or_overstated_claims": [
        "A five-minute start is an evidence-based or optimal duration for overcoming procrastination.",
        "Completing one visible first action by itself reduces procrastination, improves productivity or reliably creates momentum.",
        "Implementation intentions, sub-goals, goal-setting, time management or behavioral activation were isolated by this review as the active causal component.",
        "The broader implementation-intention literature on goal attainment directly proves effectiveness for procrastination-specific task initiation.",
        "The cognitive-behavioral subgroup estimate validates this Brali protocol or any other single self-help exercise.",
        "A short timer is equivalent to the multi-component psychological treatments included in the review.",
        "The intervention effect generalizes equally to every person, task type, workplace, educational setting or clinical presentation.",
        "Beginning briefly guarantees continuation, completion, lower distress or better performance."
      ],
      "limitations": [
        "The pooled overall effect was small and showed significant between-study heterogeneity, so the average should not be treated as one stable effect across interventions or populations.",
        "The review identified some risk of bias in every included study; blinding-related domains were frequently rated high risk.",
        "Primary procrastination measures varied, and several assessed academic rather than general procrastination.",
        "The included studies used different interventions, delivery formats, durations and participant populations, preventing component-level attribution.",
        "The review analyzed post-treatment outcomes because too few studies reported follow-up data, limiting conclusions about durability.",
        "Secondary outcomes were too heterogeneous to aggregate.",
        "The stronger cognitive-behavioral subgroup estimate depended on excluding one small outlying study and represented only three of the four cognitive-behavioral studies.",
        "The source did not test the exact five-minute interval, the exact Brali action sequence or the protocol's observable-review checkpoint."
      ],
      "target_hack_ids": [
        "stop-procrastination-fast-5-minute-reset"
      ],
      "target_protocol_ids": [
        "brali:stop-procrastination-fast-5-minute-reset"
      ],
      "risk_flags": [
        "mental-health"
      ],
      "notes": "Editorial outcome: keep the rewritten protocol practical and bounded, but do not promote it to reviewed from this source. Preserve explicit wording that five minutes is a convenient reversible default rather than an optimal dose, that stopping to clarify or remove a blocker is valid, and that the routine is not a treatment for clinically significant procrastination or related distress. Future promotion would require evidence that directly evaluates a sufficiently similar low-intensity first-action protocol and its outcome boundary."
    },
    {
      "schema_version": 1,
      "id": "testing-versus-restudy-retention-boundary-2014",
      "canonical_url": "https://brali-lifeos.github.io/evidence/testing-versus-restudy-retention-boundary-2014/",
      "json_url": "https://brali-lifeos.github.io/evidence/testing-versus-restudy-retention-boundary-2014/index.json",
      "decision": "support-existing",
      "decision_label": "Supports existing guidance",
      "reviewed_at": "2026-08-22",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The effect of testing versus restudy on retention: a meta-analytic review of the testing effect",
        "url": "https://doi.org/10.1037/a0037559",
        "type": "meta-analysis",
        "doi": "10.1037/a0037559",
        "citation_text": "Rowland CA. Psychological Bulletin. 2014;140(6):1432-1463.",
        "study_design": "Random-effects meta-analysis of testing-effect studies using a conceptually common contrast between an initial memory test and an equivalent-duration restudy control, followed by a later retention assessment. The review also analyzed methodological and theoretical moderators rather than assuming one invariant testing effect.",
        "population": "Participants from the experimental testing-effect literature included by the review, with many studies using college samples and varied learning materials, testing formats and retention intervals. The source does not justify treating the pooled estimate as one person-level effect or as a universal educational dose.",
        "intervention_or_exposure": "An initial retrieval test over previously studied information compared with additional restudy of that information, followed by a later memory assessment. This maps to Brali's retrieve-before-looking step more directly than to its separate application-practice step.",
        "outcomes": [
          "Later retention of previously learned information after testing versus restudy",
          "Moderator patterns across initial test format, final test format, retention interval, feedback, design and other study characteristics"
        ]
      },
      "supported_claim": "Across the included testing-versus-restudy literature, retrieval testing produced better later retention on average than equivalent-duration restudy. The pooled random-effects estimate was positive, but effects varied substantially across studies and testing conditions. This supports Brali's bounded recommendation to retrieve information from memory, check the answer, and use that loop when retention is the target.",
      "unsupported_or_overstated_claims": [
        "Active recall improves every kind of learning outcome.",
        "Better retention automatically produces better application, transfer, problem solving or procedural performance.",
        "One pooled testing-effect estimate describes one invariant mechanism across all retrieval-practice designs.",
        "A specific number of questions, repetitions, minutes or percentage target is evidence-based from this meta-analysis.",
        "Retrieval practice always outperforms restudy for every learner, item, delay and test format.",
        "Recognition testing, cued recall and free recall produce interchangeable effects.",
        "The meta-analysis establishes a single neurological or cognitive mechanism that Brali should present as the explanation for why retrieval works.",
        "A memory benefit demonstrated against restudy validates Brali's protocol for skills whose real criterion is performance or application."
      ],
      "limitations": [
        "Between-study heterogeneity in the primary analysis was very high, so the pooled mean does not describe one uniform effect across designs or contexts.",
        "The review deliberately restricted the quantitative synthesis to testing-versus-equivalent-restudy contrasts and does not represent every retrieval-practice paradigm.",
        "Testing benefits varied with methodological factors including initial test format, retention interval and feedback.",
        "Many contributing studies used controlled experimental learning tasks and college samples, which limits direct extrapolation to every real-world learning context.",
        "The review's main outcome is retention; it does not establish that remembered information can be applied correctly in a novel task or procedure.",
        "Theoretical findings did not provide one complete cohesive mechanism for all testing phenomena, so Brali should avoid mechanism decoration.",
        "A positive average effect does not imply that every individual effect was positive; the reviewed distribution included negative and null estimates."
      ],
      "target_hack_ids": [
        "active-recall-test-yourself"
      ],
      "target_protocol_ids": [
        "brali:active-recall-test-yourself"
      ],
      "risk_flags": [],
      "notes": "Editorial outcome: support the existing reviewed retention boundary without broadening it. Keep the public protocol's explicit split between retention and application, its retrieval-plus-checking loop, and its refusal to prescribe a universal duration or target. The existing 2026 challenge evidence remains important because the Rowland review does not erase distinctions between direct retention, later learning and actual rule or skill application."
    },
    {
      "schema_version": 1,
      "id": "writing-feedback-level-match-2024",
      "canonical_url": "https://brali-lifeos.github.io/evidence/writing-feedback-level-match-2024/",
      "json_url": "https://brali-lifeos.github.io/evidence/writing-feedback-level-match-2024/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-21",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "How effective is feedback for L1, L2, and FL learners’ writing? A meta-analysis",
        "url": "https://doi.org/10.1016/j.learninstruc.2024.101961",
        "type": "meta-analysis",
        "doi": "10.1016/j.learninstruc.2024.101961",
        "citation_text": "Scherer S, Graham S, Busse V. Learning and Instruction. 2024;93:101961.",
        "study_design": "Meta-analysis of 95 papers meeting experimental or quasi-experimental inclusion criteria, comprising 200 treatment/control comparisons and 253 effect sizes for surface- and deep-level writing outcomes.",
        "population": "Secondary-school and university or college students represented in L1, L2, and foreign-language writing studies across the included literature.",
        "intervention_or_exposure": "Writing feedback targeting surface-level features, deep-level features, or both, delivered through sources including instructors, peers, algorithm-based systems, and self-feedback depending on the included comparison.",
        "outcomes": [
          "Surface-level writing outcomes such as grammar and spelling",
          "Deep-level writing outcomes such as writing quality, ideation, organization, genre elements, composition length, and vocabulary"
        ]
      },
      "supported_claim": "In the reviewed secondary-school and university writing literature, feedback effects differed by the level of writing outcome: surface-level feedback improved surface-level outcomes, deep-level feedback improved deep-level outcomes, and combined surface-and-deep feedback could improve both. This supports matching a feedback request to the revision outcome rather than treating all feedback as interchangeable.",
      "unsupported_or_overstated_claims": [
        "More feedback is always better for writing.",
        "Deep feedback should always precede surface feedback.",
        "Peer, instructor, algorithm-based, or self-feedback is universally the best source.",
        "The meta-analytic effects generalize unchanged to professional, technical, marketing, fiction, or other creative writing outside the reviewed educational contexts.",
        "A specific number of feedback rounds, timing interval, or review duration is optimal.",
        "Receiving feedback guarantees improvement for an individual writer."
      ],
      "limitations": [
        "The evidence concerns educational writing by secondary-school and university or college learners rather than all writing contexts.",
        "Effects varied across L1, L2, and foreign-language learners, so a single population-wide estimate can conceal important differences.",
        "The included papers used experimental or quasi-experimental designs but varied in feedback treatment, source, duration, outcome measurement, and instructional context.",
        "Only about one-third of the treatment/control comparisons examined maintenance effects, limiting conclusions about durability.",
        "Many studies did not report enough learner-level variables, and evidence was comparatively scarce for some learner groups and treatment-outcome combinations.",
        "Meta-analytic categories for surface, deep, combined, and feedback source are useful abstractions but do not establish one universally optimal revision workflow."
      ],
      "target_hack_ids": [
        "ask-a-peer-to-proofread"
      ],
      "target_protocol_ids": [
        "brali:ask-a-peer-to-proofread"
      ],
      "risk_flags": [],
      "notes": "Replace the inherited generic peer-proofreading page with a conservative educational-writing workflow: name the current revision level, request feedback at that level, revise against that target, and use a separate pass when another level also matters. Keep the evidence boundary visible and do not convert subgroup findings into a universal ranking of feedback agents."
    },
    {
      "schema_version": 1,
      "id": "apology-repair-2021",
      "canonical_url": "https://brali-lifeos.github.io/evidence/apology-repair-2021/",
      "json_url": "https://brali-lifeos.github.io/evidence/apology-repair-2021/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-19",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Experimental evidence that apologies promote forgiveness by communicating relationship value",
        "url": "https://pubmed.ncbi.nlm.nih.gov/34162912/",
        "type": "randomized-experiment",
        "doi": "10.1038/s41598-021-92373-y",
        "citation_text": "Forster DE, Billingsley J, Burnette JL, et al. Scientific Reports. 2021;11:13107.",
        "study_design": "Controlled 2 × 2 double-randomization experiments manipulating apology and relationship-value information to test causal pathways to forgiveness.",
        "population": "Human participants in controlled interpersonal-transgression scenarios used in the published experiments.",
        "intervention_or_exposure": "Receiving an apology versus no apology, crossed with manipulated relationship-value information.",
        "outcomes": [
          "Forgiveness",
          "Perceived relationship value"
        ]
      },
      "supported_claim": "In the controlled experiments, apologies promoted forgiveness and operated through a causal pathway shared with perceived relationship value. This supports apology as one potentially useful repair signal after a workable interpersonal transgression.",
      "unsupported_or_overstated_claims": [
        "An apology guarantees forgiveness, reconciliation, or restored trust.",
        "The harmed person should forgive, accept the apology, or resume the relationship.",
        "One fixed apology script, cooling-off duration, or delivery deadline is scientifically established.",
        "Apology is an appropriate repair strategy for abuse, coercion, threats, stalking, or other unsafe situations.",
        "The study proves that apology alone resolves complex or repeated relationship conflict."
      ],
      "limitations": [
        "Controlled experimental transgression paradigms do not capture every feature of real-world repeated or high-stakes conflict.",
        "Forgiveness is only one outcome and is not equivalent to trust, reconciliation, safety, or durable behavior change.",
        "The effect of apology depends on context, relationship value and the nature of the transgression.",
        "The study does not validate the specific wording or step sequence used in Brali's public repair protocol."
      ],
      "target_hack_ids": [
        "relationship-repair-coach"
      ],
      "target_protocol_ids": [
        "brali:relationship-repair-coach"
      ],
      "risk_flags": [],
      "notes": "Use apology as a bounded repair action: own the behavior, avoid conditional language, offer a concrete corrective step, invite the other person's perspective, and do not demand forgiveness. Safety and boundaries take precedence over repair."
    },
    {
      "schema_version": 1,
      "id": "sleep-opportunity-extension-2021",
      "canonical_url": "https://brali-lifeos.github.io/evidence/sleep-opportunity-extension-2021/",
      "json_url": "https://brali-lifeos.github.io/evidence/sleep-opportunity-extension-2021/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-19",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Behavioral interventions to extend sleep duration: A systematic review and meta-analysis",
        "url": "https://pubmed.ncbi.nlm.nih.gov/34507028/",
        "type": "meta-analysis",
        "doi": "10.1016/j.smrv.2021.101532",
        "citation_text": "Baron KG, Duffecy J, Reutrakul S, et al. Sleep Medicine Reviews. 2021;60:101532.",
        "study_design": "Systematic review and meta-analysis of 42 behavioral sleep-extension studies, including controlled two-arm and one-arm designs.",
        "population": "Participants aged 12 years and older across adult, college-student, adolescent and other samples in 14 countries.",
        "intervention_or_exposure": "Behavioral interventions intended to increase sleep duration, ranging from direct sleep-schedule changes to coaching and educational approaches.",
        "outcomes": [
          "Sleep duration"
        ]
      },
      "supported_claim": "Behavioral interventions designed to extend sleep increased sleep duration on average compared with control or baseline. Direct interventions that specified a sleep schedule tended to have larger effects, but results were highly heterogeneous.",
      "unsupported_or_overstated_claims": [
        "A self-tracking experiment can identify one scientifically optimal personal sleep duration.",
        "A fixed bedtime or wake-time tolerance is proven to improve health or cognition for everyone.",
        "A 21-day experiment, 30-minute duration bins, or a particular alertness score is evidence-based.",
        "Sleep extension by itself is a treatment for chronic insomnia or another sleep disorder.",
        "Consumer sleep-stage estimates are required to decide how much sleep a person needs."
      ],
      "limitations": [
        "Statistical heterogeneity was very high across both two-arm and one-arm studies.",
        "Populations, intervention components, duration and measurement methods varied substantially.",
        "The primary outcome was sleep duration; the review does not establish a single personal optimum for daytime performance or health.",
        "The findings do not replace clinical assessment for persistent insomnia, excessive sleepiness, suspected sleep disorders, or other medical concerns."
      ],
      "target_hack_ids": [
        "ideal-sleep-hours-finder"
      ],
      "target_protocol_ids": [
        "brali:ideal-sleep-hours-finder"
      ],
      "risk_flags": [
        "health"
      ],
      "notes": "Promote only as a conservative protocol for protecting adequate sleep opportunity and observing trends. Pair public adult-duration wording with the AASM/SRS consensus; do not revive the inherited optimization experiment or infer causality from the separate sleep-variability observational review."
    },
    {
      "schema_version": 1,
      "id": "movement-breaks-cognition-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/movement-breaks-cognition-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/movement-breaks-cognition-2026/index.json",
      "decision": "watch",
      "decision_label": "Watch, do not prescribe",
      "reviewed_at": "2026-08-18",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Acute cognitive effects of interruptions to prolonged sitting with brief standing or physical activity breaks: a systematic review and three-level meta-analysis",
        "url": "https://link.springer.com/article/10.1186/s12966-026-01953-6",
        "type": "meta-analysis",
        "doi": "10.1186/s12966-026-01953-6",
        "citation_text": "Zhuang M, Yin M, Liu Z, et al. International Journal of Behavioral Nutrition and Physical Activity. Published 13 July 2026.",
        "study_design": "Systematic review and three-level random-effects meta-analysis of 21 randomized crossover trials, with cluster-robust variance estimation and GRADE certainty assessment.",
        "population": "433 participants across the included randomized crossover trials; exploratory subgroup findings varied by age and weight status.",
        "intervention_or_exposure": "Brief standing or physical-activity interruptions during periods of prolonged sitting, compared with uninterrupted prolonged sitting.",
        "outcomes": [
          "Executive function",
          "Memory",
          "Global cognition",
          "Information processing",
          "Attention"
        ]
      },
      "supported_claim": "Across the included trials, interruptions to prolonged sitting were associated with small acute improvements in executive function and memory on average. Evidence for global cognition, information processing, and attention was insufficient or inconsistent, and certainty was low for most outcomes.",
      "unsupported_or_overstated_claims": [
        "A specific break frequency is proven to optimize cognition.",
        "A specific break duration or intensity should be prescribed for everyone.",
        "Standing or movement breaks reliably improve attention.",
        "The exploratory moderator findings establish that one protocol is superior.",
        "Brief breaks are a treatment for cognitive impairment."
      ],
      "limitations": [
        "Only 21 randomized crossover trials with 433 participants were included.",
        "GRADE certainty was low for executive function, memory, information processing, and attention and very low for global cognition.",
        "No clear improvement was found for global cognition, information processing, or attention.",
        "Moderator and protocol-parameter findings were exploratory and the authors explicitly state they are insufficient for definitive prescriptive recommendations.",
        "The outcomes were acute cognitive effects, so the source does not establish long-term cognitive benefit from a break routine."
      ],
      "target_hack_ids": [],
      "target_protocol_ids": [],
      "risk_flags": [
        "health"
      ],
      "notes": "Keep as watch. Brali may recommend breaking up prolonged sitting for broader practical/health reasons only when supported by appropriate sources, but this paper does not justify a magic cognitive break interval or a cognition-optimization protocol."
    },
    {
      "schema_version": 1,
      "id": "retrieval-procedural-application-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/retrieval-procedural-application-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/retrieval-procedural-application-2026/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-08-18",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Mixed effects of retrieval practice on the acquisition of procedural knowledge in verb spelling education",
        "url": "https://www.sciencedirect.com/science/article/pii/S0959475226000629",
        "type": "primary-study",
        "doi": "10.1016/j.learninstruc.2026.102377",
        "citation_text": "Ophuis-Cox FHA, Catrysse L, Camp G. Learning and Instruction. 2026;104:102377.",
        "study_design": "Cluster-randomized classroom experiment comparing worked examples with worked examples plus retrieval practice, with outcomes tested after one week.",
        "population": "105 fourth-grade pupils from regular primary schools in the Netherlands, mean age 9.4 years.",
        "intervention_or_exposure": "Dutch verb-spelling instruction using worked examples alone or worked examples combined with retrieval practice across practice sessions.",
        "outcomes": [
          "Long-term retention of verb-spelling rules",
          "Correct application of verb-spelling rules"
        ]
      },
      "supported_claim": "Adding retrieval practice to worked examples improved long-term retention of the spelling rules in this classroom study, but the combined group did not outperform worked examples alone on correct rule application after one week.",
      "unsupported_or_overstated_claims": [
        "Remembering a rule better means a learner can apply it better.",
        "Retrieval practice alone is sufficient practice for procedural skill application.",
        "This one primary-school spelling study establishes a general failure of retrieval practice for transfer or skill learning.",
        "The specific classroom protocol should be copied as a universal learning recipe."
      ],
      "limitations": [
        "The sample was 105 fourth-grade pupils learning Dutch verb spelling in a small number of intact classes and schools.",
        "The comparison tested retrieval practice added to worked examples, not retrieval practice in isolation.",
        "The target was one form of procedural knowledge and the study does not establish identical boundaries in other domains or age groups.",
        "The result supports separating retention from application outcomes, not concluding that retrieval practice is ineffective for complex learning generally."
      ],
      "target_hack_ids": [
        "active-recall-test-yourself"
      ],
      "target_protocol_ids": [
        "brali:active-recall-test-yourself"
      ],
      "risk_flags": [],
      "notes": "Brali should keep active recall as a retention-oriented method but explicitly add application practice when the real goal is using a rule or procedure."
    },
    {
      "schema_version": 1,
      "id": "sleep-variability-cognition-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/sleep-variability-cognition-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/sleep-variability-cognition-2026/index.json",
      "decision": "watch",
      "decision_label": "Watch, do not prescribe",
      "reviewed_at": "2026-08-18",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "The association between sleep variability and cognitive function: A systematic review and meta analysis",
        "url": "https://www.sciencedirect.com/science/article/abs/pii/S0165032725019238",
        "type": "meta-analysis",
        "doi": "10.1016/j.jad.2025.120481",
        "citation_text": "Xia Y, He S, Sun Z, Wang F. Journal of Affective Disorders. 2026;393(Pt B):120481.",
        "study_design": "Systematic review and three-level meta-analysis of 20 published articles, 103 effect sizes, total N=15,828.",
        "population": "Human participants across age groups represented in the included observational literature; the association was stronger in older adults than in adolescents and middle-aged adults.",
        "intervention_or_exposure": "Night-to-night sleep variability or irregularity across sleep measures.",
        "outcomes": [
          "Cognitive performance"
        ]
      },
      "supported_claim": "Across the included literature, greater sleep variability was associated with poorer cognitive performance on average, with a small pooled correlation and age moderating the association.",
      "unsupported_or_overstated_claims": [
        "Sleep irregularity has been shown to cause cognitive decline.",
        "Making sleep timing more regular has been shown by this meta-analysis to improve cognition.",
        "A fixed bedtime or wake-time tolerance can be prescribed from this evidence.",
        "The association justifies treatment or dementia-prevention claims for an individual."
      ],
      "limitations": [
        "The reviewed evidence is primarily association evidence and does not establish the direction of causality.",
        "Sleep variability and cognition were operationalized in different ways across studies.",
        "The pooled association was small (r=-0.12 in the authors' analysis).",
        "Age moderated the association, so a single population-wide claim would hide meaningful differences.",
        "The publisher abstract and bibliographic record were directly reviewed for this Brali decision; a causal behavior-change protocol is deliberately not proposed from this source alone."
      ],
      "target_hack_ids": [],
      "target_protocol_ids": [],
      "risk_flags": [
        "health"
      ],
      "notes": "Use this source to justify maintaining Sleep & Circadian Rhythm as a research Topic and to motivate further intervention-focused discovery. Do not publish regular sleep improves cognition from this evidence."
    },
    {
      "schema_version": 1,
      "id": "testing-effect-direct-forward-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/testing-effect-direct-forward-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/testing-effect-direct-forward-2026/index.json",
      "decision": "challenge-existing",
      "decision_label": "Challenge existing claim",
      "reviewed_at": "2026-08-18",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Research on the testing effect routinely conflates direct and forward testing effects: A meta-analysis of testing effects with free recall",
        "url": "https://pubmed.ncbi.nlm.nih.gov/42258276/",
        "type": "meta-analysis",
        "doi": "10.1037/xlm0001634",
        "citation_text": "Gereau AD, Mulligan NW. Journal of Experimental Psychology: Learning, Memory, and Cognition. 2026. Online ahead of print.",
        "study_design": "Meta-analysis of free-recall testing effects distinguishing studies capable of isolating the direct testing effect from mixed designs susceptible to both direct and forward testing effects.",
        "population": "Human participants in the free-recall testing-effect literature included in the reviewed meta-analysis.",
        "intervention_or_exposure": "Retrieval practice/testing compared with non-testing study conditions, classified by whether the design isolates memory for tested material or also permits forward effects on subsequently learned material.",
        "outcomes": [
          "Free recall of tested material",
          "Free recall of subsequently learned material"
        ]
      },
      "supported_claim": "Retrieval practice can improve memory, but the literature commonly labeled as the testing effect mixes at least two distinct effects. In this meta-analysis, 66% of included testing effects came from mixed studies and 34% from direct studies, with mixed studies producing larger effects.",
      "unsupported_or_overstated_claims": [
        "A single undifferentiated testing-effect estimate describes one coherent memory phenomenon.",
        "Evidence from a mixed design can automatically be interpreted as a direct benefit for the tested material.",
        "Retrieval practice has the same mechanism or effect across tested material and subsequently learned material.",
        "A broad testing-effect citation is sufficient evidence for skill transfer or correct application."
      ],
      "limitations": [
        "The meta-analysis focuses on free-recall testing-effect studies rather than every form of retrieval practice or every educational outcome.",
        "Its main contribution is separation of direct and forward effects; it does not negate the broader evidence that retrieval practice can support retention.",
        "The reported mixture of study designs means older aggregate claims may need narrower interpretation rather than wholesale rejection."
      ],
      "target_hack_ids": [
        "active-recall-test-yourself"
      ],
      "target_protocol_ids": [
        "brali:active-recall-test-yourself"
      ],
      "risk_flags": [],
      "notes": "Use this decision to remove undifferentiated testing-effect language from Brali. Public guidance should say what outcome is being trained: remembering tested material, later learning, or application."
    },
    {
      "schema_version": 1,
      "id": "wakeful-rest-memory-2026",
      "canonical_url": "https://brali-lifeos.github.io/evidence/wakeful-rest-memory-2026/",
      "json_url": "https://brali-lifeos.github.io/evidence/wakeful-rest-memory-2026/index.json",
      "decision": "propose-protocol",
      "decision_label": "Protocol candidate",
      "reviewed_at": "2026-08-18",
      "reviewed_by": "Brali Evidence Reviewer",
      "source": {
        "title": "Should we all just take 10? A meta-analysis of wakeful rest",
        "url": "https://link.springer.com/article/10.3758/s13423-025-02778-3",
        "type": "meta-analysis",
        "doi": "10.3758/s13423-025-02778-3",
        "citation_text": "Parra D, Zhang Z, Radvansky G. Psychonomic Bulletin & Review. 2026;33:49.",
        "study_design": "Random-effects meta-analysis with moderator analyses and meta-regression; 142 effect sizes from 51 human experimental studies.",
        "population": "Patients with anterograde amnesia, healthy older adults, and healthy young to middle-aged adults; child studies were too sparse for the main group analysis.",
        "intervention_or_exposure": "A short period of quiet wakeful rest or minimal interference immediately after learning, compared with a post-learning interference task.",
        "outcomes": [
          "Delayed declarative-memory performance",
          "Reduced forgetting after learning"
        ]
      },
      "supported_claim": "A short period of quiet, minimally interfering wakeful rest after learning can reduce forgetting on average. The pooled effect is substantially smaller and less consistent in healthy young to middle-aged adults than in older adults or patient samples.",
      "unsupported_or_overstated_claims": [
        "Ten minutes is an optimal or universally evidence-based dose.",
        "Everyone will remember more after post-learning rest.",
        "Wakeful rest is equivalent to sleep or replaces adequate sleep.",
        "A specific posture, eye-closure rule, breathing pattern, or rehearsal instruction is required.",
        "The protocol is a treatment for memory impairment or dementia."
      ],
      "limitations": [
        "There was moderate to large between-study heterogeneity overall.",
        "The overall prediction interval included negative effects, so future studies are not guaranteed to find a benefit.",
        "Healthy young to middle-aged adults had the smallest pooled effect (g=0.20 in the authors' group analysis).",
        "Patient studies were small and showed funnel-plot asymmetry consistent with possible publication bias.",
        "Learning materials, interference tasks, test procedures, age, and clinical status varied across studies.",
        "The meta-analysis concerns declarative memory and does not justify broad claims about skill acquisition, creativity, focus, or general intelligence."
      ],
      "target_hack_ids": [
        "avoid-list-interference-memory-retention"
      ],
      "target_protocol_ids": [
        "brali:avoid-list-interference-memory-retention"
      ],
      "risk_flags": [],
      "notes": "Promote only as a conservative personal experiment. Use quiet/minimal interference after learning as the action and treat duration as user-chosen; do not turn the paper title into a universal ten-minute prescription."
    }
  ]
}
