pipeline
Classical end-to-end empirical analysis workflow in the traditional Python econometric stack — pandas + numpy + scipy + statsmodels + linearmodels + pyfixest +…
R-based econometric analysis for academic research. Use when writing R code for panel data, difference-in-differences, instrumental variables, spatial econometrics, or regression analysis. Covers data.table, fixest, sf, modelsummary, and publication-ready outputs.
$ npx -y skills add brycewang-stanford/Auto-Empirical-Research-Skills --skill econometrics-r --agent claude-codeHow it fires
How this skill gets triggered: by you, by Claude, or both.
/econometrics-rContext preview
The summary Claude sees to decide when to auto-load this skill.
R-based econometric analysis for academic research. Use when writing R code for panel data, difference-in-differences, instrumental variables, spatial econometrics, or regression analysis. Covers data.table, fixest, sf, modelsummary, and publication-ready outputs.
name: econometrics-r description: R-based econometric analysis for academic research. Use when writing R code for panel data, difference-in-differences, instrumental variables, spatial econometrics, or regression analysis. Covers data.table, fixest, sf, modelsummary, and publication-ready outputs.
library(data.table) # Data manipulation library(fixest) # Fixed effects estimation library(modelsummary) # Regression tables library(ggplot2) # Visualization library(sf) # Spatial data library(here) # Project paths
# Read and assign
dt <- fread(here("data", "raw", "file.csv"))
# Common operations
dt[, new_var := old_var * 100] # Create variable
dt[, mean_y := mean(y, na.rm = TRUE), by = group] # Group operations
dt[year >= 2000 & treated == 1] # Filter
dt[, .(mean_y = mean(y), n = .N), by = group] # Summarize
dt[other_dt, on = .(id, year)] # Merge
# Lag/lead within groups
setorder(dt, id, year)
dt[, lag_y := shift(y, 1), by = id]
dt[, lead_y := shift(y, -1), by = id]# Two-way fixed effects est1 <- feols(y ~ treatment + controls | id + year, data = dt) # Clustered standard errors (default: fixed effect groups) est2 <- feols(y ~ treatment | id + year, data = dt, cluster = ~state) # IV regression est3 <- feols(y ~ controls | id + year | endog ~ instrument, data = dt)
# Classic 2x2 DiD est_did <- feols(y ~ treated:post | id + year, data = dt) # Event study / dynamic effects dt[, rel_time := year - treatment_year] dt[, rel_time := fifelse(is.na(rel_time), -1000, rel_time)] # Never-treated est_es <- feols(y ~ i(rel_time, ref = -1) | id + year, data = dt) iplot(est_es) # Coefficient plot
# Sun-Abraham (requires cohort variable) est_sa <- feols(y ~ sunab(cohort, year) | id + year, data = dt) # Multiple estimators comparison library(did) # Callaway-Sant'Anna
models <- list(
"OLS" = est1,
"With FE" = est2,
"IV" = est3
)
modelsummary(models,
stars = c('*' = 0.1, '**' = 0.05, '***' = 0.01),
coef_omit = "Intercept",
gof_omit = "AIC|BIC|Log",
output = here("output", "tables", "main_results.tex")
)etable(est1, est2, est3,
se.below = TRUE,
keep = "treatment",
fitstat = c("n", "r2", "fe"),
tex = TRUE,
file = here("output", "tables", "results.tex")
)
# example
etable(
m1.suit, m2.suit,
dict = c(
'gruter_1' = 'Gruter Suitability 1',
'gruter_2' = 'Gruter Suitability 2',
'gruter_3' = 'Gruter Suitability 3',
'gruter_4' = 'Gruter Suitability 4',
'area_ha' = 'Orchard Size (ha)',
'yield' = 'Yield (kg/ha), 2023'
),
extralines = list(
'_Average yield (kg/ha)' = c(
round(mean(yields[area_ha > 1 & year == 2023, yield], na.rm = TRUE), 2),
round(mean(yields[area_ha > 1 & year == 2023, yield], na.rm = TRUE), 2)
),
'_Average orchard size (ha)' = c(
round(mean(yields[area_ha > 1 & year == 2023, area_ha], na.rm = TRUE), 2),
round(mean(yields[area_ha > 1 & year == 2023, area_ha], na.rm = TRUE), 2)
)
),
tex = TRUE,
style.tex = style.tex('aer'),
digits = 3,
depvar = TRUE
)coef_data <- broom::tidy(est_es, conf.int = TRUE)
ggplot(coef_data, aes(x = term, y = estimate)) +
geom_point() +
geom_errorbar(aes(ymin = conf.low, ymax = conf.high), width = 0.2) +
geom_hline(yintercept = 0, linetype = "dashed") +
theme_bw() +
labs(x = "Period", y = "Coefficient")
ggsave(here("output", "figures", "event_study.pdf"), width = 8, height = 5)library(sf)
map_data <- st_read(here("data", "raw", "shapefile.shp"))
map_data <- merge(map_data, results_dt, by = "region_id")
ggplot(map_data) +
geom_sf(aes(fill = estimate), color = "white", size = 0.1) +
scale_fill_viridis_c() +
theme_void()library(spdep) library(spatialreg) # Create spatial weights coords <- st_coordinates(st_centroid(map_data)) nb <- knn2nb(knearneigh(coords, k = 5)) W <- nb2listw(nb, style = "W") # Spatial lag model est_sar <- lagsarlm(y ~ x1 + x2, data = map_data, listw = W) # Spatial error model est_sem <- errorsarlm(y ~ x1 + x2, data = map_data, listw = W)
library(grf) # Generalized random forests # Causal forest cf <- causal_forest( X = as.matrix(dt[, .(x1, x2, x3)]), Y = dt$y, W = dt$treatment ) # Treatment effects ate <- average_treatment_effect(cf) cate <- predict(cf)$predictions
📌 文档结构(2026-07-22 起): 本文件是中文默认入口 —— banner + badges + 信任面 + 9 阶段流水线速览 + 76 行合集总表。 每个合集的完整描述、按用途分组、精确数字、验证方法在 docs/CONTENT_ZH.md(扩展正文,总表行内的 → 直接跳转到对应锚点)。 English version: README-en.md · 中文扩展正文:docs/CONTENT_ZH.md · README-zh-CN.md 已弃用(重定向占位) 🌐 语言: English |
Classical end-to-end empirical analysis workflow in the traditional Python econometric stack — pandas + numpy + scipy + statsmodels + linearmodels + pyfixest +…
Use when the user asks to run a full empirical / causal analysis in Python — by default in the style of an applied economics paper (AER / QJE / JPE / ReStud /…
Classical end-to-end empirical analysis workflow in the traditional Python econometric stack — pandas + numpy + scipy + statsmodels + linearmodels + pyfixest +…
Classical end-to-end empirical analysis workflow in the traditional Stata ecosystem — native Stata + reghdfe + ivreg2 + csdid + did_imputation +…
Classical end-to-end empirical analysis workflow in the modern tidyverse + econometrics R ecosystem — dplyr + tidyr + haven + fixest + sandwich + lmtest +…
Systematic writing framework for philosophy and interdisciplinary academic papers from optimized outline to submission-ready manuscript. Use when users want…