feat(sql): add gov_population_yearly view

Exposes one row per (year, canonical_govid) drawn from long.population.
Used by per-capita denominators and peer matching.
This commit is contained in:
2026-04-29 15:12:53 -04:00
parent a28fb2e19b
commit ed9658d267
2 changed files with 35 additions and 0 deletions
+8
View File
@@ -0,0 +1,8 @@
CREATE OR REPLACE VIEW gov_population_yearly AS
SELECT DISTINCT
year,
canonical_govid,
population,
popyear
FROM long
WHERE population IS NOT NULL;
+27
View File
@@ -57,3 +57,30 @@ test_that("spending_annotated carries category + xwalk columns", {
expect_true(nm %in% names(row), info = paste("missing column:", nm)) expect_true(nm %in% names(row), info = paste("missing column:", nm))
} }
}) })
test_that("gov_population_yearly exposes one row per (year, canonical_govid)", {
skip_if_no_corpus()
with_fixture_corpus({
con <- uscogdata:::.ensure_session()
df <- DBI::dbGetQuery(
con,
"SELECT year, canonical_govid, population, popyear
FROM gov_population_yearly
WHERE canonical_govid = '101006006'
ORDER BY year"
)
expect_setequal(df$year, c(2019L, 2020L))
expect_equal(nrow(df), 2L)
expect_true(all(!is.na(df$population)))
expect_equal(df$population[df$year == 2019L], 1935878L)
expect_equal(df$population[df$year == 2020L], 1952778L)
# Uniqueness on (year, canonical_govid) across the whole view.
dup <- DBI::dbGetQuery(
con,
"SELECT year, canonical_govid, COUNT(*) AS n
FROM gov_population_yearly
GROUP BY year, canonical_govid HAVING n > 1"
)
expect_equal(nrow(dup), 0L)
})
})