Provider-agnostic LLM registry, config check, migration runbook

[llm] assigns the bulk and editor roles by name over a [providers.*]
registry (kind = openai | anthropic, per-provider model, effort, daily
ceiling and price table); DeepseekBackend becomes OpenAiCompatibleBackend
(reasoning_effort passthrough), AnthropicBackend builds from the same
ProviderConfig, meters and provider_costs are keyed by provider name.
Gemini 3.8 Flash is declared via Google's OpenAI-compatible endpoint so
switching the editor is one line (or DAILY_EPUB_LLM__EDITOR=gemini for an
A/B dry run). Stale [deepseek]/[anthropic] tables, the top-level
max_daily_usd and the old key env vars fail loudly.

daily-epub config check validates and prints the resolved roles, models,
key presence and paths without opening the database.

docs/runbooks/curation-v2-migration.md walks the server upgrade from v1.

Registry implemented by a Claude agent from an orchestrator brief;
verified fmt/clippy(-W dead_code)/test green.

Co-Authored-By: Claude Fable 5.1 <noreply@anthropic.com>
Claude-Session: https://claude.ai/code/session_01A1rCLQeKBgnBo3oTgHuTMe
This commit is contained in:
2026-09-02 18:44:57 +00:00
co-authored by Claude Fable 5.1
parent 5edbeb509f
commit 99d1338890
19 changed files with 2136 additions and 620 deletions
+11 -11
View File
@@ -570,9 +570,9 @@ async fn llm_pipeline_runs_against_a_mock_backend() {
usage,
);
let meter = UsageMeter::new(&cfg.deepseek, cfg.max_daily_usd);
let meter = UsageMeter::for_provider(&cfg.providers["deepseek"]);
let llm = LlmClient::with_backend(
&cfg.deepseek.model,
&cfg.providers["deepseek"].model,
"You are the editor of The Daily EPUB.".into(),
meter.clone(),
backend.clone(),
@@ -642,9 +642,9 @@ async fn llm_pipeline_runs_against_a_mock_backend() {
let colophon = Colophon {
provider_costs: BTreeMap::from([("deepseek".to_string(), meter.cost_usd())]),
models: Models {
bulk: cfg.deepseek.model.clone(),
editor: format!("{} (bulk fallback)", cfg.deepseek.model),
summaries: cfg.deepseek.model.clone(),
bulk: cfg.providers["deepseek"].model.clone(),
editor: format!("{} (bulk fallback)", cfg.providers["deepseek"].model),
summaries: cfg.providers["deepseek"].model.clone(),
},
entries_fetched: 8,
feeds_seen: 8,
@@ -655,7 +655,7 @@ async fn llm_pipeline_runs_against_a_mock_backend() {
let mut lineup = lineup;
pipeline::apply_summaries(&mut lineup, &editorial_doc);
let issue = assemble_build_publish(&db, &cfg, lineup, colophon).await;
assert_eq!(issue.colophon.models.bulk, cfg.deepseek.model);
assert_eq!(issue.colophon.models.bulk, cfg.providers["deepseek"].model);
assert!(issue.colophon.cost_usd > 0.0);
}
@@ -682,9 +682,9 @@ async fn failing_deepseek_still_publishes_with_heuristic_fallbacks() {
let backend = std::sync::Arc::new(MockBackend::new());
let client = LlmClient::with_backend(
&cfg.deepseek.model,
&cfg.providers["deepseek"].model,
"reader profile".into(),
UsageMeter::new(&cfg.deepseek, cfg.max_daily_usd),
UsageMeter::for_provider(&cfg.providers["deepseek"]),
backend.clone(),
);
let curator = Curator::new(
@@ -732,9 +732,9 @@ async fn failing_deepseek_still_publishes_with_heuristic_fallbacks() {
Colophon {
provider_costs: BTreeMap::new(),
models: Models {
bulk: cfg.deepseek.model.clone(),
editor: format!("{} (bulk fallback)", cfg.deepseek.model),
summaries: cfg.deepseek.model.clone(),
bulk: cfg.providers["deepseek"].model.clone(),
editor: format!("{} (bulk fallback)", cfg.providers["deepseek"].model),
summaries: cfg.providers["deepseek"].model.clone(),
},
entries_fetched: 8,
feeds_seen: 8,