mirror of
https://github.com/ksyasuda/SubMiner.git
synced 2026-08-16 13:55:51 -07:00
Compare commits
7 Commits
| Author | SHA1 | Date | |
|---|---|---|---|
|
d220055b99
|
|||
|
65911474cf
|
|||
|
05da8cf86a
|
|||
|
e53f6d61ff
|
|||
|
48fdac62a6
|
|||
|
9ae303af7d
|
|||
|
297aa382c6
|
@@ -0,0 +1,6 @@
|
|||||||
|
type: fixed
|
||||||
|
area: stats
|
||||||
|
|
||||||
|
- Fixed Vocabulary totals and charts counting only the first browsing page instead of all tracked vocabulary, without delaying the rest of the page.
|
||||||
|
- New-word history now uses permanent daily lexical rollups, backfilled in the background and repaired when tracked material is removed or reprocessed.
|
||||||
|
- Calendar-day chart labels now preserve the recorded local date in time zones west of UTC.
|
||||||
@@ -82,7 +82,7 @@ Expandable session history with new-word activity, cumulative totals, and pause/
|
|||||||
|
|
||||||
#### Vocabulary
|
#### Vocabulary
|
||||||
|
|
||||||
Top repeated words (click a bar to open the word), new-word timeline, cross-title and frequency rank tables with Hide Known / Hide Kana filters, kanji breakdown, word exclusion list, and click-through occurrence drilldown with Mine Word / Mine Sentence / Mine Audio buttons.
|
The summary cards show all unique vocabulary and kanji recorded in the local tracking database; **New This Week** is the only weekly figure and uses a rolling seven-day window. The word and kanji tables load first while those complete totals calculate separately. Top Repeated Words and New Words by Day use complete tracking history rather than the table's browsing page; new-word history is maintained as a permanent daily lexical rollup, including retroactive corrections when tracked material is removed or reprocessed. On the first launch after upgrading, that history is built in the background and the chart refreshes when it is ready. The rest of the tab includes cross-title and frequency rank tables with Hide Known / Hide Kana filters, kanji breakdown, word exclusion list, and click-through occurrence drilldown with Mine Word / Mine Sentence / Mine Audio buttons.
|
||||||
|
|
||||||

|

|
||||||
|
|
||||||
@@ -180,6 +180,7 @@ In practice:
|
|||||||
- Anime and episode pages keep lifetime totals from summary tables while session drill-down still reads retained sessions directly. With the current defaults, both are kept forever.
|
- Anime and episode pages keep lifetime totals from summary tables while session drill-down still reads retained sessions directly. With the current defaults, both are kept forever.
|
||||||
- Trends can read the full available history because daily/monthly rollups are also kept forever by default.
|
- Trends can read the full available history because daily/monthly rollups are also kept forever by default.
|
||||||
- Vocabulary and kanji totals are cumulative and not bounded by the raw session retention knobs.
|
- Vocabulary and kanji totals are cumulative and not bounded by the raw session retention knobs.
|
||||||
|
- New-word charts use their own permanent lexical daily rollups, which are not pruned by activity-rollup retention.
|
||||||
|
|
||||||
## Storage / Performance Model
|
## Storage / Performance Model
|
||||||
|
|
||||||
@@ -349,6 +350,7 @@ Rollup tables:
|
|||||||
|
|
||||||
- `imm_daily_rollups`
|
- `imm_daily_rollups`
|
||||||
- `imm_monthly_rollups`
|
- `imm_monthly_rollups`
|
||||||
|
- `imm_lexical_daily_rollups` - permanent first-discovery counts for vocabulary and kanji chart history
|
||||||
- `imm_rollup_state` - incremental rollup progress bookkeeping
|
- `imm_rollup_state` - incremental rollup progress bookkeeping
|
||||||
|
|
||||||
Vocabulary tables:
|
Vocabulary tables:
|
||||||
|
|||||||
@@ -284,6 +284,22 @@ function createMockTracker(
|
|||||||
getSessionTimeline: async () => [],
|
getSessionTimeline: async () => [],
|
||||||
getSessionEvents: async () => [],
|
getSessionEvents: async () => [],
|
||||||
getVocabularyStats: async () => VOCABULARY_STATS,
|
getVocabularyStats: async () => VOCABULARY_STATS,
|
||||||
|
getVocabularySummary: async () => ({
|
||||||
|
uniqueWords: 501,
|
||||||
|
uniqueWordsWithoutNames: 500,
|
||||||
|
uniqueKanji: 201,
|
||||||
|
newThisWeek: 7,
|
||||||
|
newThisWeekWithoutNames: 6,
|
||||||
|
knownWordCount: 250,
|
||||||
|
knownWordCountWithoutNames: 249,
|
||||||
|
}),
|
||||||
|
getVocabularyChartData: async () => ({
|
||||||
|
ready: true,
|
||||||
|
topWords: [{ wordId: 1, headword: 'する', frequency: 50 }],
|
||||||
|
topWordsWithoutNames: [{ wordId: 1, headword: 'する', frequency: 50 }],
|
||||||
|
newWordsTimeline: [{ epochDay: 20_000, wordCount: 3 }],
|
||||||
|
newWordsTimelineWithoutNames: [{ epochDay: 20_000, wordCount: 3 }],
|
||||||
|
}),
|
||||||
getStatsExcludedWords: async () => [],
|
getStatsExcludedWords: async () => [],
|
||||||
replaceStatsExcludedWords: async () => {},
|
replaceStatsExcludedWords: async () => {},
|
||||||
getKanjiStats: async () => KANJI_STATS,
|
getKanjiStats: async () => KANJI_STATS,
|
||||||
@@ -711,6 +727,23 @@ describe('stats server API routes', () => {
|
|||||||
assert.equal(body[0].headword, 'する');
|
assert.equal(body[0].headword, 'する');
|
||||||
});
|
});
|
||||||
|
|
||||||
|
it('GET /api/stats/vocabulary/summary returns database-wide card totals', async () => {
|
||||||
|
const app = createStatsApp(createMockTracker());
|
||||||
|
|
||||||
|
const res = await app.request('/api/stats/vocabulary/summary');
|
||||||
|
|
||||||
|
assert.equal(res.status, 200);
|
||||||
|
assert.deepEqual(await res.json(), {
|
||||||
|
uniqueWords: 501,
|
||||||
|
uniqueWordsWithoutNames: 500,
|
||||||
|
uniqueKanji: 201,
|
||||||
|
newThisWeek: 7,
|
||||||
|
newThisWeekWithoutNames: 6,
|
||||||
|
knownWordCount: 250,
|
||||||
|
knownWordCountWithoutNames: 249,
|
||||||
|
});
|
||||||
|
});
|
||||||
|
|
||||||
it('GET /api/stats/kanji returns kanji frequency data', async () => {
|
it('GET /api/stats/kanji returns kanji frequency data', async () => {
|
||||||
const app = createStatsApp(createMockTracker());
|
const app = createStatsApp(createMockTracker());
|
||||||
const res = await app.request('/api/stats/kanji');
|
const res = await app.request('/api/stats/kanji');
|
||||||
|
|||||||
@@ -559,6 +559,56 @@ test('fresh tracker DB creates lifetime summary tables', async () => {
|
|||||||
}
|
}
|
||||||
});
|
});
|
||||||
|
|
||||||
|
test('fresh tracker DB skips lexical rollup backfill work', async () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
let tracker: ImmersionTrackerService | null = null;
|
||||||
|
let backfillRuns = 0;
|
||||||
|
|
||||||
|
try {
|
||||||
|
const Ctor = await loadTrackerCtor();
|
||||||
|
tracker = new Ctor({ dbPath }, {
|
||||||
|
runLexicalRollupBackfillTask: async () => {
|
||||||
|
backfillRuns += 1;
|
||||||
|
},
|
||||||
|
} as never);
|
||||||
|
|
||||||
|
assert.equal(backfillRuns, 0);
|
||||||
|
} finally {
|
||||||
|
tracker?.destroy();
|
||||||
|
cleanupDbPath(dbPath);
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('tracker starts the injected lexical rollup backfill when it is pending', async () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
let tracker: ImmersionTrackerService | null = null;
|
||||||
|
let backfillRuns = 0;
|
||||||
|
|
||||||
|
try {
|
||||||
|
const setupDb = new Database(dbPath);
|
||||||
|
const { ensureSchema } = await import('./immersion-tracker/storage');
|
||||||
|
ensureSchema(setupDb);
|
||||||
|
setupDb
|
||||||
|
.prepare(
|
||||||
|
`UPDATE imm_rollup_state SET state_value = '0' WHERE state_key = 'lexical_daily_rollups_ready'`,
|
||||||
|
)
|
||||||
|
.run();
|
||||||
|
setupDb.close();
|
||||||
|
|
||||||
|
const Ctor = await loadTrackerCtor();
|
||||||
|
tracker = new Ctor({ dbPath }, {
|
||||||
|
runLexicalRollupBackfillTask: async () => {
|
||||||
|
backfillRuns += 1;
|
||||||
|
},
|
||||||
|
} as never);
|
||||||
|
|
||||||
|
assert.equal(backfillRuns, 1);
|
||||||
|
} finally {
|
||||||
|
tracker?.destroy();
|
||||||
|
cleanupDbPath(dbPath);
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
test('startup backfills lifetime summaries when retained sessions exist but summary tables are empty', async () => {
|
test('startup backfills lifetime summaries when retained sessions exist but summary tables are empty', async () => {
|
||||||
const dbPath = makeDbPath();
|
const dbPath = makeDbPath();
|
||||||
let tracker: ImmersionTrackerService | null = null;
|
let tracker: ImmersionTrackerService | null = null;
|
||||||
|
|||||||
@@ -58,6 +58,7 @@ import {
|
|||||||
getSessionEvents,
|
getSessionEvents,
|
||||||
getSimilarWords,
|
getSimilarWords,
|
||||||
getStatsExcludedWords,
|
getStatsExcludedWords,
|
||||||
|
getVocabularyChartData,
|
||||||
getVocabularyStats,
|
getVocabularyStats,
|
||||||
replaceStatsExcludedWords,
|
replaceStatsExcludedWords,
|
||||||
searchSubtitleSentences,
|
searchSubtitleSentences,
|
||||||
@@ -96,6 +97,12 @@ import {
|
|||||||
DeleteMaintenanceWorkerRuntime,
|
DeleteMaintenanceWorkerRuntime,
|
||||||
type RunDeleteMaintenanceTask,
|
type RunDeleteMaintenanceTask,
|
||||||
} from './immersion-tracker/delete-maintenance-worker-runtime';
|
} from './immersion-tracker/delete-maintenance-worker-runtime';
|
||||||
|
import {
|
||||||
|
VocabularySummaryWorkerRuntime,
|
||||||
|
type RunVocabularySummaryTask,
|
||||||
|
} from './immersion-tracker/vocabulary-summary-worker-runtime';
|
||||||
|
import { LexicalRollupWorkerRuntime } from './immersion-tracker/lexical-rollup-worker-runtime';
|
||||||
|
import { areLexicalDailyRollupsReady } from './immersion-tracker/lexical-rollups';
|
||||||
import { DeleteMaintenanceScheduler } from './immersion-tracker/delete-maintenance-scheduler';
|
import { DeleteMaintenanceScheduler } from './immersion-tracker/delete-maintenance-scheduler';
|
||||||
import {
|
import {
|
||||||
cleanupDuplicateSubtitleLines,
|
cleanupDuplicateSubtitleLines,
|
||||||
@@ -185,6 +192,7 @@ import {
|
|||||||
type StatsExcludedWordRow,
|
type StatsExcludedWordRow,
|
||||||
type StreakCalendarRow,
|
type StreakCalendarRow,
|
||||||
type VocabularyCleanupSummary,
|
type VocabularyCleanupSummary,
|
||||||
|
type VocabularyStatsSummary,
|
||||||
type WatchTimePerAnimeRow,
|
type WatchTimePerAnimeRow,
|
||||||
type WordAnimeAppearanceRow,
|
type WordAnimeAppearanceRow,
|
||||||
type WordDetailRow,
|
type WordDetailRow,
|
||||||
@@ -407,6 +415,12 @@ export class ImmersionTrackerService {
|
|||||||
private readonly dbPath: string;
|
private readonly dbPath: string;
|
||||||
private readonly writeLock = { locked: false };
|
private readonly writeLock = { locked: false };
|
||||||
private readonly destroyDeleteMaintenanceRunner: () => void;
|
private readonly destroyDeleteMaintenanceRunner: () => void;
|
||||||
|
private readonly runVocabularySummaryTask: (
|
||||||
|
knownWords: ReadonlySet<string> | null,
|
||||||
|
) => Promise<VocabularyStatsSummary>;
|
||||||
|
private readonly destroyVocabularySummaryRunner: () => void;
|
||||||
|
private readonly runLexicalRollupBackfillTask: () => Promise<void>;
|
||||||
|
private readonly destroyLexicalRollupBackfillRunner: () => void;
|
||||||
private readonly deleteMaintenanceScheduler: DeleteMaintenanceScheduler;
|
private readonly deleteMaintenanceScheduler: DeleteMaintenanceScheduler;
|
||||||
private flushTimer: ReturnType<typeof setTimeout> | null = null;
|
private flushTimer: ReturnType<typeof setTimeout> | null = null;
|
||||||
private maintenanceTimer: ReturnType<typeof setInterval> | null = null;
|
private maintenanceTimer: ReturnType<typeof setInterval> | null = null;
|
||||||
@@ -434,6 +448,10 @@ export class ImmersionTrackerService {
|
|||||||
dependencies: {
|
dependencies: {
|
||||||
runDeleteMaintenanceTask?: RunDeleteMaintenanceTask;
|
runDeleteMaintenanceTask?: RunDeleteMaintenanceTask;
|
||||||
destroyDeleteMaintenanceRunner?: () => void;
|
destroyDeleteMaintenanceRunner?: () => void;
|
||||||
|
runVocabularySummaryTask?: RunVocabularySummaryTask;
|
||||||
|
destroyVocabularySummaryRunner?: () => void;
|
||||||
|
runLexicalRollupBackfillTask?: (dbPath: string) => Promise<void>;
|
||||||
|
destroyLexicalRollupBackfillRunner?: () => void;
|
||||||
} = {},
|
} = {},
|
||||||
) {
|
) {
|
||||||
this.dbPath = options.dbPath;
|
this.dbPath = options.dbPath;
|
||||||
@@ -460,6 +478,27 @@ export class ImmersionTrackerService {
|
|||||||
if (!this.isDestroyed && this.queue.length > 0) this.scheduleFlush(0);
|
if (!this.isDestroyed && this.queue.length > 0) this.scheduleFlush(0);
|
||||||
},
|
},
|
||||||
});
|
});
|
||||||
|
if (dependencies.runVocabularySummaryTask) {
|
||||||
|
this.runVocabularySummaryTask = (knownWords) =>
|
||||||
|
dependencies.runVocabularySummaryTask!(this.dbPath, knownWords);
|
||||||
|
this.destroyVocabularySummaryRunner =
|
||||||
|
dependencies.destroyVocabularySummaryRunner ?? (() => {});
|
||||||
|
} else {
|
||||||
|
const vocabularySummaryRuntime = new VocabularySummaryWorkerRuntime();
|
||||||
|
this.runVocabularySummaryTask = (knownWords) =>
|
||||||
|
vocabularySummaryRuntime.run(this.dbPath, knownWords);
|
||||||
|
this.destroyVocabularySummaryRunner = () => vocabularySummaryRuntime.destroy();
|
||||||
|
}
|
||||||
|
if (dependencies.runLexicalRollupBackfillTask) {
|
||||||
|
this.runLexicalRollupBackfillTask = () =>
|
||||||
|
dependencies.runLexicalRollupBackfillTask!(this.dbPath);
|
||||||
|
this.destroyLexicalRollupBackfillRunner =
|
||||||
|
dependencies.destroyLexicalRollupBackfillRunner ?? (() => {});
|
||||||
|
} else {
|
||||||
|
const lexicalRollupRuntime = new LexicalRollupWorkerRuntime();
|
||||||
|
this.runLexicalRollupBackfillTask = () => lexicalRollupRuntime.run(this.dbPath);
|
||||||
|
this.destroyLexicalRollupBackfillRunner = () => lexicalRollupRuntime.destroy();
|
||||||
|
}
|
||||||
const parentDir = path.dirname(this.dbPath);
|
const parentDir = path.dirname(this.dbPath);
|
||||||
if (!fs.existsSync(parentDir)) {
|
if (!fs.existsSync(parentDir)) {
|
||||||
fs.mkdirSync(parentDir, { recursive: true });
|
fs.mkdirSync(parentDir, { recursive: true });
|
||||||
@@ -519,6 +558,14 @@ export class ImmersionTrackerService {
|
|||||||
this.db = new Database(this.dbPath);
|
this.db = new Database(this.dbPath);
|
||||||
applyPragmas(this.db);
|
applyPragmas(this.db);
|
||||||
ensureSchema(this.db);
|
ensureSchema(this.db);
|
||||||
|
if (!areLexicalDailyRollupsReady(this.db)) {
|
||||||
|
void this.runLexicalRollupBackfillTask().catch((error: unknown) => {
|
||||||
|
this.logger.warn(
|
||||||
|
'Lexical daily rollup backfill failed; it will retry on next startup',
|
||||||
|
error,
|
||||||
|
);
|
||||||
|
});
|
||||||
|
}
|
||||||
const reconciledSessions = reconcileStaleActiveSessions(this.db);
|
const reconciledSessions = reconcileStaleActiveSessions(this.db);
|
||||||
if (reconciledSessions > 0) {
|
if (reconciledSessions > 0) {
|
||||||
this.logger.info(
|
this.logger.info(
|
||||||
@@ -565,6 +612,8 @@ export class ImmersionTrackerService {
|
|||||||
this.isDestroyed = true;
|
this.isDestroyed = true;
|
||||||
this.deleteMaintenanceScheduler.destroy();
|
this.deleteMaintenanceScheduler.destroy();
|
||||||
this.destroyDeleteMaintenanceRunner();
|
this.destroyDeleteMaintenanceRunner();
|
||||||
|
this.destroyVocabularySummaryRunner();
|
||||||
|
this.destroyLexicalRollupBackfillRunner();
|
||||||
this.db.close();
|
this.db.close();
|
||||||
}
|
}
|
||||||
|
|
||||||
@@ -634,6 +683,14 @@ export class ImmersionTrackerService {
|
|||||||
return getVocabularyStats(this.db, limit, excludePos);
|
return getVocabularyStats(this.db, limit, excludePos);
|
||||||
}
|
}
|
||||||
|
|
||||||
|
async getVocabularySummary(knownWords: ReadonlySet<string> | null) {
|
||||||
|
return this.runVocabularySummaryTask(knownWords);
|
||||||
|
}
|
||||||
|
|
||||||
|
async getVocabularyChartData() {
|
||||||
|
return getVocabularyChartData(this.db);
|
||||||
|
}
|
||||||
|
|
||||||
async getStatsExcludedWords(): Promise<StatsExcludedWordRow[]> {
|
async getStatsExcludedWords(): Promise<StatsExcludedWordRow[]> {
|
||||||
return getStatsExcludedWords(this.db);
|
return getStatsExcludedWords(this.db);
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -31,6 +31,7 @@ import {
|
|||||||
getKanjiOccurrences,
|
getKanjiOccurrences,
|
||||||
getSessionSummaries,
|
getSessionSummaries,
|
||||||
getVocabularyStats,
|
getVocabularyStats,
|
||||||
|
getVocabularySummary,
|
||||||
getKanjiStats,
|
getKanjiStats,
|
||||||
getSessionEvents,
|
getSessionEvents,
|
||||||
getSessionTimeline,
|
getSessionTimeline,
|
||||||
@@ -1875,6 +1876,88 @@ test('getVocabularyStats returns rows ordered by frequency descending', () => {
|
|||||||
}
|
}
|
||||||
});
|
});
|
||||||
|
|
||||||
|
test('getVocabularySummary counts every tracked vocabulary row instead of a display page', () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
const db = openTestDb(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
const nowSec = Math.floor(Date.now() / 1000);
|
||||||
|
const insertWord = db.prepare(`
|
||||||
|
INSERT INTO imm_words (
|
||||||
|
headword, word, reading, part_of_speech, pos1, pos2, pos3,
|
||||||
|
first_seen, last_seen, frequency
|
||||||
|
) VALUES (?, ?, '', 'noun', '名詞', '一般', '', ?, ?, 1)
|
||||||
|
`);
|
||||||
|
const insertKanji = db.prepare(`
|
||||||
|
INSERT INTO imm_kanji (kanji, first_seen, last_seen, frequency)
|
||||||
|
VALUES (?, ?, ?, 1)
|
||||||
|
`);
|
||||||
|
|
||||||
|
for (let index = 0; index < 501; index += 1) {
|
||||||
|
insertWord.run(`単語${index}`, `単語${index}`, nowSec - 8 * 86_400, nowSec - 8 * 86_400);
|
||||||
|
}
|
||||||
|
for (let index = 0; index < 201; index += 1) {
|
||||||
|
insertKanji.run(
|
||||||
|
String.fromCodePoint(0x4e00 + index),
|
||||||
|
nowSec - 8 * 86_400,
|
||||||
|
nowSec - 8 * 86_400,
|
||||||
|
);
|
||||||
|
}
|
||||||
|
insertWord.run('今週', '今週', nowSec - 86_400, nowSec - 86_400);
|
||||||
|
|
||||||
|
assert.deepEqual(getVocabularySummary(db, new Set(['単語0', '今週']), nowSec * 1000), {
|
||||||
|
uniqueWords: 502,
|
||||||
|
uniqueWordsWithoutNames: 502,
|
||||||
|
uniqueKanji: 201,
|
||||||
|
newThisWeek: 1,
|
||||||
|
newThisWeekWithoutNames: 1,
|
||||||
|
knownWordCount: 2,
|
||||||
|
knownWordCountWithoutNames: 2,
|
||||||
|
});
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
cleanupDbPath(dbPath);
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('getVocabularySummary applies vocabulary exclusions and Hide Names totals', () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
const db = openTestDb(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
const insertWord = db.prepare(`
|
||||||
|
INSERT INTO imm_words (
|
||||||
|
headword, word, reading, part_of_speech, pos1, pos2, pos3,
|
||||||
|
first_seen, last_seen, frequency
|
||||||
|
) VALUES (?, ?, '', 'noun', '名詞', ?, '', 1, 1, 1)
|
||||||
|
`);
|
||||||
|
insertWord.run('猫', '猫', '一般');
|
||||||
|
insertWord.run('太郎', '太郎', '固有名詞');
|
||||||
|
insertWord.run('東京', '東京都', '一般');
|
||||||
|
db.prepare(
|
||||||
|
`
|
||||||
|
INSERT INTO imm_stats_excluded_words (headword, word, reading)
|
||||||
|
VALUES ('東京', '東京', '')
|
||||||
|
`,
|
||||||
|
).run();
|
||||||
|
|
||||||
|
assert.deepEqual(getVocabularySummary(db, new Set(['猫', '太郎', '東京']), 9 * 86_400_000), {
|
||||||
|
uniqueWords: 2,
|
||||||
|
uniqueWordsWithoutNames: 1,
|
||||||
|
uniqueKanji: 0,
|
||||||
|
newThisWeek: 0,
|
||||||
|
newThisWeekWithoutNames: 0,
|
||||||
|
knownWordCount: 2,
|
||||||
|
knownWordCountWithoutNames: 1,
|
||||||
|
});
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
cleanupDbPath(dbPath);
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
test('getVocabularyStats filters rows that fail tokenizer vocabulary rules', () => {
|
test('getVocabularyStats filters rows that fail tokenizer vocabulary rules', () => {
|
||||||
const dbPath = makeDbPath();
|
const dbPath = makeDbPath();
|
||||||
const db = openTestDb(dbPath);
|
const db = openTestDb(dbPath);
|
||||||
|
|||||||
@@ -0,0 +1,103 @@
|
|||||||
|
import assert from 'node:assert/strict';
|
||||||
|
import fs from 'node:fs';
|
||||||
|
import os from 'node:os';
|
||||||
|
import path from 'node:path';
|
||||||
|
import test from 'node:test';
|
||||||
|
import {
|
||||||
|
LexicalRollupWorkerRuntime,
|
||||||
|
resolveLexicalRollupWorkerPath,
|
||||||
|
} from './lexical-rollup-worker-runtime';
|
||||||
|
import { areLexicalDailyRollupsReady } from './lexical-rollups';
|
||||||
|
import { Database } from './sqlite';
|
||||||
|
import { applyPragmas, ensureSchema } from './storage';
|
||||||
|
|
||||||
|
test('lexical rollup worker backfills without using the tracker connection', async () => {
|
||||||
|
const directory = fs.mkdtempSync(path.join(os.tmpdir(), 'subminer-lexical-rollup-runtime-'));
|
||||||
|
const dbPath = path.join(directory, 'immersion.sqlite');
|
||||||
|
const runtime = new LexicalRollupWorkerRuntime();
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
applyPragmas(db);
|
||||||
|
ensureSchema(db);
|
||||||
|
db.prepare(
|
||||||
|
`INSERT INTO imm_words(headword, word, reading, first_seen, last_seen, frequency)
|
||||||
|
VALUES ('鳥', '鳥', 'とり', 1700000000, 1700000000, 1)`,
|
||||||
|
).run();
|
||||||
|
db.exec('DELETE FROM imm_lexical_daily_rollups');
|
||||||
|
db.prepare(`UPDATE imm_rollup_state SET state_value = '0' WHERE state_key = ?`).run(
|
||||||
|
'lexical_daily_rollups_ready',
|
||||||
|
);
|
||||||
|
db.close();
|
||||||
|
|
||||||
|
await runtime.run(dbPath);
|
||||||
|
|
||||||
|
const checkDb = new Database(dbPath);
|
||||||
|
try {
|
||||||
|
assert.equal(areLexicalDailyRollupsReady(checkDb), true);
|
||||||
|
} finally {
|
||||||
|
checkDb.close();
|
||||||
|
}
|
||||||
|
} finally {
|
||||||
|
runtime.destroy();
|
||||||
|
try {
|
||||||
|
db.close();
|
||||||
|
} catch {
|
||||||
|
// Closed before the worker starts.
|
||||||
|
}
|
||||||
|
fs.rmSync(directory, { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('lexical rollup worker module resolves in the current layout', () => {
|
||||||
|
const workerPath = resolveLexicalRollupWorkerPath();
|
||||||
|
assert.ok(workerPath, 'expected the lexical rollup worker module to resolve');
|
||||||
|
assert.ok(workerPath.endsWith(__filename.endsWith('.ts') ? '.ts' : '.js'));
|
||||||
|
});
|
||||||
|
|
||||||
|
test('lexical rollup worker leaves a backfill pending when no worker can start', async () => {
|
||||||
|
const runtime = new LexicalRollupWorkerRuntime({
|
||||||
|
resolveWorkerPath: () => null,
|
||||||
|
warn: () => {},
|
||||||
|
} as never);
|
||||||
|
|
||||||
|
try {
|
||||||
|
await assert.doesNotReject(runtime.run('/tmp/not-used.sqlite'));
|
||||||
|
} finally {
|
||||||
|
runtime.destroy();
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('lexical rollup worker absorbs termination failures after settling', async () => {
|
||||||
|
let sendMessage: ((message: { ok: boolean }) => void) | null = null;
|
||||||
|
const runtime = new LexicalRollupWorkerRuntime({
|
||||||
|
resolveWorkerPath: () => '/tmp/fake-worker.js',
|
||||||
|
createWorker: async () => ({
|
||||||
|
once(event: string, listener: (value: never) => void) {
|
||||||
|
if (event === 'message') sendMessage = listener as (message: { ok: boolean }) => void;
|
||||||
|
return this;
|
||||||
|
},
|
||||||
|
terminate: async () => {
|
||||||
|
throw new Error('termination failed');
|
||||||
|
},
|
||||||
|
}),
|
||||||
|
warn: () => {},
|
||||||
|
} as never);
|
||||||
|
|
||||||
|
const unhandled: unknown[] = [];
|
||||||
|
const captureUnhandled = (reason: unknown) => unhandled.push(reason);
|
||||||
|
process.on('unhandledRejection', captureUnhandled);
|
||||||
|
try {
|
||||||
|
const task = runtime.run('/tmp/not-used.sqlite');
|
||||||
|
await new Promise((resolve) => setImmediate(resolve));
|
||||||
|
const notify = sendMessage as ((message: { ok: boolean }) => void) | null;
|
||||||
|
assert.ok(notify);
|
||||||
|
notify({ ok: true });
|
||||||
|
await task;
|
||||||
|
await new Promise((resolve) => setImmediate(resolve));
|
||||||
|
assert.deepEqual(unhandled, []);
|
||||||
|
} finally {
|
||||||
|
process.off('unhandledRejection', captureUnhandled);
|
||||||
|
runtime.destroy();
|
||||||
|
}
|
||||||
|
});
|
||||||
@@ -0,0 +1,109 @@
|
|||||||
|
import fs from 'node:fs';
|
||||||
|
import path from 'node:path';
|
||||||
|
import { createLogger } from '../../../logger';
|
||||||
|
|
||||||
|
interface WorkerResponse {
|
||||||
|
ok?: boolean;
|
||||||
|
error?: unknown;
|
||||||
|
}
|
||||||
|
|
||||||
|
interface WorkerHandle {
|
||||||
|
once(event: 'message', listener: (message: WorkerResponse) => void): this;
|
||||||
|
once(event: 'error', listener: (error: Error) => void): this;
|
||||||
|
once(event: 'exit', listener: (code: number) => void): this;
|
||||||
|
terminate(): Promise<number>;
|
||||||
|
}
|
||||||
|
|
||||||
|
interface LexicalRollupWorkerRuntimeOptions {
|
||||||
|
resolveWorkerPath?: () => string | null;
|
||||||
|
createWorker?: (workerPath: string, workerData: { dbPath: string }) => Promise<WorkerHandle>;
|
||||||
|
warn?: (message: string, ...meta: unknown[]) => void;
|
||||||
|
}
|
||||||
|
|
||||||
|
const logger = createLogger('main:immersion-tracker:lexical-rollup-worker');
|
||||||
|
|
||||||
|
export function resolveLexicalRollupWorkerPath(): string | null {
|
||||||
|
const fileName = __filename.endsWith('.ts')
|
||||||
|
? 'lexical-rollup-worker-thread.ts'
|
||||||
|
: 'lexical-rollup-worker-thread.js';
|
||||||
|
const workerPath = path.join(__dirname, fileName);
|
||||||
|
return fs.existsSync(workerPath) ? workerPath : null;
|
||||||
|
}
|
||||||
|
|
||||||
|
export class LexicalRollupWorkerRuntime {
|
||||||
|
private readonly activeWorkers = new Set<WorkerHandle>();
|
||||||
|
private destroyed = false;
|
||||||
|
|
||||||
|
constructor(private readonly options: LexicalRollupWorkerRuntimeOptions = {}) {}
|
||||||
|
|
||||||
|
async run(dbPath: string): Promise<void> {
|
||||||
|
if (this.destroyed) throw new Error('Lexical rollup worker is shut down');
|
||||||
|
let worker: WorkerHandle;
|
||||||
|
try {
|
||||||
|
const workerPath = (this.options.resolveWorkerPath ?? resolveLexicalRollupWorkerPath)();
|
||||||
|
if (!workerPath) throw new Error('Emitted lexical rollup worker module was not found');
|
||||||
|
const createWorker =
|
||||||
|
this.options.createWorker ??
|
||||||
|
(async (resolvedPath, workerData) => {
|
||||||
|
const { Worker } = await import('node:worker_threads');
|
||||||
|
return new Worker(resolvedPath, { workerData });
|
||||||
|
});
|
||||||
|
worker = await createWorker(workerPath, { dbPath });
|
||||||
|
} catch (error) {
|
||||||
|
if (this.destroyed) throw new Error('Lexical rollup worker is shut down');
|
||||||
|
(this.options.warn ?? logger.warn)(
|
||||||
|
'Lexical rollup worker unavailable; leaving backfill pending for a later startup',
|
||||||
|
error,
|
||||||
|
);
|
||||||
|
return;
|
||||||
|
}
|
||||||
|
|
||||||
|
if (this.destroyed) {
|
||||||
|
await worker.terminate().catch(() => undefined);
|
||||||
|
throw new Error('Lexical rollup worker is shut down');
|
||||||
|
}
|
||||||
|
|
||||||
|
return new Promise<void>((resolve, reject) => {
|
||||||
|
let settled = false;
|
||||||
|
this.activeWorkers.add(worker);
|
||||||
|
const settle = (error?: Error) => {
|
||||||
|
if (settled) return;
|
||||||
|
settled = true;
|
||||||
|
this.activeWorkers.delete(worker);
|
||||||
|
void worker.terminate().catch(() => undefined);
|
||||||
|
if (error) reject(error);
|
||||||
|
else resolve();
|
||||||
|
};
|
||||||
|
worker.once('message', (message) => {
|
||||||
|
if (message.ok) settle();
|
||||||
|
else
|
||||||
|
settle(
|
||||||
|
new Error(
|
||||||
|
`Lexical rollup backfill failed: ${String(message.error ?? 'unknown error')}`,
|
||||||
|
),
|
||||||
|
);
|
||||||
|
});
|
||||||
|
worker.once('error', (error) => settle(error));
|
||||||
|
worker.once('exit', (code) => {
|
||||||
|
if (!settled) {
|
||||||
|
settle(
|
||||||
|
new Error(
|
||||||
|
code === 0
|
||||||
|
? 'Lexical rollup worker exited without a response'
|
||||||
|
: `Lexical rollup worker exited with code ${code}`,
|
||||||
|
),
|
||||||
|
);
|
||||||
|
}
|
||||||
|
});
|
||||||
|
});
|
||||||
|
}
|
||||||
|
|
||||||
|
destroy(): void {
|
||||||
|
if (this.destroyed) return;
|
||||||
|
this.destroyed = true;
|
||||||
|
for (const worker of this.activeWorkers) {
|
||||||
|
void worker.terminate().catch(() => undefined);
|
||||||
|
}
|
||||||
|
this.activeWorkers.clear();
|
||||||
|
}
|
||||||
|
}
|
||||||
@@ -0,0 +1,11 @@
|
|||||||
|
import { parentPort, workerData } from 'node:worker_threads';
|
||||||
|
import { executeLexicalRollupBackfillTask } from './lexical-rollup-worker';
|
||||||
|
|
||||||
|
if (!parentPort) throw new Error('lexical rollup worker missing parent port');
|
||||||
|
|
||||||
|
try {
|
||||||
|
executeLexicalRollupBackfillTask((workerData as { dbPath: string }).dbPath);
|
||||||
|
parentPort.postMessage({ ok: true });
|
||||||
|
} catch (error) {
|
||||||
|
parentPort.postMessage({ error: error instanceof Error ? error.message : String(error) });
|
||||||
|
}
|
||||||
@@ -0,0 +1,35 @@
|
|||||||
|
import assert from 'node:assert/strict';
|
||||||
|
import fs from 'node:fs';
|
||||||
|
import os from 'node:os';
|
||||||
|
import path from 'node:path';
|
||||||
|
import test from 'node:test';
|
||||||
|
import { areLexicalDailyRollupsReady, getLexicalDailyRollups } from './lexical-rollups';
|
||||||
|
import { executeLexicalRollupBackfillTask } from './lexical-rollup-worker';
|
||||||
|
import { Database } from './sqlite';
|
||||||
|
import { ensureSchema } from './storage';
|
||||||
|
|
||||||
|
test('lexical rollup backfill materializes pre-existing vocabulary off the caller DB connection', () => {
|
||||||
|
const directory = fs.mkdtempSync(path.join(os.tmpdir(), 'subminer-lexical-rollup-worker-'));
|
||||||
|
const dbPath = path.join(directory, 'immersion.sqlite');
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
db.prepare(
|
||||||
|
`INSERT INTO imm_words(headword, word, reading, first_seen, last_seen, frequency)
|
||||||
|
VALUES (?, ?, ?, ?, ?, 1)`,
|
||||||
|
).run('犬', '犬', 'いぬ', 1_700_000_000, 1_700_000_000);
|
||||||
|
db.exec('DELETE FROM imm_lexical_daily_rollups');
|
||||||
|
db.prepare(`UPDATE imm_rollup_state SET state_value = '0' WHERE state_key = ?`).run(
|
||||||
|
'lexical_daily_rollups_ready',
|
||||||
|
);
|
||||||
|
|
||||||
|
executeLexicalRollupBackfillTask(dbPath);
|
||||||
|
|
||||||
|
assert.equal(areLexicalDailyRollupsReady(db), true);
|
||||||
|
assert.equal(getLexicalDailyRollups(db)[0]?.wordCount, 1);
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
fs.rmSync(directory, { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
@@ -0,0 +1,15 @@
|
|||||||
|
import { areLexicalDailyRollupsReady, rebuildLexicalDailyRollups } from './lexical-rollups';
|
||||||
|
import { Database } from './sqlite';
|
||||||
|
import { applyPragmas } from './storage';
|
||||||
|
|
||||||
|
export function executeLexicalRollupBackfillTask(dbPath: string): void {
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
try {
|
||||||
|
applyPragmas(db);
|
||||||
|
if (!areLexicalDailyRollupsReady(db)) {
|
||||||
|
rebuildLexicalDailyRollups(db);
|
||||||
|
}
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
}
|
||||||
|
}
|
||||||
@@ -0,0 +1,173 @@
|
|||||||
|
import assert from 'node:assert/strict';
|
||||||
|
import fs from 'node:fs';
|
||||||
|
import os from 'node:os';
|
||||||
|
import path from 'node:path';
|
||||||
|
import test from 'node:test';
|
||||||
|
import { getLexicalDailyRollups, rebuildLexicalDailyRollups } from './lexical-rollups';
|
||||||
|
import { getTrendsDashboard } from './query-trends';
|
||||||
|
import { getVocabularyChartData, replaceStatsExcludedWords } from './query-lexical';
|
||||||
|
import { Database } from './sqlite';
|
||||||
|
import type { DatabaseSync } from './sqlite';
|
||||||
|
import { ensureSchema } from './storage';
|
||||||
|
|
||||||
|
function makeDbPath(): string {
|
||||||
|
const dir = fs.mkdtempSync(path.join(os.tmpdir(), 'subminer-lexical-rollups-'));
|
||||||
|
return path.join(dir, 'immersion.sqlite');
|
||||||
|
}
|
||||||
|
|
||||||
|
test('lexical daily rollups follow first-seen corrections and deletions', () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
const firstDay = 19_500;
|
||||||
|
const correctedDay = firstDay + 2;
|
||||||
|
const firstSeen = firstDay * 86_400 + 43_200;
|
||||||
|
const correctedSeen = correctedDay * 86_400 + 43_200;
|
||||||
|
|
||||||
|
db.prepare(
|
||||||
|
`INSERT INTO imm_words(headword, word, reading, first_seen, last_seen, frequency)
|
||||||
|
VALUES (?, ?, ?, ?, ?, 1)`,
|
||||||
|
).run('猫', '猫', 'ねこ', firstSeen, firstSeen);
|
||||||
|
db.prepare(
|
||||||
|
`INSERT INTO imm_kanji(kanji, first_seen, last_seen, frequency)
|
||||||
|
VALUES (?, ?, ?, 1)`,
|
||||||
|
).run('猫', firstSeen, firstSeen);
|
||||||
|
|
||||||
|
assert.deepEqual(getLexicalDailyRollups(db), [
|
||||||
|
{ epochDay: firstDay, wordCount: 1, wordCountWithoutNames: 1, kanjiCount: 1 },
|
||||||
|
]);
|
||||||
|
|
||||||
|
db.prepare(`UPDATE imm_words SET first_seen = ? WHERE headword = ?`).run(correctedSeen, '猫');
|
||||||
|
db.prepare(`DELETE FROM imm_kanji WHERE kanji = ?`).run('猫');
|
||||||
|
|
||||||
|
assert.deepEqual(getLexicalDailyRollups(db), [
|
||||||
|
{ epochDay: correctedDay, wordCount: 1, wordCountWithoutNames: 1, kanjiCount: 0 },
|
||||||
|
]);
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
fs.rmSync(path.dirname(dbPath), { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('vocabulary charts use complete top-word and lexical rollup data', () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
const insertWord = db.prepare(
|
||||||
|
`INSERT INTO imm_words(headword, word, reading, first_seen, last_seen, frequency)
|
||||||
|
VALUES (?, ?, '', 1700000000, 1700000000, ?)`,
|
||||||
|
);
|
||||||
|
for (let index = 0; index < 501; index += 1) {
|
||||||
|
insertWord.run(`語${index}`, `語${index}`, index === 500 ? 10_000 : 1);
|
||||||
|
}
|
||||||
|
|
||||||
|
const charts = getVocabularyChartData(db);
|
||||||
|
|
||||||
|
assert.equal(charts.topWords[0]?.headword, '語500');
|
||||||
|
assert.equal(charts.topWords[0]?.frequency, 10_000);
|
||||||
|
assert.equal(charts.newWordsTimeline[0]?.wordCount, 501);
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
fs.rmSync(path.dirname(dbPath), { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('vocabulary charts find full top-word sets beyond excluded and name rows', () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
const insertWord = db.prepare(
|
||||||
|
`INSERT INTO imm_words(headword, word, reading, pos2, first_seen, last_seen, frequency)
|
||||||
|
VALUES (?, ?, '', ?, 1700000000, 1700000000, ?)`,
|
||||||
|
);
|
||||||
|
const exclusions = [];
|
||||||
|
for (let index = 0; index < 100; index += 1) {
|
||||||
|
const headword = `語${index}`;
|
||||||
|
insertWord.run(
|
||||||
|
headword,
|
||||||
|
headword,
|
||||||
|
index < 80 && index >= 60 ? '固有名詞' : '一般',
|
||||||
|
100 - index,
|
||||||
|
);
|
||||||
|
if (index < 60) exclusions.push({ headword, word: headword, reading: '' });
|
||||||
|
}
|
||||||
|
replaceStatsExcludedWords(db, exclusions);
|
||||||
|
|
||||||
|
const charts = getVocabularyChartData(db);
|
||||||
|
|
||||||
|
assert.equal(charts.topWords.length, 12);
|
||||||
|
assert.equal(charts.topWords[0]?.headword, '語60');
|
||||||
|
assert.equal(charts.topWordsWithoutNames.length, 12);
|
||||||
|
assert.equal(charts.topWordsWithoutNames[0]?.headword, '語80');
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
fs.rmSync(path.dirname(dbPath), { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('vocabulary charts handle exclusion lists above one SQLite variable batch', () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
db.prepare(
|
||||||
|
`INSERT INTO imm_words(headword, word, reading, first_seen, last_seen, frequency)
|
||||||
|
VALUES ('語0', '語0', '', 1700000000, 1700000000, 1)`,
|
||||||
|
).run();
|
||||||
|
const exclusions = Array.from({ length: 10_923 }, (_, index) => ({
|
||||||
|
headword: `語${index}`,
|
||||||
|
word: `語${index}`,
|
||||||
|
reading: '',
|
||||||
|
}));
|
||||||
|
replaceStatsExcludedWords(db, exclusions);
|
||||||
|
|
||||||
|
const charts = getVocabularyChartData(db);
|
||||||
|
|
||||||
|
assert.deepEqual(charts.topWords, []);
|
||||||
|
assert.deepEqual(charts.newWordsTimeline, []);
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
fs.rmSync(path.dirname(dbPath), { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('lexical rollup rebuild preserves the original error when rollback also fails', () => {
|
||||||
|
const originalError = new Error('rebuild failed');
|
||||||
|
const db = {
|
||||||
|
exec(sql: string) {
|
||||||
|
if (sql === 'BEGIN IMMEDIATE') return;
|
||||||
|
if (sql === 'ROLLBACK') throw new Error('rollback failed');
|
||||||
|
throw originalError;
|
||||||
|
},
|
||||||
|
} as unknown as DatabaseSync;
|
||||||
|
|
||||||
|
assert.throws(() => rebuildLexicalDailyRollups(db), originalError);
|
||||||
|
});
|
||||||
|
|
||||||
|
test('trends read historical new-word buckets from lexical rollups', () => {
|
||||||
|
const dbPath = makeDbPath();
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
ensureSchema(db);
|
||||||
|
db.prepare(
|
||||||
|
`INSERT INTO imm_words(headword, word, reading, first_seen, last_seen, frequency)
|
||||||
|
VALUES ('海', '海', 'うみ', 1700000000, 1700000000, 1)`,
|
||||||
|
).run();
|
||||||
|
db.prepare(`UPDATE imm_lexical_daily_rollups SET word_count = 9`).run();
|
||||||
|
|
||||||
|
const dashboard = getTrendsDashboard(db, 'all', 'day', false);
|
||||||
|
|
||||||
|
assert.equal(dashboard.progress.newWords[0]?.value, 9);
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
fs.rmSync(path.dirname(dbPath), { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
@@ -0,0 +1,178 @@
|
|||||||
|
import type { DatabaseSync } from './sqlite';
|
||||||
|
|
||||||
|
export interface LexicalDailyRollup {
|
||||||
|
epochDay: number;
|
||||||
|
wordCount: number;
|
||||||
|
wordCountWithoutNames: number;
|
||||||
|
kanjiCount: number;
|
||||||
|
}
|
||||||
|
|
||||||
|
const LOCAL_EPOCH_DAY_SQL = `
|
||||||
|
CAST(julianday(CAST(%VALUE% AS REAL), 'unixepoch', 'localtime') - 2440587.5 AS INTEGER)
|
||||||
|
`;
|
||||||
|
|
||||||
|
export function localEpochDaySql(value: string): string {
|
||||||
|
return LOCAL_EPOCH_DAY_SQL.replace('%VALUE%', value);
|
||||||
|
}
|
||||||
|
|
||||||
|
function createWordRollupTriggers(db: DatabaseSync): void {
|
||||||
|
const dayForNew = localEpochDaySql('NEW.first_seen');
|
||||||
|
const dayForOld = localEpochDaySql('OLD.first_seen');
|
||||||
|
|
||||||
|
db.exec(`
|
||||||
|
CREATE TRIGGER IF NOT EXISTS imm_words_lexical_rollup_insert
|
||||||
|
AFTER INSERT ON imm_words
|
||||||
|
WHEN NEW.first_seen IS NOT NULL
|
||||||
|
BEGIN
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
VALUES (${dayForNew}, 1, CASE WHEN NEW.pos2 = '固有名詞' THEN 0 ELSE 1 END, 0)
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET
|
||||||
|
word_count = word_count + 1,
|
||||||
|
word_count_without_names = word_count_without_names + excluded.word_count_without_names;
|
||||||
|
END;
|
||||||
|
|
||||||
|
CREATE TRIGGER IF NOT EXISTS imm_words_lexical_rollup_delete
|
||||||
|
AFTER DELETE ON imm_words
|
||||||
|
WHEN OLD.first_seen IS NOT NULL
|
||||||
|
BEGIN
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
VALUES (${dayForOld}, -1, CASE WHEN OLD.pos2 = '固有名詞' THEN 0 ELSE -1 END, 0)
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET
|
||||||
|
word_count = word_count - 1,
|
||||||
|
word_count_without_names = word_count_without_names + excluded.word_count_without_names;
|
||||||
|
DELETE FROM imm_lexical_daily_rollups
|
||||||
|
WHERE epoch_day = ${dayForOld} AND word_count = 0 AND kanji_count = 0;
|
||||||
|
END;
|
||||||
|
|
||||||
|
CREATE TRIGGER IF NOT EXISTS imm_words_lexical_rollup_first_seen_update
|
||||||
|
AFTER UPDATE OF first_seen, pos2 ON imm_words
|
||||||
|
WHEN OLD.first_seen IS NOT NEW.first_seen OR OLD.pos2 IS NOT NEW.pos2
|
||||||
|
BEGIN
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
SELECT ${dayForOld}, -1, CASE WHEN OLD.pos2 = '固有名詞' THEN 0 ELSE -1 END, 0
|
||||||
|
WHERE OLD.first_seen IS NOT NULL
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET
|
||||||
|
word_count = word_count - 1,
|
||||||
|
word_count_without_names = word_count_without_names + excluded.word_count_without_names;
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
SELECT ${dayForNew}, 1, CASE WHEN NEW.pos2 = '固有名詞' THEN 0 ELSE 1 END, 0
|
||||||
|
WHERE NEW.first_seen IS NOT NULL
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET
|
||||||
|
word_count = word_count + 1,
|
||||||
|
word_count_without_names = word_count_without_names + excluded.word_count_without_names;
|
||||||
|
DELETE FROM imm_lexical_daily_rollups
|
||||||
|
WHERE word_count = 0 AND kanji_count = 0;
|
||||||
|
END;
|
||||||
|
`);
|
||||||
|
}
|
||||||
|
|
||||||
|
function createKanjiRollupTriggers(db: DatabaseSync): void {
|
||||||
|
const dayForNew = localEpochDaySql('NEW.first_seen');
|
||||||
|
const dayForOld = localEpochDaySql('OLD.first_seen');
|
||||||
|
db.exec(`
|
||||||
|
CREATE TRIGGER IF NOT EXISTS imm_kanji_lexical_rollup_insert
|
||||||
|
AFTER INSERT ON imm_kanji WHEN NEW.first_seen IS NOT NULL
|
||||||
|
BEGIN
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
VALUES (${dayForNew}, 0, 0, 1)
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET kanji_count = kanji_count + 1;
|
||||||
|
END;
|
||||||
|
CREATE TRIGGER IF NOT EXISTS imm_kanji_lexical_rollup_delete
|
||||||
|
AFTER DELETE ON imm_kanji WHEN OLD.first_seen IS NOT NULL
|
||||||
|
BEGIN
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
VALUES (${dayForOld}, 0, 0, -1)
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET kanji_count = kanji_count - 1;
|
||||||
|
DELETE FROM imm_lexical_daily_rollups
|
||||||
|
WHERE epoch_day = ${dayForOld} AND word_count = 0 AND kanji_count = 0;
|
||||||
|
END;
|
||||||
|
CREATE TRIGGER IF NOT EXISTS imm_kanji_lexical_rollup_first_seen_update
|
||||||
|
AFTER UPDATE OF first_seen ON imm_kanji WHEN OLD.first_seen IS NOT NEW.first_seen
|
||||||
|
BEGIN
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
SELECT ${dayForOld}, 0, 0, -1 WHERE OLD.first_seen IS NOT NULL
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET kanji_count = kanji_count - 1;
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
SELECT ${dayForNew}, 0, 0, 1 WHERE NEW.first_seen IS NOT NULL
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET kanji_count = kanji_count + 1;
|
||||||
|
DELETE FROM imm_lexical_daily_rollups WHERE word_count = 0 AND kanji_count = 0;
|
||||||
|
END;
|
||||||
|
`);
|
||||||
|
}
|
||||||
|
|
||||||
|
export function ensureLexicalDailyRollupTables(db: DatabaseSync): void {
|
||||||
|
db.exec(`
|
||||||
|
CREATE TABLE IF NOT EXISTS imm_lexical_daily_rollups(
|
||||||
|
epoch_day INTEGER PRIMARY KEY,
|
||||||
|
word_count INTEGER NOT NULL DEFAULT 0,
|
||||||
|
word_count_without_names INTEGER NOT NULL DEFAULT 0,
|
||||||
|
kanji_count INTEGER NOT NULL DEFAULT 0
|
||||||
|
);
|
||||||
|
INSERT INTO imm_rollup_state(state_key, state_value)
|
||||||
|
VALUES ('lexical_daily_rollups_ready', '0')
|
||||||
|
ON CONFLICT(state_key) DO NOTHING;
|
||||||
|
`);
|
||||||
|
createWordRollupTriggers(db);
|
||||||
|
createKanjiRollupTriggers(db);
|
||||||
|
}
|
||||||
|
|
||||||
|
export function areLexicalDailyRollupsReady(db: DatabaseSync): boolean {
|
||||||
|
const row = db
|
||||||
|
.prepare(`SELECT state_value AS value FROM imm_rollup_state WHERE state_key = ?`)
|
||||||
|
.get('lexical_daily_rollups_ready') as { value: string } | null;
|
||||||
|
return row?.value === '1';
|
||||||
|
}
|
||||||
|
|
||||||
|
export function markLexicalDailyRollupsReady(db: DatabaseSync): void {
|
||||||
|
db.prepare(`UPDATE imm_rollup_state SET state_value = '1' WHERE state_key = ?`).run(
|
||||||
|
'lexical_daily_rollups_ready',
|
||||||
|
);
|
||||||
|
}
|
||||||
|
|
||||||
|
/** Rebuild from the first-seen source of truth; run off the UI/main DB thread. */
|
||||||
|
export function rebuildLexicalDailyRollups(db: DatabaseSync): void {
|
||||||
|
let transactionStarted = false;
|
||||||
|
try {
|
||||||
|
db.exec('BEGIN IMMEDIATE');
|
||||||
|
transactionStarted = true;
|
||||||
|
db.exec('DELETE FROM imm_lexical_daily_rollups');
|
||||||
|
db.exec(`
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
SELECT ${localEpochDaySql('first_seen')}, COUNT(*),
|
||||||
|
SUM(CASE WHEN pos2 = '固有名詞' THEN 0 ELSE 1 END), 0
|
||||||
|
FROM imm_words
|
||||||
|
WHERE first_seen IS NOT NULL
|
||||||
|
GROUP BY ${localEpochDaySql('first_seen')};
|
||||||
|
INSERT INTO imm_lexical_daily_rollups(epoch_day, word_count, word_count_without_names, kanji_count)
|
||||||
|
SELECT ${localEpochDaySql('first_seen')}, 0, 0, COUNT(*)
|
||||||
|
FROM imm_kanji
|
||||||
|
WHERE first_seen IS NOT NULL
|
||||||
|
GROUP BY ${localEpochDaySql('first_seen')}
|
||||||
|
ON CONFLICT(epoch_day) DO UPDATE SET kanji_count = kanji_count + excluded.kanji_count;
|
||||||
|
`);
|
||||||
|
markLexicalDailyRollupsReady(db);
|
||||||
|
db.exec('COMMIT');
|
||||||
|
} catch (error) {
|
||||||
|
if (transactionStarted) {
|
||||||
|
try {
|
||||||
|
db.exec('ROLLBACK');
|
||||||
|
} catch {
|
||||||
|
// Preserve the rebuild failure; it is the actionable cause.
|
||||||
|
}
|
||||||
|
}
|
||||||
|
throw error;
|
||||||
|
}
|
||||||
|
}
|
||||||
|
|
||||||
|
export function getLexicalDailyRollups(db: DatabaseSync): LexicalDailyRollup[] {
|
||||||
|
return db
|
||||||
|
.prepare(
|
||||||
|
`
|
||||||
|
SELECT epoch_day AS epochDay, word_count AS wordCount,
|
||||||
|
word_count_without_names AS wordCountWithoutNames, kanji_count AS kanjiCount
|
||||||
|
FROM imm_lexical_daily_rollups
|
||||||
|
ORDER BY epoch_day ASC
|
||||||
|
`,
|
||||||
|
)
|
||||||
|
.all() as LexicalDailyRollup[];
|
||||||
|
}
|
||||||
@@ -13,19 +13,36 @@ import type {
|
|||||||
SimilarWordRow,
|
SimilarWordRow,
|
||||||
StatsExcludedWordRow,
|
StatsExcludedWordRow,
|
||||||
VocabularyStatsRow,
|
VocabularyStatsRow,
|
||||||
|
VocabularyStatsSummary,
|
||||||
WordAnimeAppearanceRow,
|
WordAnimeAppearanceRow,
|
||||||
WordDetailRow,
|
WordDetailRow,
|
||||||
WordOccurrenceRow,
|
WordOccurrenceRow,
|
||||||
} from './types';
|
} from './types';
|
||||||
import { fromDbTimestamp, toDbTimestamp } from './query-shared';
|
import { fromDbTimestamp, toDbTimestamp } from './query-shared';
|
||||||
import { nowMs } from './time';
|
import { nowMs } from './time';
|
||||||
|
import {
|
||||||
|
areLexicalDailyRollupsReady,
|
||||||
|
getLexicalDailyRollups,
|
||||||
|
localEpochDaySql,
|
||||||
|
} from './lexical-rollups';
|
||||||
|
|
||||||
const VOCABULARY_STATS_FILTER_OVERSAMPLE_FACTOR = 4;
|
const VOCABULARY_STATS_FILTER_OVERSAMPLE_FACTOR = 4;
|
||||||
const VOCABULARY_STATS_FILTER_OVERSAMPLE_MIN = 100;
|
const VOCABULARY_STATS_FILTER_OVERSAMPLE_MIN = 100;
|
||||||
|
const VOCABULARY_CHART_LIMIT = 12;
|
||||||
|
const VOCABULARY_CHART_PAGE_SIZE = 100;
|
||||||
|
const EXCLUSION_ALIAS_BATCH_SIZE = 300;
|
||||||
const SENTENCE_SEARCH_DEFAULT_LIMIT = 50;
|
const SENTENCE_SEARCH_DEFAULT_LIMIT = 50;
|
||||||
const SENTENCE_SEARCH_MAX_LIMIT = 100;
|
const SENTENCE_SEARCH_MAX_LIMIT = 100;
|
||||||
const KANJI_PATTERN = /\p{Script=Han}/gu;
|
const KANJI_PATTERN = /\p{Script=Han}/gu;
|
||||||
|
|
||||||
|
export interface VocabularyChartData {
|
||||||
|
ready: boolean;
|
||||||
|
topWords: Array<{ wordId: number; headword: string; frequency: number }>;
|
||||||
|
topWordsWithoutNames: Array<{ wordId: number; headword: string; frequency: number }>;
|
||||||
|
newWordsTimeline: Array<{ epochDay: number; wordCount: number }>;
|
||||||
|
newWordsTimelineWithoutNames: Array<{ epochDay: number; wordCount: number }>;
|
||||||
|
}
|
||||||
|
|
||||||
function resolveSentenceSearchLimit(limit: number): number {
|
function resolveSentenceSearchLimit(limit: number): number {
|
||||||
if (!Number.isFinite(limit)) return SENTENCE_SEARCH_DEFAULT_LIMIT;
|
if (!Number.isFinite(limit)) return SENTENCE_SEARCH_DEFAULT_LIMIT;
|
||||||
const normalized = Math.floor(limit);
|
const normalized = Math.floor(limit);
|
||||||
@@ -153,6 +170,182 @@ export function getVocabularyStats(
|
|||||||
return visibleRows.slice(0, limit);
|
return visibleRows.slice(0, limit);
|
||||||
}
|
}
|
||||||
|
|
||||||
|
/**
|
||||||
|
* Chart data is intentionally independent of the paginated vocabulary tables.
|
||||||
|
* Top words use the frequency index; new-word history reads permanent daily
|
||||||
|
* lexical rollups rather than loading every vocabulary row into the dashboard.
|
||||||
|
*/
|
||||||
|
export function getVocabularyChartData(db: DatabaseSync): VocabularyChartData {
|
||||||
|
const ready = areLexicalDailyRollupsReady(db);
|
||||||
|
const excludedAliases = new Set(
|
||||||
|
getStatsExcludedWords(db).flatMap((word) => excludedVocabularyAliases(word)),
|
||||||
|
);
|
||||||
|
const isExcluded = (word: Pick<VocabularyStatsRow, 'headword' | 'word' | 'reading'>): boolean =>
|
||||||
|
excludedVocabularyAliases(word).some((alias) => excludedAliases.has(alias));
|
||||||
|
const topWords = getTopVocabularyChartWords(db, isExcluded);
|
||||||
|
const rollups = ready ? getLexicalDailyRollups(db) : [];
|
||||||
|
const timeline = new Map(rollups.map((row) => [row.epochDay, { ...row }]));
|
||||||
|
if (excludedAliases.size > 0 && ready) {
|
||||||
|
const aliases = [...excludedAliases];
|
||||||
|
const excludedRows = new Map<
|
||||||
|
number,
|
||||||
|
Pick<VocabularyStatsRow, 'headword' | 'word' | 'reading' | 'pos2'> & {
|
||||||
|
wordId: number;
|
||||||
|
epochDay: number;
|
||||||
|
}
|
||||||
|
>();
|
||||||
|
for (let offset = 0; offset < aliases.length; offset += EXCLUSION_ALIAS_BATCH_SIZE) {
|
||||||
|
const batch = aliases.slice(offset, offset + EXCLUSION_ALIAS_BATCH_SIZE);
|
||||||
|
const placeholders = batch.map(() => '?').join(', ');
|
||||||
|
const rows = db
|
||||||
|
.prepare(
|
||||||
|
`
|
||||||
|
SELECT id AS wordId, headword, word, reading, pos2,
|
||||||
|
${localEpochDaySql('first_seen')} AS epochDay
|
||||||
|
FROM imm_words
|
||||||
|
WHERE headword IN (${placeholders}) OR word IN (${placeholders}) OR reading IN (${placeholders})
|
||||||
|
`,
|
||||||
|
)
|
||||||
|
.all(...batch, ...batch, ...batch) as Array<
|
||||||
|
Pick<VocabularyStatsRow, 'headword' | 'word' | 'reading' | 'pos2'> & {
|
||||||
|
wordId: number;
|
||||||
|
epochDay: number;
|
||||||
|
}
|
||||||
|
>;
|
||||||
|
for (const row of rows) excludedRows.set(row.wordId, row);
|
||||||
|
}
|
||||||
|
for (const word of excludedRows.values()) {
|
||||||
|
if (!isExcluded(word)) continue;
|
||||||
|
const rollup = timeline.get(word.epochDay);
|
||||||
|
if (!rollup) continue;
|
||||||
|
rollup.wordCount -= 1;
|
||||||
|
if (word.pos2 !== '固有名詞') rollup.wordCountWithoutNames -= 1;
|
||||||
|
}
|
||||||
|
}
|
||||||
|
return {
|
||||||
|
ready,
|
||||||
|
topWords: topWords.all.map((word) => ({
|
||||||
|
wordId: word.wordId,
|
||||||
|
headword: word.headword,
|
||||||
|
frequency: word.frequency,
|
||||||
|
})),
|
||||||
|
topWordsWithoutNames: topWords.withoutNames.map((word) => ({
|
||||||
|
wordId: word.wordId,
|
||||||
|
headword: word.headword,
|
||||||
|
frequency: word.frequency,
|
||||||
|
})),
|
||||||
|
newWordsTimeline: [...timeline.values()]
|
||||||
|
.filter((row) => row.wordCount > 0)
|
||||||
|
.map((row) => ({ epochDay: row.epochDay, wordCount: row.wordCount })),
|
||||||
|
newWordsTimelineWithoutNames: [...timeline.values()]
|
||||||
|
.filter((row) => row.wordCountWithoutNames > 0)
|
||||||
|
.map((row) => ({ epochDay: row.epochDay, wordCount: row.wordCountWithoutNames })),
|
||||||
|
};
|
||||||
|
}
|
||||||
|
|
||||||
|
function getTopVocabularyChartWords(
|
||||||
|
db: DatabaseSync,
|
||||||
|
isExcluded: (word: Pick<VocabularyStatsRow, 'headword' | 'word' | 'reading'>) => boolean,
|
||||||
|
): { all: VocabularyStatsRow[]; withoutNames: VocabularyStatsRow[] } {
|
||||||
|
const stmt = db.prepare(`
|
||||||
|
SELECT id AS wordId, headword, word, reading,
|
||||||
|
part_of_speech AS partOfSpeech, pos1, pos2, pos3,
|
||||||
|
frequency, frequency_rank AS frequencyRank,
|
||||||
|
first_seen AS firstSeen, last_seen AS lastSeen,
|
||||||
|
0 AS animeCount
|
||||||
|
FROM imm_words
|
||||||
|
ORDER BY frequency DESC, id
|
||||||
|
LIMIT ? OFFSET ?
|
||||||
|
`);
|
||||||
|
const all: VocabularyStatsRow[] = [];
|
||||||
|
const withoutNames: VocabularyStatsRow[] = [];
|
||||||
|
let offset = 0;
|
||||||
|
|
||||||
|
while (all.length < VOCABULARY_CHART_LIMIT || withoutNames.length < VOCABULARY_CHART_LIMIT) {
|
||||||
|
const page = stmt.all(VOCABULARY_CHART_PAGE_SIZE, offset) as VocabularyStatsRow[];
|
||||||
|
if (page.length === 0) break;
|
||||||
|
for (const word of page) {
|
||||||
|
if (!isVocabularyStatsRowVisible(word) || isExcluded(word)) continue;
|
||||||
|
if (all.length < VOCABULARY_CHART_LIMIT) all.push(word);
|
||||||
|
if (word.pos2 !== '固有名詞' && withoutNames.length < VOCABULARY_CHART_LIMIT) {
|
||||||
|
withoutNames.push(word);
|
||||||
|
}
|
||||||
|
}
|
||||||
|
offset += page.length;
|
||||||
|
}
|
||||||
|
|
||||||
|
return { all, withoutNames };
|
||||||
|
}
|
||||||
|
|
||||||
|
function excludedVocabularyAliases(
|
||||||
|
word: Pick<VocabularyStatsRow, 'headword' | 'word' | 'reading'>,
|
||||||
|
): string[] {
|
||||||
|
const aliases = [word.headword.trim(), word.word.trim()].filter(Boolean);
|
||||||
|
if (aliases.length === 0) aliases.push(word.reading.trim());
|
||||||
|
return [...new Set(aliases)];
|
||||||
|
}
|
||||||
|
|
||||||
|
function timestampSeconds(timestamp: number): number {
|
||||||
|
return timestamp < 10_000_000_000 ? timestamp : Math.floor(timestamp / 1000);
|
||||||
|
}
|
||||||
|
|
||||||
|
export function getVocabularySummary(
|
||||||
|
db: DatabaseSync,
|
||||||
|
knownWords: ReadonlySet<string> | null,
|
||||||
|
nowMs: number = Date.now(),
|
||||||
|
): VocabularyStatsSummary {
|
||||||
|
const words = db
|
||||||
|
.prepare(
|
||||||
|
`
|
||||||
|
SELECT id AS wordId, headword, word, reading,
|
||||||
|
part_of_speech AS partOfSpeech, pos1, pos2, pos3,
|
||||||
|
frequency, frequency_rank AS frequencyRank,
|
||||||
|
first_seen AS firstSeen, last_seen AS lastSeen,
|
||||||
|
0 AS animeCount
|
||||||
|
FROM imm_words
|
||||||
|
`,
|
||||||
|
)
|
||||||
|
.all() as VocabularyStatsRow[];
|
||||||
|
const excludedAliases = new Set(
|
||||||
|
getStatsExcludedWords(db).flatMap((word) => excludedVocabularyAliases(word)),
|
||||||
|
);
|
||||||
|
const weekAgoSec = nowMs / 1000 - 7 * 86_400;
|
||||||
|
const summary: VocabularyStatsSummary = {
|
||||||
|
uniqueWords: 0,
|
||||||
|
uniqueWordsWithoutNames: 0,
|
||||||
|
uniqueKanji: (db.prepare('SELECT COUNT(*) AS count FROM imm_kanji').get() as { count: number })
|
||||||
|
.count,
|
||||||
|
newThisWeek: 0,
|
||||||
|
newThisWeekWithoutNames: 0,
|
||||||
|
knownWordCount: knownWords ? 0 : null,
|
||||||
|
knownWordCountWithoutNames: knownWords ? 0 : null,
|
||||||
|
};
|
||||||
|
|
||||||
|
for (const word of words) {
|
||||||
|
if (
|
||||||
|
!isVocabularyStatsRowVisible(word) ||
|
||||||
|
excludedVocabularyAliases(word).some((alias) => excludedAliases.has(alias))
|
||||||
|
) {
|
||||||
|
continue;
|
||||||
|
}
|
||||||
|
const isName = word.pos2 === '固有名詞';
|
||||||
|
const isNewThisWeek = timestampSeconds(fromDbTimestamp(word.firstSeen) ?? 0) >= weekAgoSec;
|
||||||
|
const isKnown = knownWords?.has(word.headword) ?? false;
|
||||||
|
summary.uniqueWords += 1;
|
||||||
|
if (!isName) summary.uniqueWordsWithoutNames += 1;
|
||||||
|
if (isNewThisWeek) {
|
||||||
|
summary.newThisWeek += 1;
|
||||||
|
if (!isName) summary.newThisWeekWithoutNames += 1;
|
||||||
|
}
|
||||||
|
if (isKnown) {
|
||||||
|
summary.knownWordCount! += 1;
|
||||||
|
if (!isName) summary.knownWordCountWithoutNames! += 1;
|
||||||
|
}
|
||||||
|
}
|
||||||
|
|
||||||
|
return summary;
|
||||||
|
}
|
||||||
|
|
||||||
export function getStatsExcludedWords(db: DatabaseSync): StatsExcludedWordRow[] {
|
export function getStatsExcludedWords(db: DatabaseSync): StatsExcludedWordRow[] {
|
||||||
return db
|
return db
|
||||||
.prepare(
|
.prepare(
|
||||||
|
|||||||
@@ -13,6 +13,7 @@ import {
|
|||||||
toDbTimestamp,
|
toDbTimestamp,
|
||||||
} from './query-shared';
|
} from './query-shared';
|
||||||
import { getDailyRollups, getMonthlyRollups } from './query-sessions';
|
import { getDailyRollups, getMonthlyRollups } from './query-sessions';
|
||||||
|
import { areLexicalDailyRollupsReady, getLexicalDailyRollups } from './lexical-rollups';
|
||||||
|
|
||||||
type TrendRange = '7d' | '30d' | '90d' | '365d' | 'all';
|
type TrendRange = '7d' | '30d' | '90d' | '365d' | 'all';
|
||||||
type TrendGroupBy = 'day' | 'month';
|
type TrendGroupBy = 'day' | 'month';
|
||||||
@@ -660,6 +661,16 @@ function buildNewWordsPerDay(
|
|||||||
cutoffMs: string | null,
|
cutoffMs: string | null,
|
||||||
axis: number[] | null,
|
axis: number[] | null,
|
||||||
): TrendChartPoint[] {
|
): TrendChartPoint[] {
|
||||||
|
if (areLexicalDailyRollupsReady(db)) {
|
||||||
|
// A trend range is defined in calendar buckets, so the rollup includes the
|
||||||
|
// complete local cutoff day rather than applying a time-of-day boundary.
|
||||||
|
const cutoffDay = cutoffMs === null ? null : getLocalEpochDay(db, cutoffMs);
|
||||||
|
const rows = getLexicalDailyRollups(db).filter(
|
||||||
|
(row) => cutoffDay === null || row.epochDay >= cutoffDay,
|
||||||
|
);
|
||||||
|
return fillAxisPoints(axis, new Map(rows.map((row) => [row.epochDay, row.wordCount])));
|
||||||
|
}
|
||||||
|
|
||||||
const whereClause = cutoffMs === null ? '' : 'AND first_seen >= ?';
|
const whereClause = cutoffMs === null ? '' : 'AND first_seen >= ?';
|
||||||
const prepared = db.prepare(`
|
const prepared = db.prepare(`
|
||||||
SELECT
|
SELECT
|
||||||
@@ -691,6 +702,18 @@ function buildNewWordsPerMonth(
|
|||||||
cutoffMs: string | null,
|
cutoffMs: string | null,
|
||||||
axis: number[] | null,
|
axis: number[] | null,
|
||||||
): TrendChartPoint[] {
|
): TrendChartPoint[] {
|
||||||
|
if (areLexicalDailyRollupsReady(db)) {
|
||||||
|
const cutoffDay = cutoffMs === null ? null : getLocalEpochDay(db, cutoffMs);
|
||||||
|
const byMonth = new Map<number, number>();
|
||||||
|
for (const row of getLexicalDailyRollups(db)) {
|
||||||
|
if (cutoffDay !== null && row.epochDay < cutoffDay) continue;
|
||||||
|
const { year, month } = dayPartsFromEpochDay(row.epochDay);
|
||||||
|
const monthKey = year * 100 + month;
|
||||||
|
byMonth.set(monthKey, (byMonth.get(monthKey) ?? 0) + row.wordCount);
|
||||||
|
}
|
||||||
|
return fillAxisPoints(axis, byMonth);
|
||||||
|
}
|
||||||
|
|
||||||
const whereClause = cutoffMs === null ? '' : 'AND first_seen >= ?';
|
const whereClause = cutoffMs === null ? '' : 'AND first_seen >= ?';
|
||||||
const prepared = db.prepare(`
|
const prepared = db.prepare(`
|
||||||
SELECT
|
SELECT
|
||||||
|
|||||||
@@ -4,6 +4,7 @@ import { parseMediaInfo } from '../../../jimaku/utils';
|
|||||||
import { normalizeTitleIdentity } from '../../utils/title-normalization';
|
import { normalizeTitleIdentity } from '../../utils/title-normalization';
|
||||||
import type { DatabaseSync } from './sqlite';
|
import type { DatabaseSync } from './sqlite';
|
||||||
import { nowMs } from './time';
|
import { nowMs } from './time';
|
||||||
|
import { ensureLexicalDailyRollupTables, markLexicalDailyRollupsReady } from './lexical-rollups';
|
||||||
import { SCHEMA_VERSION } from './types';
|
import { SCHEMA_VERSION } from './types';
|
||||||
import type { QueuedWrite, VideoMetadata, YoutubeVideoMetadata } from './types';
|
import type { QueuedWrite, VideoMetadata, YoutubeVideoMetadata } from './types';
|
||||||
import { toDbMs, toDbTimestamp } from './query-shared';
|
import { toDbMs, toDbTimestamp } from './query-shared';
|
||||||
@@ -890,11 +891,11 @@ export function ensureSchema(db: DatabaseSync): void {
|
|||||||
VALUES ('last_rollup_sample_ms', 0)
|
VALUES ('last_rollup_sample_ms', 0)
|
||||||
ON CONFLICT(state_key) DO NOTHING
|
ON CONFLICT(state_key) DO NOTHING
|
||||||
`);
|
`);
|
||||||
|
|
||||||
const currentVersion = db
|
const currentVersion = db
|
||||||
.prepare('SELECT schema_version FROM imm_schema_version ORDER BY schema_version DESC LIMIT 1')
|
.prepare('SELECT schema_version FROM imm_schema_version ORDER BY schema_version DESC LIMIT 1')
|
||||||
.get() as { schema_version: number } | null;
|
.get() as { schema_version: number } | null;
|
||||||
if (currentVersion?.schema_version === SCHEMA_VERSION) {
|
if (currentVersion?.schema_version === SCHEMA_VERSION) {
|
||||||
|
ensureLexicalDailyRollupTables(db);
|
||||||
ensureLifetimeSummaryTables(db);
|
ensureLifetimeSummaryTables(db);
|
||||||
ensureStatsExcludedWordsTable(db);
|
ensureStatsExcludedWordsTable(db);
|
||||||
ensureAnimeMergeTables(db);
|
ensureAnimeMergeTables(db);
|
||||||
@@ -1453,6 +1454,7 @@ export function ensureSchema(db: DatabaseSync): void {
|
|||||||
|
|
||||||
migrateSessionEventTimestampsToText(db);
|
migrateSessionEventTimestampsToText(db);
|
||||||
|
|
||||||
|
ensureLexicalDailyRollupTables(db);
|
||||||
ensureLifetimeSummaryTables(db);
|
ensureLifetimeSummaryTables(db);
|
||||||
ensureStatsExcludedWordsTable(db);
|
ensureStatsExcludedWordsTable(db);
|
||||||
|
|
||||||
@@ -1585,6 +1587,12 @@ export function ensureSchema(db: DatabaseSync): void {
|
|||||||
VALUES (${SCHEMA_VERSION}, ${toDbTimestamp(nowMs())})
|
VALUES (${SCHEMA_VERSION}, ${toDbTimestamp(nowMs())})
|
||||||
ON CONFLICT DO NOTHING
|
ON CONFLICT DO NOTHING
|
||||||
`);
|
`);
|
||||||
|
|
||||||
|
// A new database has no history to materialize. Upgrades are populated by the
|
||||||
|
// background worker so startup never scans the existing vocabulary table.
|
||||||
|
if (!currentVersion) {
|
||||||
|
markLexicalDailyRollupsReady(db);
|
||||||
|
}
|
||||||
}
|
}
|
||||||
|
|
||||||
export function createTrackerPreparedStatements(db: DatabaseSync): TrackerPreparedStatements {
|
export function createTrackerPreparedStatements(db: DatabaseSync): TrackerPreparedStatements {
|
||||||
|
|||||||
@@ -1,4 +1,4 @@
|
|||||||
export const SCHEMA_VERSION = 21;
|
export const SCHEMA_VERSION = 22;
|
||||||
export const DEFAULT_QUEUE_CAP = 1_000;
|
export const DEFAULT_QUEUE_CAP = 1_000;
|
||||||
export const DEFAULT_BATCH_SIZE = 25;
|
export const DEFAULT_BATCH_SIZE = 25;
|
||||||
export const DEFAULT_FLUSH_INTERVAL_MS = 500;
|
export const DEFAULT_FLUSH_INTERVAL_MS = 500;
|
||||||
@@ -306,6 +306,16 @@ export interface VocabularyStatsRow {
|
|||||||
lastSeen: number;
|
lastSeen: number;
|
||||||
}
|
}
|
||||||
|
|
||||||
|
export interface VocabularyStatsSummary {
|
||||||
|
uniqueWords: number;
|
||||||
|
uniqueWordsWithoutNames: number;
|
||||||
|
uniqueKanji: number;
|
||||||
|
newThisWeek: number;
|
||||||
|
newThisWeekWithoutNames: number;
|
||||||
|
knownWordCount: number | null;
|
||||||
|
knownWordCountWithoutNames: number | null;
|
||||||
|
}
|
||||||
|
|
||||||
export interface StatsExcludedWordRow {
|
export interface StatsExcludedWordRow {
|
||||||
headword: string;
|
headword: string;
|
||||||
word: string;
|
word: string;
|
||||||
|
|||||||
@@ -0,0 +1,67 @@
|
|||||||
|
import assert from 'node:assert/strict';
|
||||||
|
import fs from 'node:fs';
|
||||||
|
import os from 'node:os';
|
||||||
|
import path from 'node:path';
|
||||||
|
import test from 'node:test';
|
||||||
|
import {
|
||||||
|
resolveVocabularySummaryWorkerPath,
|
||||||
|
VocabularySummaryWorkerRuntime,
|
||||||
|
} from './vocabulary-summary-worker-runtime';
|
||||||
|
import { Database } from './sqlite';
|
||||||
|
import { applyPragmas, ensureSchema } from './storage';
|
||||||
|
|
||||||
|
test('vocabulary summary worker reads the database from a separate connection', async () => {
|
||||||
|
const tempDir = fs.mkdtempSync(path.join(os.tmpdir(), 'subminer-vocabulary-summary-worker-'));
|
||||||
|
const dbPath = path.join(tempDir, 'immersion.sqlite');
|
||||||
|
const runtime = new VocabularySummaryWorkerRuntime();
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
|
||||||
|
try {
|
||||||
|
applyPragmas(db);
|
||||||
|
ensureSchema(db);
|
||||||
|
db.prepare(
|
||||||
|
`
|
||||||
|
INSERT INTO imm_words (
|
||||||
|
headword, word, reading, part_of_speech, pos1, pos2, pos3,
|
||||||
|
first_seen, last_seen, frequency
|
||||||
|
) VALUES ('猫', '猫', 'ねこ', 'noun', '名詞', '一般', '', 1, 1, 1)
|
||||||
|
`,
|
||||||
|
).run();
|
||||||
|
db.close();
|
||||||
|
|
||||||
|
const summary = await runtime.run(dbPath, new Set(['猫']));
|
||||||
|
|
||||||
|
assert.equal(summary.uniqueWords, 1);
|
||||||
|
assert.equal(summary.knownWordCount, 1);
|
||||||
|
} finally {
|
||||||
|
runtime.destroy();
|
||||||
|
try {
|
||||||
|
db.close();
|
||||||
|
} catch {
|
||||||
|
// The worker needs the setup connection closed before it starts.
|
||||||
|
}
|
||||||
|
fs.rmSync(tempDir, { recursive: true, force: true });
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
|
test('vocabulary summary worker module resolves in the current layout', () => {
|
||||||
|
const workerPath = resolveVocabularySummaryWorkerPath();
|
||||||
|
assert.ok(workerPath, 'expected the vocabulary summary worker module to resolve');
|
||||||
|
assert.ok(workerPath.endsWith(__filename.endsWith('.ts') ? '.ts' : '.js'));
|
||||||
|
});
|
||||||
|
|
||||||
|
test('vocabulary summary worker never falls back to the caller thread', async () => {
|
||||||
|
const runtime = new VocabularySummaryWorkerRuntime({
|
||||||
|
resolveWorkerPath: () => null,
|
||||||
|
warn: () => {},
|
||||||
|
});
|
||||||
|
|
||||||
|
try {
|
||||||
|
await assert.rejects(
|
||||||
|
runtime.run('/tmp/subminer-summary-worker-not-used.sqlite', null),
|
||||||
|
/worker unavailable/i,
|
||||||
|
);
|
||||||
|
} finally {
|
||||||
|
runtime.destroy();
|
||||||
|
}
|
||||||
|
});
|
||||||
@@ -0,0 +1,125 @@
|
|||||||
|
import fs from 'node:fs';
|
||||||
|
import path from 'node:path';
|
||||||
|
import { createLogger } from '../../../logger';
|
||||||
|
import type { VocabularyStatsSummary } from './types';
|
||||||
|
|
||||||
|
interface VocabularySummaryWorkerResponse {
|
||||||
|
summary?: VocabularyStatsSummary;
|
||||||
|
error?: unknown;
|
||||||
|
}
|
||||||
|
|
||||||
|
interface VocabularySummaryWorkerHandle {
|
||||||
|
once(event: 'message', listener: (message: VocabularySummaryWorkerResponse) => void): this;
|
||||||
|
once(event: 'error', listener: (error: Error) => void): this;
|
||||||
|
once(event: 'exit', listener: (code: number) => void): this;
|
||||||
|
terminate(): Promise<number>;
|
||||||
|
}
|
||||||
|
|
||||||
|
interface VocabularySummaryWorkerRuntimeOptions {
|
||||||
|
resolveWorkerPath?: () => string | null;
|
||||||
|
createWorker?: (
|
||||||
|
workerPath: string,
|
||||||
|
workerData: { dbPath: string; knownWords: string[] | null },
|
||||||
|
) => Promise<VocabularySummaryWorkerHandle>;
|
||||||
|
warn?: (message: string, ...meta: unknown[]) => void;
|
||||||
|
}
|
||||||
|
|
||||||
|
export type RunVocabularySummaryTask = (
|
||||||
|
dbPath: string,
|
||||||
|
knownWords: ReadonlySet<string> | null,
|
||||||
|
) => Promise<VocabularyStatsSummary>;
|
||||||
|
|
||||||
|
export function resolveVocabularySummaryWorkerPath(): string | null {
|
||||||
|
const fileName = __filename.endsWith('.ts')
|
||||||
|
? 'vocabulary-summary-worker-thread.ts'
|
||||||
|
: 'vocabulary-summary-worker-thread.js';
|
||||||
|
const workerPath = path.join(__dirname, fileName);
|
||||||
|
return fs.existsSync(workerPath) ? workerPath : null;
|
||||||
|
}
|
||||||
|
|
||||||
|
const logger = createLogger('main:immersion-tracker:vocabulary-summary-worker');
|
||||||
|
|
||||||
|
export class VocabularySummaryWorkerRuntime {
|
||||||
|
private readonly activeWorkers = new Set<VocabularySummaryWorkerHandle>();
|
||||||
|
private destroyed = false;
|
||||||
|
|
||||||
|
constructor(private readonly options: VocabularySummaryWorkerRuntimeOptions = {}) {}
|
||||||
|
|
||||||
|
async run(
|
||||||
|
dbPath: string,
|
||||||
|
knownWords: ReadonlySet<string> | null,
|
||||||
|
): Promise<VocabularyStatsSummary> {
|
||||||
|
if (this.destroyed) throw new Error('Vocabulary summary worker is shut down');
|
||||||
|
const workerData = { dbPath, knownWords: knownWords ? [...knownWords] : null };
|
||||||
|
let worker: VocabularySummaryWorkerHandle;
|
||||||
|
try {
|
||||||
|
const workerPath = (this.options.resolveWorkerPath ?? resolveVocabularySummaryWorkerPath)();
|
||||||
|
if (!workerPath) throw new Error('Emitted vocabulary summary worker module was not found');
|
||||||
|
const createWorker =
|
||||||
|
this.options.createWorker ??
|
||||||
|
(async (resolvedPath, data) => {
|
||||||
|
const { Worker } = await import('node:worker_threads');
|
||||||
|
return new Worker(resolvedPath, { workerData: data });
|
||||||
|
});
|
||||||
|
worker = await createWorker(workerPath, workerData);
|
||||||
|
} catch (error) {
|
||||||
|
if (this.destroyed) throw new Error('Vocabulary summary worker is shut down');
|
||||||
|
(this.options.warn ?? logger.warn)(
|
||||||
|
'Vocabulary summary worker unavailable; refusing to scan vocabulary on the current thread',
|
||||||
|
error,
|
||||||
|
);
|
||||||
|
throw new Error('Vocabulary summary worker unavailable');
|
||||||
|
}
|
||||||
|
|
||||||
|
if (this.destroyed) {
|
||||||
|
await worker.terminate().catch(() => undefined);
|
||||||
|
throw new Error('Vocabulary summary worker is shut down');
|
||||||
|
}
|
||||||
|
|
||||||
|
return new Promise<VocabularyStatsSummary>((resolve, reject) => {
|
||||||
|
let settled = false;
|
||||||
|
this.activeWorkers.add(worker);
|
||||||
|
const settle = (result: VocabularyStatsSummary | Error) => {
|
||||||
|
if (settled) return;
|
||||||
|
settled = true;
|
||||||
|
this.activeWorkers.delete(worker);
|
||||||
|
void worker.terminate().catch(() => undefined);
|
||||||
|
if (result instanceof Error) reject(result);
|
||||||
|
else resolve(result);
|
||||||
|
};
|
||||||
|
|
||||||
|
worker.once('message', (message) => {
|
||||||
|
if (message.summary) {
|
||||||
|
settle(message.summary);
|
||||||
|
return;
|
||||||
|
}
|
||||||
|
settle(
|
||||||
|
new Error(
|
||||||
|
`Vocabulary summary failed: ${String(message.error ?? 'unknown worker error')}`,
|
||||||
|
),
|
||||||
|
);
|
||||||
|
});
|
||||||
|
worker.once('error', (error) => settle(error));
|
||||||
|
worker.once('exit', (code) => {
|
||||||
|
if (!settled) {
|
||||||
|
settle(
|
||||||
|
new Error(
|
||||||
|
code === 0
|
||||||
|
? 'Vocabulary summary worker exited without a response'
|
||||||
|
: `Vocabulary summary worker exited with code ${code}`,
|
||||||
|
),
|
||||||
|
);
|
||||||
|
}
|
||||||
|
});
|
||||||
|
});
|
||||||
|
}
|
||||||
|
|
||||||
|
destroy(): void {
|
||||||
|
if (this.destroyed) return;
|
||||||
|
this.destroyed = true;
|
||||||
|
for (const worker of this.activeWorkers) {
|
||||||
|
void worker.terminate().catch(() => undefined);
|
||||||
|
}
|
||||||
|
this.activeWorkers.clear();
|
||||||
|
}
|
||||||
|
}
|
||||||
@@ -0,0 +1,19 @@
|
|||||||
|
import { parentPort, workerData } from 'node:worker_threads';
|
||||||
|
import { executeVocabularySummaryTask } from './vocabulary-summary-worker';
|
||||||
|
|
||||||
|
interface VocabularySummaryWorkerData {
|
||||||
|
dbPath: string;
|
||||||
|
knownWords: string[] | null;
|
||||||
|
}
|
||||||
|
|
||||||
|
if (!parentPort) throw new Error('vocabulary summary worker missing parent port');
|
||||||
|
|
||||||
|
const request = workerData as VocabularySummaryWorkerData;
|
||||||
|
|
||||||
|
try {
|
||||||
|
parentPort.postMessage({
|
||||||
|
summary: executeVocabularySummaryTask(request.dbPath, request.knownWords),
|
||||||
|
});
|
||||||
|
} catch (error) {
|
||||||
|
parentPort.postMessage({ error: error instanceof Error ? error.message : String(error) });
|
||||||
|
}
|
||||||
@@ -0,0 +1,17 @@
|
|||||||
|
import { getVocabularySummary } from './query-lexical';
|
||||||
|
import { Database } from './sqlite';
|
||||||
|
import { applyPragmas } from './storage';
|
||||||
|
import type { VocabularyStatsSummary } from './types';
|
||||||
|
|
||||||
|
export function executeVocabularySummaryTask(
|
||||||
|
dbPath: string,
|
||||||
|
knownWords: string[] | null,
|
||||||
|
): VocabularyStatsSummary {
|
||||||
|
const db = new Database(dbPath);
|
||||||
|
try {
|
||||||
|
applyPragmas(db);
|
||||||
|
return getVocabularySummary(db, knownWords ? new Set(knownWords) : null);
|
||||||
|
} finally {
|
||||||
|
db.close();
|
||||||
|
}
|
||||||
|
}
|
||||||
@@ -10,6 +10,7 @@ import {
|
|||||||
parseExcludedWordsBody,
|
parseExcludedWordsBody,
|
||||||
parseIntQuery,
|
parseIntQuery,
|
||||||
parsePositiveIdList,
|
parsePositiveIdList,
|
||||||
|
loadKnownWordsSet,
|
||||||
} from './route-support.js';
|
} from './route-support.js';
|
||||||
|
|
||||||
export function registerStatsLibraryRoutes(
|
export function registerStatsLibraryRoutes(
|
||||||
@@ -31,6 +32,17 @@ export function registerStatsLibraryRoutes(
|
|||||||
return c.json(statsJson('vocabulary', vocab));
|
return c.json(statsJson('vocabulary', vocab));
|
||||||
});
|
});
|
||||||
|
|
||||||
|
app.get('/api/stats/vocabulary/summary', async (c) => {
|
||||||
|
const summary = await tracker.getVocabularySummary(
|
||||||
|
loadKnownWordsSet(options?.knownWordCachePath),
|
||||||
|
);
|
||||||
|
return c.json(statsJson('vocabularySummary', summary));
|
||||||
|
});
|
||||||
|
|
||||||
|
app.get('/api/stats/vocabulary/charts', async (c) => {
|
||||||
|
return c.json(statsJson('vocabularyCharts', await tracker.getVocabularyChartData()));
|
||||||
|
});
|
||||||
|
|
||||||
app.get('/api/stats/excluded-words', async (c) => {
|
app.get('/api/stats/excluded-words', async (c) => {
|
||||||
return c.json(statsJson('excludedWords', await tracker.getStatsExcludedWords()));
|
return c.json(statsJson('excludedWords', await tracker.getStatsExcludedWords()));
|
||||||
});
|
});
|
||||||
|
|||||||
@@ -50,6 +50,24 @@ export interface StatsKnownWordsSummary {
|
|||||||
knownWordCount: number;
|
knownWordCount: number;
|
||||||
}
|
}
|
||||||
|
|
||||||
|
export interface StatsVocabularySummary {
|
||||||
|
uniqueWords: number;
|
||||||
|
uniqueWordsWithoutNames: number;
|
||||||
|
uniqueKanji: number;
|
||||||
|
newThisWeek: number;
|
||||||
|
newThisWeekWithoutNames: number;
|
||||||
|
knownWordCount: number | null;
|
||||||
|
knownWordCountWithoutNames: number | null;
|
||||||
|
}
|
||||||
|
|
||||||
|
export interface StatsVocabularyCharts {
|
||||||
|
ready: boolean;
|
||||||
|
topWords: Array<{ wordId: number; headword: string; frequency: number }>;
|
||||||
|
topWordsWithoutNames: Array<{ wordId: number; headword: string; frequency: number }>;
|
||||||
|
newWordsTimeline: Array<{ epochDay: number; wordCount: number }>;
|
||||||
|
newWordsTimelineWithoutNames: Array<{ epochDay: number; wordCount: number }>;
|
||||||
|
}
|
||||||
|
|
||||||
export interface StatsAnilistSearchResult {
|
export interface StatsAnilistSearchResult {
|
||||||
id: number;
|
id: number;
|
||||||
episodes: number | null;
|
episodes: number | null;
|
||||||
@@ -164,6 +182,8 @@ export interface StatsJsonResponseMap {
|
|||||||
sessionEvents: SessionEvent[];
|
sessionEvents: SessionEvent[];
|
||||||
sessionKnownWordsTimeline: StatsSessionKnownWordsTimelinePoint[];
|
sessionKnownWordsTimeline: StatsSessionKnownWordsTimelinePoint[];
|
||||||
vocabulary: VocabularyEntry[];
|
vocabulary: VocabularyEntry[];
|
||||||
|
vocabularySummary: StatsVocabularySummary;
|
||||||
|
vocabularyCharts: StatsVocabularyCharts;
|
||||||
excludedWords: StatsExcludedWord[];
|
excludedWords: StatsExcludedWord[];
|
||||||
setExcludedWords: StatsOkResponse;
|
setExcludedWords: StatsOkResponse;
|
||||||
duplicateLineCleanup: StatsDuplicateLineCleanupResult;
|
duplicateLineCleanup: StatsDuplicateLineCleanupResult;
|
||||||
@@ -222,6 +242,8 @@ export interface StatsHttpClient {
|
|||||||
getSessionEvents: (id: number, limit?: number, eventTypes?: number[]) => Promise<SessionEvent[]>;
|
getSessionEvents: (id: number, limit?: number, eventTypes?: number[]) => Promise<SessionEvent[]>;
|
||||||
getSessionKnownWordsTimeline: (id: number) => Promise<StatsSessionKnownWordsTimelinePoint[]>;
|
getSessionKnownWordsTimeline: (id: number) => Promise<StatsSessionKnownWordsTimelinePoint[]>;
|
||||||
getVocabulary: (limit?: number) => Promise<VocabularyEntry[]>;
|
getVocabulary: (limit?: number) => Promise<VocabularyEntry[]>;
|
||||||
|
getVocabularySummary: () => Promise<StatsVocabularySummary>;
|
||||||
|
getVocabularyCharts: () => Promise<StatsVocabularyCharts>;
|
||||||
getExcludedWords: () => Promise<StatsExcludedWord[]>;
|
getExcludedWords: () => Promise<StatsExcludedWord[]>;
|
||||||
setExcludedWords: (words: StatsExcludedWord[]) => Promise<void>;
|
setExcludedWords: (words: StatsExcludedWord[]) => Promise<void>;
|
||||||
cleanupDuplicateLines: (
|
cleanupDuplicateLines: (
|
||||||
|
|||||||
@@ -98,10 +98,10 @@ export function DuplicateLineCleanup({ onClose, onCleaned }: DuplicateLineCleanu
|
|||||||
|
|
||||||
<div className="space-y-4 px-5 py-4">
|
<div className="space-y-4 px-5 py-4">
|
||||||
<p className="text-xs leading-relaxed text-ctp-subtext0">
|
<p className="text-xs leading-relaxed text-ctp-subtext0">
|
||||||
Typeset subtitles — karaoke openings, animated signs — are authored as one event per
|
Karaoke openings and animated signs are typeset as one subtitle event per animation
|
||||||
animation frame, and older versions counted every frame as its own line. This finds
|
frame, and older versions counted every frame as a line. This collapses those runs back
|
||||||
those runs and collapses each one back to a single line, giving back the word and kanji
|
to one line and drops the word and kanji counts they added. Repeated dialogue is left
|
||||||
counts they inflated. Ordinary repeated dialogue is left alone.
|
alone.
|
||||||
</p>
|
</p>
|
||||||
|
|
||||||
<div>
|
<div>
|
||||||
@@ -192,8 +192,8 @@ export function DuplicateLineCleanup({ onClose, onCleaned }: DuplicateLineCleanu
|
|||||||
</button>
|
</button>
|
||||||
</div>
|
</div>
|
||||||
<p className="text-[11px] text-ctp-overlay1">
|
<p className="text-[11px] text-ctp-overlay1">
|
||||||
Scan first: cleanup removes rows and cannot be undone. Session watch time and lines-seen
|
Scan first: cleanup deletes rows and can't be undone. Watch time and lines-seen totals
|
||||||
totals are left untouched.
|
stay as they are.
|
||||||
</p>
|
</p>
|
||||||
</div>
|
</div>
|
||||||
</div>
|
</div>
|
||||||
|
|||||||
@@ -6,11 +6,10 @@ import { KanjiBreakdown } from './KanjiBreakdown';
|
|||||||
import { KanjiDetailPanel } from './KanjiDetailPanel';
|
import { KanjiDetailPanel } from './KanjiDetailPanel';
|
||||||
import { ExclusionManager } from './ExclusionManager';
|
import { ExclusionManager } from './ExclusionManager';
|
||||||
import { DuplicateLineCleanup } from './DuplicateLineCleanup';
|
import { DuplicateLineCleanup } from './DuplicateLineCleanup';
|
||||||
import { formatNumber } from '../../lib/formatters';
|
import { epochDayToDate, formatNumber } from '../../lib/formatters';
|
||||||
import { TrendChart } from '../trends/TrendChart';
|
import { TrendChart } from '../trends/TrendChart';
|
||||||
import { FrequencyRankTable } from './FrequencyRankTable';
|
import { FrequencyRankTable } from './FrequencyRankTable';
|
||||||
import { CrossAnimeWordsTable } from './CrossAnimeWordsTable';
|
import { CrossAnimeWordsTable } from './CrossAnimeWordsTable';
|
||||||
import { buildVocabularySummary } from '../../lib/dashboard-data';
|
|
||||||
import type { ExcludedWord } from '../../hooks/useExcludedWords';
|
import type { ExcludedWord } from '../../hooks/useExcludedWords';
|
||||||
import type { KanjiEntry, VocabularyEntry } from '../../types/stats';
|
import type { KanjiEntry, VocabularyEntry } from '../../types/stats';
|
||||||
|
|
||||||
@@ -35,7 +34,7 @@ export function VocabularyTab({
|
|||||||
onRemoveExclusion,
|
onRemoveExclusion,
|
||||||
onClearExclusions,
|
onClearExclusions,
|
||||||
}: VocabularyTabProps) {
|
}: VocabularyTabProps) {
|
||||||
const { words, kanji, knownWords, loading, error, reload } = useVocabulary();
|
const { words, kanji, knownWords, summary, charts, loading, error, reload } = useVocabulary();
|
||||||
const [selectedKanjiId, setSelectedKanjiId] = useState<number | null>(null);
|
const [selectedKanjiId, setSelectedKanjiId] = useState<number | null>(null);
|
||||||
const [hideNames, setHideNames] = useState(false);
|
const [hideNames, setHideNames] = useState(false);
|
||||||
const [showExclusionManager, setShowExclusionManager] = useState(false);
|
const [showExclusionManager, setShowExclusionManager] = useState(false);
|
||||||
@@ -48,19 +47,26 @@ export function VocabularyTab({
|
|||||||
if (excluded.length > 0) result = result.filter((w) => !isExcluded(w));
|
if (excluded.length > 0) result = result.filter((w) => !isExcluded(w));
|
||||||
return result;
|
return result;
|
||||||
}, [words, hideNames, excluded, isExcluded]);
|
}, [words, hideNames, excluded, isExcluded]);
|
||||||
const summary = useMemo(
|
const chartData = useMemo(
|
||||||
() => buildVocabularySummary(filteredWords, kanji),
|
() => ({
|
||||||
[filteredWords, kanji],
|
topWords:
|
||||||
|
((hideNames ? charts?.topWordsWithoutNames : charts?.topWords) ?? []).map((word) => ({
|
||||||
|
label: word.headword,
|
||||||
|
value: word.frequency,
|
||||||
|
})) ?? [],
|
||||||
|
newWordsTimeline:
|
||||||
|
((hideNames ? charts?.newWordsTimelineWithoutNames : charts?.newWordsTimeline) ?? []).map(
|
||||||
|
(point) => ({
|
||||||
|
label: epochDayToDate(point.epochDay).toLocaleDateString(undefined, {
|
||||||
|
month: 'short',
|
||||||
|
day: 'numeric',
|
||||||
|
}),
|
||||||
|
value: point.wordCount,
|
||||||
|
}),
|
||||||
|
) ?? [],
|
||||||
|
}),
|
||||||
|
[charts, hideNames],
|
||||||
);
|
);
|
||||||
const knownWordCount = useMemo(() => {
|
|
||||||
if (knownWords.size === 0) return 0;
|
|
||||||
|
|
||||||
let count = 0;
|
|
||||||
for (const w of filteredWords) {
|
|
||||||
if (knownWords.has(w.headword)) count += 1;
|
|
||||||
}
|
|
||||||
return count;
|
|
||||||
}, [filteredWords, knownWords]);
|
|
||||||
|
|
||||||
if (loading) {
|
if (loading) {
|
||||||
return (
|
return (
|
||||||
@@ -82,7 +88,9 @@ export function VocabularyTab({
|
|||||||
};
|
};
|
||||||
|
|
||||||
const handleBarClick = (headword: string): void => {
|
const handleBarClick = (headword: string): void => {
|
||||||
const match = filteredWords.find((w) => w.headword === headword);
|
const match = (hideNames ? charts?.topWordsWithoutNames : charts?.topWords)?.find(
|
||||||
|
(word) => word.headword === headword,
|
||||||
|
);
|
||||||
if (match) onOpenWordDetail?.(match.wordId);
|
if (match) onOpenWordDetail?.(match.wordId);
|
||||||
};
|
};
|
||||||
|
|
||||||
@@ -90,29 +98,43 @@ export function VocabularyTab({
|
|||||||
setSelectedKanjiId(entry.kanjiId);
|
setSelectedKanjiId(entry.kanjiId);
|
||||||
};
|
};
|
||||||
|
|
||||||
|
const displayedSummary = hideNames
|
||||||
|
? {
|
||||||
|
uniqueWords: summary?.uniqueWordsWithoutNames ?? 0,
|
||||||
|
newThisWeek: summary?.newThisWeekWithoutNames ?? 0,
|
||||||
|
knownWordCount: summary?.knownWordCountWithoutNames ?? null,
|
||||||
|
}
|
||||||
|
: {
|
||||||
|
uniqueWords: summary?.uniqueWords ?? 0,
|
||||||
|
newThisWeek: summary?.newThisWeek ?? 0,
|
||||||
|
knownWordCount: summary?.knownWordCount ?? null,
|
||||||
|
};
|
||||||
|
|
||||||
return (
|
return (
|
||||||
<div className="space-y-4">
|
<div className="space-y-4">
|
||||||
<div className="grid grid-cols-2 xl:grid-cols-4 gap-3">
|
<div className="grid grid-cols-2 xl:grid-cols-4 gap-3">
|
||||||
<StatCard
|
<StatCard
|
||||||
label="Unique Words"
|
label="Unique Words"
|
||||||
value={formatNumber(summary.uniqueWords)}
|
value={summary ? formatNumber(displayedSummary.uniqueWords) : '…'}
|
||||||
color="text-ctp-blue"
|
color="text-ctp-blue"
|
||||||
/>
|
/>
|
||||||
{knownWords.size > 0 && (
|
{displayedSummary.knownWordCount !== null ? (
|
||||||
<StatCard
|
<StatCard
|
||||||
label="Known Words"
|
label="Known Words"
|
||||||
value={`${formatNumber(knownWordCount)} (${summary.uniqueWords > 0 ? Math.round((knownWordCount / summary.uniqueWords) * 100) : 0}%)`}
|
value={`${formatNumber(displayedSummary.knownWordCount)} (${displayedSummary.uniqueWords > 0 ? Math.round((displayedSummary.knownWordCount / displayedSummary.uniqueWords) * 100) : 0}%)`}
|
||||||
color="text-ctp-green"
|
color="text-ctp-green"
|
||||||
/>
|
/>
|
||||||
)}
|
) : knownWords.size > 0 ? (
|
||||||
|
<StatCard label="Known Words" value="…" color="text-ctp-green" />
|
||||||
|
) : null}
|
||||||
<StatCard
|
<StatCard
|
||||||
label="Unique Kanji"
|
label="Unique Kanji"
|
||||||
value={formatNumber(summary.uniqueKanji)}
|
value={summary ? formatNumber(summary.uniqueKanji) : '…'}
|
||||||
color="text-ctp-teal"
|
color="text-ctp-teal"
|
||||||
/>
|
/>
|
||||||
<StatCard
|
<StatCard
|
||||||
label="New This Week"
|
label="New This Week"
|
||||||
value={`+${formatNumber(summary.newThisWeek)}`}
|
value={summary ? `+${formatNumber(displayedSummary.newThisWeek)}` : '…'}
|
||||||
color="text-ctp-mauve"
|
color="text-ctp-mauve"
|
||||||
/>
|
/>
|
||||||
</div>
|
</div>
|
||||||
@@ -154,19 +176,25 @@ export function VocabularyTab({
|
|||||||
<div className="grid grid-cols-1 xl:grid-cols-2 gap-4">
|
<div className="grid grid-cols-1 xl:grid-cols-2 gap-4">
|
||||||
<TrendChart
|
<TrendChart
|
||||||
title="Top Repeated Words"
|
title="Top Repeated Words"
|
||||||
data={summary.topWords}
|
data={chartData.topWords}
|
||||||
color="#8aadf4"
|
color="#8aadf4"
|
||||||
type="bar"
|
type="bar"
|
||||||
onBarClick={handleBarClick}
|
onBarClick={handleBarClick}
|
||||||
/>
|
/>
|
||||||
<TrendChart
|
<TrendChart
|
||||||
title="New Words by Day"
|
title="New Words by Day"
|
||||||
data={summary.newWordsTimeline}
|
data={chartData.newWordsTimeline}
|
||||||
color="#c6a0f6"
|
color="#c6a0f6"
|
||||||
type="line"
|
type="line"
|
||||||
/>
|
/>
|
||||||
</div>
|
</div>
|
||||||
|
|
||||||
|
{charts && !charts.ready && (
|
||||||
|
<p className="text-xs text-ctp-overlay1" role="status">
|
||||||
|
Building vocabulary history in the background…
|
||||||
|
</p>
|
||||||
|
)}
|
||||||
|
|
||||||
<FrequencyRankTable
|
<FrequencyRankTable
|
||||||
words={filteredWords}
|
words={filteredWords}
|
||||||
knownWords={knownWords}
|
knownWords={knownWords}
|
||||||
|
|||||||
@@ -1,11 +1,18 @@
|
|||||||
import { useState, useEffect, useCallback } from 'react';
|
import { useState, useEffect, useCallback } from 'react';
|
||||||
import { getStatsClient } from './useStatsApi';
|
import { getStatsClient } from './useStatsApi';
|
||||||
import type { VocabularyEntry, KanjiEntry } from '../types/stats';
|
import type {
|
||||||
|
VocabularyEntry,
|
||||||
|
KanjiEntry,
|
||||||
|
StatsVocabularyCharts,
|
||||||
|
StatsVocabularySummary,
|
||||||
|
} from '../types/stats';
|
||||||
|
|
||||||
export function useVocabulary() {
|
export function useVocabulary() {
|
||||||
const [words, setWords] = useState<VocabularyEntry[]>([]);
|
const [words, setWords] = useState<VocabularyEntry[]>([]);
|
||||||
const [kanji, setKanji] = useState<KanjiEntry[]>([]);
|
const [kanji, setKanji] = useState<KanjiEntry[]>([]);
|
||||||
const [knownWords, setKnownWords] = useState<Set<string>>(new Set());
|
const [knownWords, setKnownWords] = useState<Set<string>>(new Set());
|
||||||
|
const [summary, setSummary] = useState<StatsVocabularySummary | null>(null);
|
||||||
|
const [charts, setCharts] = useState<StatsVocabularyCharts | null>(null);
|
||||||
const [loading, setLoading] = useState(true);
|
const [loading, setLoading] = useState(true);
|
||||||
const [error, setError] = useState<string | null>(null);
|
const [error, setError] = useState<string | null>(null);
|
||||||
// Bumped by `reload` after maintenance rewrites the vocabulary tables.
|
// Bumped by `reload` after maintenance rewrites the vocabulary tables.
|
||||||
@@ -16,6 +23,8 @@ export function useVocabulary() {
|
|||||||
let cancelled = false;
|
let cancelled = false;
|
||||||
setLoading(true);
|
setLoading(true);
|
||||||
setError(null);
|
setError(null);
|
||||||
|
setSummary(null);
|
||||||
|
setCharts(null);
|
||||||
const client = getStatsClient();
|
const client = getStatsClient();
|
||||||
Promise.allSettled([client.getVocabulary(500), client.getKanji(200), client.getKnownWords()])
|
Promise.allSettled([client.getVocabulary(500), client.getKanji(200), client.getKnownWords()])
|
||||||
.then(([wordsResult, kanjiResult, knownResult]) => {
|
.then(([wordsResult, kanjiResult, knownResult]) => {
|
||||||
@@ -46,10 +55,36 @@ export function useVocabulary() {
|
|||||||
if (cancelled) return;
|
if (cancelled) return;
|
||||||
setLoading(false);
|
setLoading(false);
|
||||||
});
|
});
|
||||||
|
void client
|
||||||
|
.getVocabularySummary()
|
||||||
|
.then((nextSummary) => {
|
||||||
|
if (!cancelled) setSummary(nextSummary);
|
||||||
|
})
|
||||||
|
.catch((summaryError: unknown) => {
|
||||||
|
console.error('Failed to load vocabulary summary', summaryError);
|
||||||
|
});
|
||||||
|
let chartRetryTimer: ReturnType<typeof setTimeout> | null = null;
|
||||||
|
const loadCharts = (): void => {
|
||||||
|
void client
|
||||||
|
.getVocabularyCharts()
|
||||||
|
.then((nextCharts) => {
|
||||||
|
if (cancelled) return;
|
||||||
|
setCharts(nextCharts);
|
||||||
|
if (!nextCharts.ready) {
|
||||||
|
chartRetryTimer = setTimeout(loadCharts, 1_000);
|
||||||
|
}
|
||||||
|
})
|
||||||
|
.catch((chartError: unknown) => {
|
||||||
|
console.error('Failed to load vocabulary charts', chartError);
|
||||||
|
if (!cancelled) chartRetryTimer = setTimeout(loadCharts, 1_000);
|
||||||
|
});
|
||||||
|
};
|
||||||
|
loadCharts();
|
||||||
return () => {
|
return () => {
|
||||||
cancelled = true;
|
cancelled = true;
|
||||||
|
if (chartRetryTimer) clearTimeout(chartRetryTimer);
|
||||||
};
|
};
|
||||||
}, [reloadToken]);
|
}, [reloadToken]);
|
||||||
|
|
||||||
return { words, kanji, knownWords, loading, error, reload };
|
return { words, kanji, knownWords, summary, charts, loading, error, reload };
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -100,6 +100,8 @@ export const apiClient = {
|
|||||||
getSessionKnownWordsTimeline: (id: number) =>
|
getSessionKnownWordsTimeline: (id: number) =>
|
||||||
fetchJson('sessionKnownWordsTimeline', `/api/stats/sessions/${id}/known-words-timeline`),
|
fetchJson('sessionKnownWordsTimeline', `/api/stats/sessions/${id}/known-words-timeline`),
|
||||||
getVocabulary: (limit = 100) => fetchJson('vocabulary', `/api/stats/vocabulary?limit=${limit}`),
|
getVocabulary: (limit = 100) => fetchJson('vocabulary', `/api/stats/vocabulary?limit=${limit}`),
|
||||||
|
getVocabularySummary: () => fetchJson('vocabularySummary', '/api/stats/vocabulary/summary'),
|
||||||
|
getVocabularyCharts: () => fetchJson('vocabularyCharts', '/api/stats/vocabulary/charts'),
|
||||||
getExcludedWords: () => fetchJson('excludedWords', '/api/stats/excluded-words'),
|
getExcludedWords: () => fetchJson('excludedWords', '/api/stats/excluded-words'),
|
||||||
setExcludedWords: async (words: StatsExcludedWord[]): Promise<void> => {
|
setExcludedWords: async (words: StatsExcludedWord[]): Promise<void> => {
|
||||||
await fetchResponse('/api/stats/excluded-words', {
|
await fetchResponse('/api/stats/excluded-words', {
|
||||||
|
|||||||
@@ -1,7 +1,12 @@
|
|||||||
import assert from 'node:assert/strict';
|
import assert from 'node:assert/strict';
|
||||||
import test from 'node:test';
|
import test from 'node:test';
|
||||||
|
|
||||||
import { epochMsFromDbTimestamp, formatRelativeDate, formatSessionDayLabel } from './formatters';
|
import {
|
||||||
|
epochDayToDate,
|
||||||
|
epochMsFromDbTimestamp,
|
||||||
|
formatRelativeDate,
|
||||||
|
formatSessionDayLabel,
|
||||||
|
} from './formatters';
|
||||||
|
|
||||||
const FIXED_NOW = new Date(2026, 2, 16, 12, 0, 0).getTime();
|
const FIXED_NOW = new Date(2026, 2, 16, 12, 0, 0).getTime();
|
||||||
|
|
||||||
@@ -108,6 +113,19 @@ test('epochMsFromDbTimestamp keeps ms timestamps as-is', () => {
|
|||||||
assert.equal(epochMsFromDbTimestamp(1_700_000_000_000), 1_700_000_000_000);
|
assert.equal(epochMsFromDbTimestamp(1_700_000_000_000), 1_700_000_000_000);
|
||||||
});
|
});
|
||||||
|
|
||||||
|
test('epochDayToDate preserves the calendar day west of UTC', () => {
|
||||||
|
const previousTimezone = process.env.TZ;
|
||||||
|
process.env.TZ = 'America/Los_Angeles';
|
||||||
|
try {
|
||||||
|
const epochDay = Math.floor(Date.UTC(2026, 2, 16) / 86_400_000);
|
||||||
|
const date = epochDayToDate(epochDay);
|
||||||
|
assert.deepEqual([date.getFullYear(), date.getMonth(), date.getDate()], [2026, 2, 16]);
|
||||||
|
} finally {
|
||||||
|
if (previousTimezone === undefined) delete process.env.TZ;
|
||||||
|
else process.env.TZ = previousTimezone;
|
||||||
|
}
|
||||||
|
});
|
||||||
|
|
||||||
test('formatSessionDayLabel formats today and yesterday', () => {
|
test('formatSessionDayLabel formats today and yesterday', () => {
|
||||||
withFixedNow((now) => {
|
withFixedNow((now) => {
|
||||||
const oneDayMs = 24 * 60 * 60_000;
|
const oneDayMs = 24 * 60 * 60_000;
|
||||||
|
|||||||
@@ -38,7 +38,8 @@ export function formatRelativeDate(ms: number): string {
|
|||||||
}
|
}
|
||||||
|
|
||||||
export function epochDayToDate(epochDay: number): Date {
|
export function epochDayToDate(epochDay: number): Date {
|
||||||
return new Date(epochDay * 86_400_000);
|
const utcDate = new Date(epochDay * 86_400_000);
|
||||||
|
return new Date(utcDate.getUTCFullYear(), utcDate.getUTCMonth(), utcDate.getUTCDate());
|
||||||
}
|
}
|
||||||
|
|
||||||
export function localDayFromMs(ms: number): number {
|
export function localDayFromMs(ms: number): number {
|
||||||
|
|||||||
@@ -6,6 +6,7 @@ import { fileURLToPath } from 'node:url';
|
|||||||
const VOCABULARY_TAB_PATH = fileURLToPath(
|
const VOCABULARY_TAB_PATH = fileURLToPath(
|
||||||
new URL('../components/vocabulary/VocabularyTab.tsx', import.meta.url),
|
new URL('../components/vocabulary/VocabularyTab.tsx', import.meta.url),
|
||||||
);
|
);
|
||||||
|
const VOCABULARY_HOOK_PATH = fileURLToPath(new URL('../hooks/useVocabulary.ts', import.meta.url));
|
||||||
|
|
||||||
test('VocabularyTab declares all hooks before loading and error early returns', () => {
|
test('VocabularyTab declares all hooks before loading and error early returns', () => {
|
||||||
const source = fs.readFileSync(VOCABULARY_TAB_PATH, 'utf8');
|
const source = fs.readFileSync(VOCABULARY_TAB_PATH, 'utf8');
|
||||||
@@ -20,15 +21,28 @@ test('VocabularyTab declares all hooks before loading and error early returns',
|
|||||||
assert.deepEqual(hooksAfterLoadingGuard ?? [], []);
|
assert.deepEqual(hooksAfterLoadingGuard ?? [], []);
|
||||||
});
|
});
|
||||||
|
|
||||||
test('VocabularyTab memoizes summary and known-word aggregate calculations', () => {
|
test('VocabularyTab uses uncapped server-side data for its charts and card totals', () => {
|
||||||
const source = fs.readFileSync(VOCABULARY_TAB_PATH, 'utf8');
|
const source = fs.readFileSync(VOCABULARY_TAB_PATH, 'utf8');
|
||||||
|
|
||||||
assert.match(
|
assert.match(
|
||||||
source,
|
source,
|
||||||
/const summary = useMemo\([\s\S]*buildVocabularySummary\(filteredWords, kanji\)[\s\S]*\[filteredWords, kanji\][\s\S]*\);/,
|
/const \{ words, kanji, knownWords, summary, charts, loading, error, reload \} = useVocabulary\(\);/,
|
||||||
);
|
);
|
||||||
|
assert.match(source, /charts\?\.topWordsWithoutNames/);
|
||||||
|
assert.match(source, /charts\?\.newWordsTimelineWithoutNames/);
|
||||||
|
assert.doesNotMatch(source, /buildVocabularySummary\(/);
|
||||||
|
assert.match(source, /uniqueWords: summary\?\.uniqueWordsWithoutNames \?\? 0/);
|
||||||
|
assert.match(source, /uniqueWords: summary\?\.uniqueWords \?\? 0/);
|
||||||
|
assert.match(source, /value=\{summary \? formatNumber\(summary\.uniqueKanji\) : '…'\}/);
|
||||||
|
});
|
||||||
|
|
||||||
|
test('useVocabulary loads exact card totals without holding up the vocabulary tables', () => {
|
||||||
|
const source = fs.readFileSync(VOCABULARY_HOOK_PATH, 'utf8');
|
||||||
|
|
||||||
assert.match(
|
assert.match(
|
||||||
source,
|
source,
|
||||||
/const knownWordCount = useMemo\(\(\) => \{[\s\S]*for \(const w of filteredWords\) \{[\s\S]*knownWords\.has\(w\.headword\)[\s\S]*\}\s*return count;\s*\}, \[filteredWords, knownWords\]\);/,
|
/Promise\.allSettled\(\[\s*client\.getVocabulary\(500\),\s*client\.getKanji\(200\),\s*client\.getKnownWords\(\),?\s*\]\)/,
|
||||||
);
|
);
|
||||||
|
assert.match(source, /void client\s*\.getVocabularySummary\(\)\s*\.then\(/);
|
||||||
|
assert.match(source, /client\s*\.getVocabularyCharts\(\)/);
|
||||||
});
|
});
|
||||||
|
|||||||
Reference in New Issue
Block a user