Threshold Majority Queries
Time O(nlogn + qlogq + (n + q) * sqrt(n) * logn) · Space O(q + n) · Official statement on LeetCode
Solutions
// Time: O(nlogn + qlogq + (n + q) * sqrt(n) + q * n)
// Space: O(n + q)
// sort, coordinate compression, mo's algorithm
class Solution {
public:
vector<int> subarrayMajority(vector<int>& nums, vector<vector<int>>& queries) {
vector<int> sorted_nums(nums);
sort(begin(sorted_nums), end(sorted_nums));
sorted_nums.erase(unique(begin(sorted_nums), end(sorted_nums)), end(sorted_nums));
unordered_map<int, int> num_to_idx;
for (int i = 0; i < size(sorted_nums); ++i) {
num_to_idx[sorted_nums[i]] = i;
}
// reference: https://cp-algorithms.com/data_structures/sqrt_decomposition.html
const auto& mo_s_algorithm = [&]() {
vector<int> cnt(size(num_to_idx));
vector<int> cnt2(size(nums) + 1);
int max_freq = 0;
const auto& add = [&](int i) {
const auto& idx = num_to_idx[nums[i]];
if (cnt[idx]) {
--cnt2[cnt[idx]];
}
++cnt[idx];
++cnt2[cnt[idx]];
max_freq = max(max_freq, cnt[idx]);
};
const auto& remove = [&](int i) {
const auto& idx = num_to_idx[nums[i]];
--cnt2[cnt[idx]];
if (!cnt2[max_freq]) {
--max_freq;
}
--cnt[idx];
if (cnt[idx]) {
++cnt2[cnt[idx]];
}
};
const auto& get_ans = [&](int t) {
if (max_freq < t) {
return -1;
}
int i = 0;
for (; i < size(cnt); ++i) {
if (cnt[i] == max_freq) {
break;
}
}
return sorted_nums[i];
};
vector<int> result(size(queries), -1);
const int block_size = sqrt(size(nums)) + 1;
vector<int> idxs(size(queries));
iota(begin(idxs), end(idxs), 0);
sort(begin(idxs), end(idxs), [&](const auto& a, const auto& b) {
const auto& i = queries[a][0] / block_size;
const auto& j = queries[b][0] / block_size;
return i != j ? i < j : (i & 1 ? queries[a][1] < queries[b][1] : queries[a][1] > queries[b][1]);
});
int left = 0, right = -1;
for (const auto& i : idxs) {
const auto& l = queries[i][0];
const auto& r = queries[i][1];
const auto& t = queries[i][2];
while (left > l) {
left -= 1;
add(left);
}
while (right < r) {
++right;
add(right);
}
while (left < l) {
remove(left);
++left;
}
while (right > r) {
remove(right);
--right;
}
result[i] = get_ans(t);
}
return result;
};
return mo_s_algorithm();
}
};
// Time: O(nlogn + qlogq + (n + q) * sqrt(n) * logn)
// Space: O(n + q)
// sort, coordinate compression, mo's algorithm, bst
class Solution2 {
public:
vector<int> subarrayMajority(vector<int>& nums, vector<vector<int>>& queries) {
vector<int> sorted_nums(nums);
sort(begin(sorted_nums), end(sorted_nums));
sorted_nums.erase(unique(begin(sorted_nums), end(sorted_nums)), end(sorted_nums));
unordered_map<int, int> num_to_idx;
for (int i = 0; i < size(sorted_nums); ++i) {
num_to_idx[sorted_nums[i]] = i;
}
// reference: https://cp-algorithms.com/data_structures/sqrt_decomposition.html
const auto& mo_s_algorithm = [&]() {
vector<int> cnt(size(num_to_idx));
vector<multiset<int>> lookup(size(nums) + 1);
int max_freq = 0;
const auto& add = [&](int i) {
const auto& idx = num_to_idx[nums[i]];
if (cnt[idx]) {
auto it = lookup[cnt[idx]].find(nums[i]);
lookup[cnt[idx]].erase(it);
}
++cnt[idx];
lookup[cnt[idx]].emplace(nums[i]);
max_freq = max(max_freq, cnt[idx]);
};
const auto& remove = [&](int i) {
const auto& idx = num_to_idx[nums[i]];
auto it = lookup[cnt[idx]].find(nums[i]);
lookup[cnt[idx]].erase(it);
if (empty(lookup[max_freq])) {
--max_freq;
}
--cnt[idx];
if (cnt[idx]) {
lookup[cnt[idx]].emplace(nums[i]);
}
};
const auto& get_ans = [&](int t) {
return max_freq >= t ? *begin(lookup[max_freq]) : -1;
};
vector<int> result(size(queries), -1);
const int block_size = sqrt(size(nums)) + 1;
vector<int> idxs(size(queries));
iota(begin(idxs), end(idxs), 0);
sort(begin(idxs), end(idxs), [&](const auto& a, const auto& b) {
const auto& i = queries[a][0] / block_size;
const auto& j = queries[b][0] / block_size;
return i != j ? i < j : (i & 1 ? queries[a][1] < queries[b][1] : queries[a][1] > queries[b][1]);
});
int left = 0, right = -1;
for (const auto& i : idxs) {
const auto& l = queries[i][0];
const auto& r = queries[i][1];
const auto& t = queries[i][2];
while (left > l) {
left -= 1;
add(left);
}
while (right < r) {
++right;
add(right);
}
while (left < l) {
remove(left);
++left;
}
while (right > r) {
remove(right);
--right;
}
result[i] = get_ans(t);
}
return result;
};
return mo_s_algorithm();
}
};
Beginner Explanation
What is Threshold Majority Queries?
Threshold Majority Queries (LeetCode #3636) is a Hard problem that primarily trains hash table.
How to think about it
- Restate the goal in your own words before coding.
- Work a tiny example by hand so the invariant becomes obvious.
- Identify the pattern — this problem aligns with sort, mos algorithm, and sorted list.
- Only then translate the idea into code.
Why this problem matters
Hard problems force you to combine patterns and prove complexity carefully — interview gold. Official solution notes mention: Sort, Coordinate Compression, Sqrt Decomposition, Mo's Algorithm.
AlgoForge explanations are original teaching notes. Always open the official problem statement on LeetCode for constraints and examples.
Interview Walkthrough
Interview approach for Threshold Majority Queries
Opening (30–60 seconds)
- Clarify inputs/outputs and edge cases (empty input, single element, duplicates, overflow).
- State a brute force so the interviewer knows you can solve it naively.
- Propose the optimal direction tied to sort, mos algorithm, and sorted list.
Core solution narrative
- Define the state you track (pointers, DP cell, set membership, stack top, etc.).
- Explain the transition when you process the next element.
- Call out time (O(nlogn + qlogq + (n + q) * sqrt(n) * logn)) and space (O(q + n)) before coding.
- Code cleanly; narrate variable names.
What interviewers listen for
- Correctness on edge cases
- Complexity honesty
- Ability to discuss trade-offs (e.g., hash map space vs. sort + two pointers)
Follow-up questions they may ask
- Can you solve it with less memory?
- What if the input stream is infinite / doesn't fit in RAM?
- How would tests look for adversarial inputs?
Optimized Approach
Optimized solution notes
The reference solutions on AlgoForge target O(nlogn + qlogq + (n + q) * sqrt(n) * logn) time and O(q + n) space.
Pattern focus: sort, mos algorithm, and sorted list
Use the pattern as a checklist:
- sort — confirm the invariant holds after each step
- mos algorithm — confirm the invariant holds after each step
- sorted list — confirm the invariant holds after each step
Multiple methods appear in the source solutions — compare them and explain when each is preferable.
Implementation tips
- Prefer readable names over micro-optimizations in interviews.
- Extract helpers only when they clarify (e.g., expand-around-center, DFS visit).
- After AC-level logic, re-scan for off-by-one and null checks.
Complexity Analysis
Complexity
| Measure | Bound |
|---|---|
| Time | O(nlogn + qlogq + (n + q) * sqrt(n) * logn) |
| Space | O(q + n) |
How to justify this in an interview
- Time: count loops, map/set operations, and recursive branching; state average vs worst case if relevant.
- Space: include hash maps, recursion stack, and output allocation when the problem asks for it.
If your implementation differs from the reference, re-derive big-O from your code — never memorize a complexity you cannot defend.
Common Mistakes
Common mistakes on Threshold Majority Queries
- Skipping edge cases — empty collections, single-element inputs, max constraints.
- Wrong invariant for sort, mos algorithm, and sorted list — updating state too early or too late.
- Mutating input unexpectedly when the problem forbids it.
- Off-by-one in windows, ranges, or binary search bounds.
- Ignoring overflow / precision for integer arithmetic problems.
- Overengineering — jumping to an advanced structure when a simpler approach works.
Alternative Approaches
Alternatives
The source file includes more than one method. Compare:
- Primary optimized path — best complexity for typical interviews.
- Secondary approach — often brute force, sorting-based, or space-optimized variant.
Practice articulating when you would pick each (constraints, readability, follow-ups).
Edge Cases
Edge cases checklist
- Minimum input size
- Maximum input size / time limits
- Duplicates and already-sorted input
- Negative numbers / zeros (if applicable)
- Disconnected structures (graphs/trees)
- Single path vs branching recursion depth
Pattern Recognition
Spotting this pattern
Signal phrases that point to sort, mos algorithm, and sorted list:
- Sorted input or ability to sort without changing the answer class
- Need for contiguous subarray / substring → consider sliding window
- Need for O(1) membership → hash set/map
- Optimal substructure + overlapping subproblems → DP
- Connectivity / components → graph DFS/BFS or Union-Find
Primary topics: hash table.
Follow-up Interview Questions
Follow-ups
- How does the solution change if the input is a stream?
- Can you solve it in-place?
- What if duplicates must be handled differently?
- How would you parallelize the approach?
- Design tests that would break a buggy implementation.
Practice Recommendations
What to practice next
- Re-solve Threshold Majority Queries in a second language (cpp, python).
- Drill 3–5 more problems tagged hash table.
- Teach the solution out loud in under 5 minutes.
- Add this problem to your revision calendar in 3 days and 14 days.
Visualization
Study checklist
- Read the official problem statement on LeetCode
- Solve on paper / whiteboard first
- Implement the sort, mos algorithm, and sorted list approach
- Verify edge cases from the checklist
- State time and space complexity aloud
- Compare with the AlgoForge reference solution
- Schedule a revision session
Revision notes
Threshold Majority Queries (#3636) — Hard. Pattern: sort, mos algorithm, and sorted list. Complexity: O(nlogn + qlogq + (n + q) * sqrt(n) * logn) time / O(q + n) space. Re-derive the invariant before coding.
FAQs
What is the time complexity of Threshold Majority Queries?+
The reference solutions aim for O(nlogn + qlogq + (n + q) * sqrt(n) * logn) time and O(q + n) space. Always re-derive complexity from the code you write in the interview.
What pattern does Threshold Majority Queries use?+
It primarily maps to sort, mos algorithm, and sorted list, within the broader topic of hash table.
Is Threshold Majority Queries good for interviews?+
Yes — as a Hard problem it is a solid practice target. Pair it with related problems in the same pattern family for spaced repetition.
Where can I read the official statement?+
Open the official LeetCode page for constraints and examples: https://leetcode.com/problems/threshold-majority-queries/