\documentclass[12pt,oneside,openany]{book}
\usepackage[T1]{fontenc}
\usepackage[utf8]{inputenc}
\usepackage{amsmath,amssymb,amsthm,mathtools,mathrsfs}
\usepackage{booktabs,array}
\usepackage[numbers,sort&compress]{natbib}
\usepackage{imakeidx}% required by \indexsetup in shared/trilogy-style.tex
\usepackage{environ}
\PassOptionsToPackage{hypertexnames=false}{hyperref}
\input{shared/trilogy-style}% shared design
\usepackage[nameinlink,capitalize,noabbrev]{cleveref}
\usepackage{aliascnt}
\hypersetup{pdftitle={Nearly Minimax Rates for Functional Estimation Under Rough Random Design},pdfauthor={P. M. Aronow, Nathan Kallus, and Patrick Lopatto}}
% The paper is a single part. Its top-level units are sections: they are set with the
% design's chapter headings, labeled "Section", and listed as such in the contents; their
% subsections are the design's numbered sections. The appendices are the chapters after
% \appendix: they are headed and referred to as "Appendix A", and their sections as
% "Appendix A.1".
\renewcommand{\chaptername}{Section}
\numberwithin{equation}{chapter}
\theoremstyle{plain}
\newtheorem{theorem}{Theorem}[chapter]
\newaliascnt{proposition}{theorem}
\newtheorem{proposition}[proposition]{Proposition}
\aliascntresetthe{proposition}
\newaliascnt{lemma}{theorem}
\newtheorem{lemma}[lemma]{Lemma}
\aliascntresetthe{lemma}
\newaliascnt{corollary}{theorem}
\newtheorem{corollary}[corollary]{Corollary}
\aliascntresetthe{corollary}
\theoremstyle{definition}
\newaliascnt{definition}{theorem}
\newtheorem{definition}[definition]{Definition}
\aliascntresetthe{definition}
\newaliascnt{example}{theorem}
\newtheorem{example}[example]{Example}
\aliascntresetthe{example}
\theoremstyle{remark}
\newaliascnt{remark}{theorem}
\newtheorem{remark}[remark]{Remark}
\aliascntresetthe{remark}
\crefname{theorem}{Theorem}{Theorems}
\crefname{proposition}{Proposition}{Propositions}
\crefname{lemma}{Lemma}{Lemmas}
\crefname{corollary}{Corollary}{Corollaries}
\crefname{definition}{Definition}{Definitions}
\crefname{example}{Example}{Examples}
\crefname{remark}{Remark}{Remarks}
\crefname{chapter}{Section}{Sections}
\crefname{section}{Section}{Sections}
\crefname{subsection}{Section}{Sections}
% After \appendix, cleveref files chapters as "appendix" and their sections as "subappendix".
\crefname{appendix}{Appendix}{Appendices}
\crefname{subappendix}{Appendix}{Appendices}
\crefname{subsubappendix}{Appendix}{Appendices}
\crefname{table}{Table}{Tables}
% Equations are referred to by their number in parentheses, as with \eqref.
\crefname{equation}{}{}
\crefformat{equation}{(#2#1#3)}
\Crefformat{equation}{(#2#1#3)}
\crefrangeformat{equation}{(#3#1#4)--(#5#2#6)}
\Crefrangeformat{equation}{(#3#1#4)--(#5#2#6)}
\crefmultiformat{equation}{(#2#1#3)}{ and~(#2#1#3)}{, (#2#1#3)}{ and~(#2#1#3)}
\Crefmultiformat{equation}{(#2#1#3)}{ and~(#2#1#3)}{, (#2#1#3)}{ and~(#2#1#3)}
% A centered block of small type on a narrower measure, with an optional spaced
% small-capitals label: the authors' note and the abstract.
\makeatletter
\newcommand{\npp@noteblock}[2]{%
\ifx\relax#1\relax\else{\centering\trilogy@spacedsc{#1}\par}\vspace{6pt}\fi
{\centering
\begin{minipage}{0.84\textwidth}\small\normalfont\setlength{\parindent}{1.5em}%
\noindent\ignorespaces#2\end{minipage}\par}}
% The authors' note and the abstract each have a page of their own, after the title page,
% set a little above the middle of the page and with no head or foot.
\NewEnviron{authorsnote}{%
\clearpage\thispagestyle{empty}%
\vspace*{\stretch{2}}%
\npp@noteblock{Authors' note}{\BODY}%
\vspace*{\stretch{3}}%
\clearpage}
\NewEnviron{abstractpage}{%
\clearpage\thispagestyle{empty}%
\vspace*{\stretch{2}}%
\npp@noteblock{Abstract}{\BODY}%
\vspace*{\stretch{3}}%
\clearpage}
% Title page: the title at a fixed position, and the
% authors, one to a line, in the lower half of the page.
% \npptitlepage{
}{}
\newcommand{\npptitlepage}[2]{%
\hypersetup{pageanchor=false}%
\begin{titlepage}
\centering\normalfont
\vspace*{1.7in}
{\large\strut\par}
\vspace{1.05in}
{\fontsize{28}{35}\selectfont #1\par}
\vspace{1.1in}
{\Large\def\\{\par\vspace{10pt}}#2\par}
\vfill
\end{titlepage}%
\hypersetup{pageanchor=true}}
\makeatother
\newcommand{\E}{\mathbb E}
\newcommand{\Pp}{\mathbb P}
\newcommand{\Var}{\operatorname{Var}}
\newcommand{\Cov}{\operatorname{Cov}}
\newcommand{\TV}{\operatorname{TV}}
\newcommand{\supp}{\operatorname{supp}}
\newcommand{\cP}{\mathcal P}
\newcommand{\cH}{\mathcal H}
\newcommand{\dd}{\,\mathrm d}
\newcommand{\risk}{r_n}
% Short running heads for section titles that do not fit beside their chapter title.
\shorthead{Embedding the priors in the generic model}{Embedding the priors}
\shorthead{Proof of \texorpdfstring{\cref{cor:applications}}{the corollary} for these targets}{Proof for these targets}
\begin{document}
\frontmatter
\npptitlepage{Nearly Minimax Rates\\ for Functional Estimation\\ Under Rough Random Design}{P.\ M.\ Aronow\\ Nathan Kallus\\ Patrick Lopatto}
\begin{authorsnote}
This paper was written by the authors with the assistance of large language
models, which were used to suggest mathematical arguments, draft and revise
exposition, and write code for computational verification. We have taken care
to supervise this writing to ensure that the proofs meet a minimum standard of
readability and that previous literature is appropriately cited. However, the
resulting document remains lacking in exposition and overall coherence. We are
nonetheless releasing this manuscript on Hexagon to disseminate the result as
quickly as possible, because we believe it may be of interest to the statistics
community. We welcome suggestions concerning mathematical corrections, omitted
citations, and attribution of ideas.
P.L.\ was partially supported by NSF grant DMS-2450004.
\end{authorsnote}
\begin{abstractpage}
We establish nearly minimax bounds for missing-at-random means,
treatment effects, and expected conditional covariances under rough
random design.
For two nuisance functions with average H\"older smoothness
$s$ in dimension $d$, the minimax root-mean-square error is
$n^{-2s/d+o(1)}$ when $s3.0.CO;2-#}.
\bibitem[Robins et~al.(1994)Robins, Rotnitzky, and Zhao]{rrz}
J.~M. Robins, A.~Rotnitzky, and L.~P. Zhao.
\newblock Estimation of regression coefficients when some regressors are not
always observed.
\newblock \emph{Journal of the American Statistical Association}, 89\penalty0
(427):\penalty0 846--866, 1994.
\newblock \doi{10.1080/01621459.1994.10476818}.
\bibitem[Robins et~al.(2015)Robins, Li, Liu, Mukherjee, Tchetgen~Tchetgen, and
van~der Vaart]{robins}
J.~M. Robins, L.~Li, L.~Liu, R.~Mukherjee, E.~Tchetgen~Tchetgen, and A.~van~der
Vaart.
\newblock Minimax estimation of a functional on a structured high-dimensional
model (corrected version), 2015.
\newblock \href{https://arxiv.org/abs/1512.02174v3}{arXiv:1512.02174v3}.
\bibitem[Robins et~al.(2017)Robins, Li, Mukherjee, Tchetgen~Tchetgen, and
van~der Vaart]{robins2017}
J.~M. Robins, L.~Li, R.~Mukherjee, E.~Tchetgen~Tchetgen, and A.~van~der Vaart.
\newblock Minimax estimation of a functional on a structured high-dimensional
model.
\newblock \emph{The Annals of Statistics}, 45\penalty0 (5):\penalty0
1951--1987, 2017.
\newblock \doi{10.1214/16-AOS1515}.
\bibitem[Robinson(1988)]{robinson}
P.~M. Robinson.
\newblock Root-{$N$}-consistent semiparametric regression.
\newblock \emph{Econometrica}, 56\penalty0 (4):\penalty0 931--954, 1988.
\newblock \doi{10.2307/1912705}.
\bibitem[Rotnitzky et~al.(2021)Rotnitzky, Smucler, and Robins]{mixedbias}
A.~Rotnitzky, E.~Smucler, and J.~M. Robins.
\newblock Characterization of parameters with a mixed bias property.
\newblock \emph{Biometrika}, 108\penalty0 (1):\penalty0 231--238, 2021.
\newblock \doi{10.1093/biomet/asaa054}.
\bibitem[Rubin(1976)]{rubin}
D.~B. Rubin.
\newblock Inference and missing data.
\newblock \emph{Biometrika}, 63\penalty0 (3):\penalty0 581--592, 1976.
\newblock \doi{10.1093/biomet/63.3.581}.
\bibitem[S{\'a}nchez-Becerra(2023)]{sanchezbecerra2023}
A.~S{\'a}nchez-Becerra.
\newblock Robust inference for the treatment effect variance in experiments
using machine learning, 2023.
\newblock \href{https://arxiv.org/abs/2306.03363v1}{arXiv:2306.03363v1}.
\bibitem[Serfling(1980)]{serfling}
R.~J. Serfling.
\newblock \emph{Approximation Theorems of Mathematical Statistics}.
\newblock Wiley, New York, 1980.
\newblock \doi{10.1002/9780470316481}.
\bibitem[Shen et~al.(2020)Shen, Gao, Witten, and Han]{shen}
Y.~Shen, C.~Gao, D.~Witten, and F.~Han.
\newblock Optimal estimation of variance in nonparametric regression with
random design.
\newblock \emph{The Annals of Statistics}, 48\penalty0 (6):\penalty0
3589--3618, 2020.
\newblock \doi{10.1214/20-AOS1944}.
\bibitem[Song(2026)]{song2026}
X.~Song.
\newblock Batched and complete {U}-statistics for trace-polynomial estimation
from classical shadows, 2026.
\newblock \href{https://arxiv.org/abs/2608.22962v1}{arXiv:2608.22962v1}.
\bibitem[Stein and Shakarchi(2003)]{steinshakarchi}
E.~M. Stein and R.~Shakarchi.
\newblock \emph{Fourier Analysis: An Introduction}, volume~1 of \emph{Princeton
Lectures in Analysis}.
\newblock Princeton University Press, Princeton, 2003.
\bibitem[Stone(1982)]{stone}
C.~J. Stone.
\newblock Optimal global rates of convergence for nonparametric regression.
\newblock \emph{The Annals of Statistics}, 10\penalty0 (4):\penalty0
1040--1053, 1982.
\newblock \doi{10.1214/aos/1176345969}.
\bibitem[Tian et~al.(2014)Tian, Alizadeh, Gentles, and Tibshirani]{tian2014}
L.~Tian, A.~A. Alizadeh, A.~J. Gentles, and R.~Tibshirani.
\newblock A simple method for estimating interactions between a treatment and a
large number of covariates.
\newblock \emph{Journal of the American Statistical Association}, 109\penalty0
(508):\penalty0 1517--1532, 2014.
\newblock \doi{10.1080/01621459.2014.951443}.
\bibitem[Trefethen(2019)]{trefethen}
L.~N. Trefethen.
\newblock \emph{Approximation Theory and Approximation Practice}.
\newblock SIAM, Philadelphia, extended edition, 2019.
\newblock \doi{10.1137/1.9781611975949}.
\bibitem[Trefethen and Bau(1997)]{trefethenbau}
L.~N. Trefethen and D.~Bau, III.
\newblock \emph{Numerical Linear Algebra}.
\newblock SIAM, Philadelphia, 1997.
\newblock \doi{10.1137/1.9780898719574}.
\bibitem[Tsybakov(2009)]{tsybakov}
A.~B. Tsybakov.
\newblock \emph{Introduction to Nonparametric Estimation}.
\newblock Springer Series in Statistics. Springer, New York, 2009.
\newblock \doi{10.1007/b13794}.
\bibitem[van~der Vaart(1998)]{vandervaart}
A.~W. van~der Vaart.
\newblock \emph{Asymptotic Statistics}.
\newblock Cambridge University Press, Cambridge, 1998.
\newblock \doi{10.1017/CBO9780511802256}.
\bibitem[Vatedka et~al.(2015)Vatedka, Kashyap, and Thangaraj]{vatedka2015}
S.~Vatedka, N.~Kashyap, and A.~Thangaraj.
\newblock Secure compute-and-forward in a bidirectional relay.
\newblock \emph{IEEE Transactions on Information Theory}, 61\penalty0
(5):\penalty0 2531--2556, 2015.
\newblock \doi{10.1109/TIT.2015.2412114}.
\bibitem[Verzelen and Gassiat(2018)]{verzelengassiat2018}
N.~Verzelen and E.~Gassiat.
\newblock Adaptive estimation of high-dimensional signal-to-noise ratios.
\newblock \emph{Bernoulli}, 24\penalty0 (4B):\penalty0 3683--3710, 2018.
\newblock \doi{10.3150/17-BEJ975}.
\bibitem[Wang and Tchetgen~Tchetgen(2018)]{wangtchetgen}
L.~Wang and E.~Tchetgen~Tchetgen.
\newblock Bounded, efficient and multiply robust estimation of average
treatment effects using instrumental variables.
\newblock \emph{Journal of the Royal Statistical Society Series B: Statistical
Methodology}, 80\penalty0 (3):\penalty0 531--550, 2018.
\newblock \doi{10.1111/rssb.12262}.
\bibitem[Wang et~al.(2008)Wang, Brown, Cai, and Levine]{wangvariance}
L.~Wang, L.~D. Brown, T.~T. Cai, and M.~Levine.
\newblock Effect of mean on variance function estimation in nonparametric
regression.
\newblock \emph{The Annals of Statistics}, 36\penalty0 (2):\penalty0 646--664,
2008.
\newblock \doi{10.1214/009053607000000901}.
\bibitem[Williamson et~al.(2023)Williamson, Gilbert, Simon, and
Carone]{williamson}
B.~D. Williamson, P.~B. Gilbert, N.~R. Simon, and M.~Carone.
\newblock A general framework for inference on algorithm-agnostic variable
importance.
\newblock \emph{Journal of the American Statistical Association}, 118\penalty0
(543):\penalty0 1645--1658, 2023.
\newblock \doi{10.1080/01621459.2021.2003200}.
\bibitem[Wu and Yang(2016)]{wuyangentropy}
Y.~Wu and P.~Yang.
\newblock Minimax rates of entropy estimation on large alphabets via best
polynomial approximation.
\newblock \emph{IEEE Transactions on Information Theory}, 62\penalty0
(6):\penalty0 3702--3720, 2016.
\newblock \doi{10.1109/TIT.2016.2548468}.
\bibitem[Wu and Yang(2019)]{wuyang}
Y.~Wu and P.~Yang.
\newblock {Chebyshev} polynomials, moment matching, and optimal estimation of
the unseen.
\newblock \emph{The Annals of Statistics}, 47\penalty0 (2):\penalty0 857--883,
2019.
\newblock \doi{10.1214/17-AOS1665}.
\bibitem[Zeng et~al.(2024)Zeng, Balakrishnan, Han, and Kennedy]{zeng}
Z.~Zeng, S.~Balakrishnan, Y.~Han, and E.~H. Kennedy.
\newblock Causal inference with high-dimensional discrete covariates, 2024.
\newblock \href{https://arxiv.org/abs/2405.00118v3}{arXiv:2405.00118v3}.
\bibitem[Zhang et~al.(2026)Zhang, Liu, and Zhang]{generalhoif}
Y.~Zhang, L.~Liu, and Z.~Zhang.
\newblock Higher-order debiased estimators for general treatment models.
\newblock \emph{Econometric Theory}, pages 1--43, 2026.
\newblock \doi{10.1017/S0266466626100516}.
\newblock First View; remark numbering refers to
\href{https://arxiv.org/abs/2606.01706v2}{arXiv:2606.01706v2}.
\end{thebibliography}
\end{document}