xref: /libCEED/rust/libceed-sys/c-src/backends/blocked/ceed-blocked-operator.c (revision 16911fdad6ca4b1dd37f9d3206958ee664667dbe)
14a2e7687Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC.
24a2e7687Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707.
34a2e7687Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details.
44a2e7687Sjeremylt //
54a2e7687Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software
64a2e7687Sjeremylt // libraries and APIs for efficient high-order finite element and spectral
74a2e7687Sjeremylt // element discretizations for exascale applications. For more information and
84a2e7687Sjeremylt // source code availability see http://github.com/ceed.
94a2e7687Sjeremylt //
104a2e7687Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC,
114a2e7687Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office
124a2e7687Sjeremylt // of Science and the National Nuclear Security Administration) responsible for
134a2e7687Sjeremylt // the planning and preparation of a capable exascale ecosystem, including
144a2e7687Sjeremylt // software, applications, hardware, advanced system engineering and early
154a2e7687Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative.
164a2e7687Sjeremylt 
174a2e7687Sjeremylt #include "ceed-blocked.h"
184a2e7687Sjeremylt #include "../ref/ceed-ref.h"
194a2e7687Sjeremylt 
204a2e7687Sjeremylt static int CeedOperatorDestroy_Blocked(CeedOperator op) {
214a2e7687Sjeremylt   int ierr;
224ce2993fSjeremylt   CeedOperator_Blocked *impl;
234ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
244a2e7687Sjeremylt 
254a2e7687Sjeremylt   for (CeedInt i=0; i<impl->numein+impl->numeout; i++) {
264a2e7687Sjeremylt     ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChk(ierr);
274a2e7687Sjeremylt     ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChk(ierr);
284a2e7687Sjeremylt   }
294a2e7687Sjeremylt   ierr = CeedFree(&impl->blkrestr); CeedChk(ierr);
304a2e7687Sjeremylt   ierr = CeedFree(&impl->evecs); CeedChk(ierr);
314a2e7687Sjeremylt   ierr = CeedFree(&impl->edata); CeedChk(ierr);
3216c359e6Sjeremylt   ierr = CeedFree(&impl->inputstate); CeedChk(ierr);
334a2e7687Sjeremylt 
34aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numein; i++) {
3591703d3fSjeremylt     ierr = CeedVectorDestroy(&impl->evecsin[i]); CeedChk(ierr);
36aedaa0e5Sjeremylt     ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChk(ierr);
374a2e7687Sjeremylt   }
3891703d3fSjeremylt   ierr = CeedFree(&impl->evecsin); CeedChk(ierr);
39aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsin); CeedChk(ierr);
404a2e7687Sjeremylt 
41aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numeout; i++) {
4291703d3fSjeremylt     ierr = CeedVectorDestroy(&impl->evecsout[i]); CeedChk(ierr);
43*16911fdaSjeremylt     if (!impl->identityqf) {
44aedaa0e5Sjeremylt       ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr);
45aedaa0e5Sjeremylt     }
46*16911fdaSjeremylt   }
4791703d3fSjeremylt   ierr = CeedFree(&impl->evecsout); CeedChk(ierr);
48aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsout); CeedChk(ierr);
494a2e7687Sjeremylt 
50fe2413ffSjeremylt   ierr = CeedFree(&impl); CeedChk(ierr);
514a2e7687Sjeremylt   return 0;
524a2e7687Sjeremylt }
534a2e7687Sjeremylt 
544a2e7687Sjeremylt /*
554a2e7687Sjeremylt   Setup infields or outfields
564a2e7687Sjeremylt  */
5789c6efa4Sjeremylt static int CeedOperatorSetupFields_Blocked(CeedQFunction qf,
5889c6efa4Sjeremylt     CeedOperator op, bool inOrOut,
594a2e7687Sjeremylt     CeedElemRestriction *blkrestr,
6091703d3fSjeremylt     CeedVector *fullevecs, CeedVector *evecs,
61aedaa0e5Sjeremylt     CeedVector *qvecs, CeedInt starte,
624a2e7687Sjeremylt     CeedInt numfields, CeedInt Q) {
634d537eeaSYohann   CeedInt dim, ierr, ncomp, size, P;
64aedaa0e5Sjeremylt   Ceed ceed;
65aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
66d1bcdac9Sjeremylt   CeedBasis basis;
67d1bcdac9Sjeremylt   CeedElemRestriction r;
68aedaa0e5Sjeremylt   CeedOperatorField *opfields;
69aedaa0e5Sjeremylt   CeedQFunctionField *qffields;
70fe2413ffSjeremylt   if (inOrOut) {
71aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, NULL, &opfields);
72fe2413ffSjeremylt     CeedChk(ierr);
73aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, NULL, &qffields);
74fe2413ffSjeremylt     CeedChk(ierr);
75fe2413ffSjeremylt   } else {
76aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, &opfields, NULL);
77fe2413ffSjeremylt     CeedChk(ierr);
78aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, &qffields, NULL);
79fe2413ffSjeremylt     CeedChk(ierr);
80fe2413ffSjeremylt   }
814a2e7687Sjeremylt   const CeedInt blksize = 8;
824a2e7687Sjeremylt 
834a2e7687Sjeremylt   // Loop over fields
844a2e7687Sjeremylt   for (CeedInt i=0; i<numfields; i++) {
85d1bcdac9Sjeremylt     CeedEvalMode emode;
86aedaa0e5Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChk(ierr);
874a2e7687Sjeremylt 
884a2e7687Sjeremylt     if (emode != CEED_EVAL_WEIGHT) {
89aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r);
90d1bcdac9Sjeremylt       CeedChk(ierr);
91fe2413ffSjeremylt       CeedElemRestriction_Ref *data;
92de686571SJeremy L Thompson       ierr = CeedElemRestrictionGetData(r, (void *)&data); CeedChk(ierr);
934ce2993fSjeremylt       Ceed ceed;
944ce2993fSjeremylt       ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChk(ierr);
958795c945Sjeremylt       CeedInt nelem, elemsize, nnodes;
964ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChk(ierr);
974ce2993fSjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChk(ierr);
988795c945Sjeremylt       ierr = CeedElemRestrictionGetNumNodes(r, &nnodes); CeedChk(ierr);
994ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChk(ierr);
1004ce2993fSjeremylt       ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize,
1018795c945Sjeremylt                                               blksize, nnodes, ncomp,
1024a2e7687Sjeremylt                                               CEED_MEM_HOST, CEED_COPY_VALUES,
103aedaa0e5Sjeremylt                                               data->indices, &blkrestr[i+starte]);
1044a2e7687Sjeremylt       CeedChk(ierr);
105aedaa0e5Sjeremylt       ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL,
10691703d3fSjeremylt                                              &fullevecs[i+starte]);
1074a2e7687Sjeremylt       CeedChk(ierr);
1084a2e7687Sjeremylt     }
1094a2e7687Sjeremylt 
1104a2e7687Sjeremylt     switch(emode) {
1114a2e7687Sjeremylt     case CEED_EVAL_NONE:
1124d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
1134d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
114aedaa0e5Sjeremylt       break;
115aedaa0e5Sjeremylt     case CEED_EVAL_INTERP:
1164d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
1174d1cd9fcSJeremy L Thompson       ierr = CeedElemRestrictionGetElementSize(r, &P);
1184d1cd9fcSJeremy L Thompson       CeedChk(ierr);
1194d537eeaSYohann       ierr = CeedVectorCreate(ceed, P*size*blksize, &evecs[i]); CeedChk(ierr);
1204d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
1214a2e7687Sjeremylt       break;
1224a2e7687Sjeremylt     case CEED_EVAL_GRAD:
123aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
1244d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
125d1bcdac9Sjeremylt       ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
1264d1cd9fcSJeremy L Thompson       ierr = CeedElemRestrictionGetElementSize(r, &P);
1274d1cd9fcSJeremy L Thompson       CeedChk(ierr);
1284d537eeaSYohann       ierr = CeedVectorCreate(ceed, P*size/dim*blksize, &evecs[i]); CeedChk(ierr);
1294d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
1304a2e7687Sjeremylt       break;
1314a2e7687Sjeremylt     case CEED_EVAL_WEIGHT: // Only on input fields
132aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
133aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChk(ierr);
134d1bcdac9Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
135aedaa0e5Sjeremylt                             CEED_EVAL_WEIGHT, NULL, qvecs[i]); CeedChk(ierr);
136aedaa0e5Sjeremylt 
1374a2e7687Sjeremylt       break;
1384a2e7687Sjeremylt     case CEED_EVAL_DIV:
1394d537eeaSYohann       break; // Not implemented
1404a2e7687Sjeremylt     case CEED_EVAL_CURL:
1414d537eeaSYohann       break; // Not implemented
1424a2e7687Sjeremylt     }
1434a2e7687Sjeremylt   }
1444a2e7687Sjeremylt   return 0;
1454a2e7687Sjeremylt }
1464a2e7687Sjeremylt 
1474a2e7687Sjeremylt /*
1484a2e7687Sjeremylt   CeedOperator needs to connect all the named fields (be they active or passive)
1494a2e7687Sjeremylt   to the named inputs and outputs of its CeedQFunction.
1504a2e7687Sjeremylt  */
1514a2e7687Sjeremylt static int CeedOperatorSetup_Blocked(CeedOperator op) {
1524a2e7687Sjeremylt   int ierr;
1534ce2993fSjeremylt   bool setupdone;
1544ce2993fSjeremylt   ierr = CeedOperatorGetSetupStatus(op, &setupdone); CeedChk(ierr);
1554ce2993fSjeremylt   if (setupdone) return 0;
156aedaa0e5Sjeremylt   Ceed ceed;
157aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
1584ce2993fSjeremylt   CeedOperator_Blocked *impl;
1594ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
1604ce2993fSjeremylt   CeedQFunction qf;
1614ce2993fSjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
1624ce2993fSjeremylt   CeedInt Q, numinputfields, numoutputfields;
1634ce2993fSjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
164*16911fdaSjeremylt   ierr = CeedQFunctionGetIdentityStatus(qf, &impl->identityqf); CeedChk(ierr);
1654a2e7687Sjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
1664a2e7687Sjeremylt   CeedChk(ierr);
167d1bcdac9Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
168d1bcdac9Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
169d1bcdac9Sjeremylt   CeedChk(ierr);
170d1bcdac9Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
171d1bcdac9Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
172d1bcdac9Sjeremylt   CeedChk(ierr);
1734a2e7687Sjeremylt 
1744a2e7687Sjeremylt   // Allocate
175aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr);
1764a2e7687Sjeremylt   CeedChk(ierr);
177aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs);
1784a2e7687Sjeremylt   CeedChk(ierr);
179aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata);
1804a2e7687Sjeremylt   CeedChk(ierr);
1814a2e7687Sjeremylt 
18216c359e6Sjeremylt   ierr = CeedCalloc(16, &impl->inputstate); CeedChk(ierr);
18391703d3fSjeremylt   ierr = CeedCalloc(16, &impl->evecsin); CeedChk(ierr);
18491703d3fSjeremylt   ierr = CeedCalloc(16, &impl->evecsout); CeedChk(ierr);
185aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsin); CeedChk(ierr);
186aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsout); CeedChk(ierr);
1874a2e7687Sjeremylt 
188aedaa0e5Sjeremylt   impl->numein = numinputfields; impl->numeout = numoutputfields;
189aedaa0e5Sjeremylt 
1904a2e7687Sjeremylt   // Set up infield and outfield pointer arrays
1914a2e7687Sjeremylt   // Infields
192aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 0, impl->blkrestr,
19391703d3fSjeremylt                                          impl->evecs, impl->evecsin,
19491703d3fSjeremylt                                          impl->qvecsin, 0,
195aedaa0e5Sjeremylt                                          numinputfields, Q);
1964a2e7687Sjeremylt   CeedChk(ierr);
1974a2e7687Sjeremylt   // Outfields
198aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 1, impl->blkrestr,
19991703d3fSjeremylt                                          impl->evecs, impl->evecsout,
20091703d3fSjeremylt                                          impl->qvecsout, numinputfields,
20191703d3fSjeremylt                                          numoutputfields, Q);
2024a2e7687Sjeremylt   CeedChk(ierr);
203aedaa0e5Sjeremylt 
204*16911fdaSjeremylt   // Identity QFunctions
205*16911fdaSjeremylt   if (impl->identityqf) {
206*16911fdaSjeremylt     CeedEvalMode inmode, outmode;
207*16911fdaSjeremylt     CeedQFunctionField *infields, *outfields;
208*16911fdaSjeremylt     ierr = CeedQFunctionGetFields(qf, &infields, &outfields); CeedChk(ierr);
209*16911fdaSjeremylt 
210*16911fdaSjeremylt     for (CeedInt i=0; i<numinputfields; i++) {
211*16911fdaSjeremylt       ierr = CeedQFunctionFieldGetEvalMode(infields[i], &inmode);
212*16911fdaSjeremylt       CeedChk(ierr);
213*16911fdaSjeremylt       ierr = CeedQFunctionFieldGetEvalMode(outfields[i], &outmode);
214*16911fdaSjeremylt       CeedChk(ierr);
215*16911fdaSjeremylt 
216*16911fdaSjeremylt       if (inmode == CEED_EVAL_NONE && outmode == CEED_EVAL_NONE)
217*16911fdaSjeremylt        // LCOV_EXCL_START
218*16911fdaSjeremylt         return CeedError(ceed, 1, "CEED_EVAL_NONE for a matching input and "
219*16911fdaSjeremylt                          "output does not make sense with identity QFunction");
220*16911fdaSjeremylt       // LCOV_EXCL_STOP
221*16911fdaSjeremylt 
222*16911fdaSjeremylt       ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr);
223*16911fdaSjeremylt       impl->qvecsout[i] = impl->qvecsin[i];
224*16911fdaSjeremylt     }
225*16911fdaSjeremylt   }
226*16911fdaSjeremylt 
2274ce2993fSjeremylt   ierr = CeedOperatorSetSetupDone(op); CeedChk(ierr);
2284a2e7687Sjeremylt 
2294a2e7687Sjeremylt   return 0;
2304a2e7687Sjeremylt }
2314a2e7687Sjeremylt 
2321d102b48SJeremy L Thompson // Setup Input fields
2331d102b48SJeremy L Thompson static inline int CeedOperatorSetupInputs_Blocked(CeedInt numinputfields,
2341d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
2351d102b48SJeremy L Thompson     CeedVector invec, bool skipactive, CeedOperator_Blocked *impl,
23689c6efa4Sjeremylt     CeedRequest *request) {
2371d102b48SJeremy L Thompson   CeedInt ierr;
238d1bcdac9Sjeremylt   CeedEvalMode emode;
239d1bcdac9Sjeremylt   CeedVector vec;
2401d102b48SJeremy L Thompson   CeedTransposeMode lmode;
24116c359e6Sjeremylt   uint64_t state;
2424a2e7687Sjeremylt 
2434a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
2441d102b48SJeremy L Thompson     // Get input vector
2451d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
2461d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
2471d102b48SJeremy L Thompson       if (skipactive)
2481d102b48SJeremy L Thompson         continue;
2491d102b48SJeremy L Thompson       else
2501d102b48SJeremy L Thompson         vec = invec;
2511d102b48SJeremy L Thompson     }
2521d102b48SJeremy L Thompson 
253d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
254d1bcdac9Sjeremylt     CeedChk(ierr);
2554a2e7687Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
2564a2e7687Sjeremylt     } else {
2574a2e7687Sjeremylt       // Restrict
25816c359e6Sjeremylt       ierr = CeedVectorGetState(vec, &state); CeedChk(ierr);
25989c6efa4Sjeremylt       if (state != impl->inputstate[i] || vec == invec) {
2601d102b48SJeremy L Thompson         ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode);
2611d102b48SJeremy L Thompson         CeedChk(ierr);
2624a2e7687Sjeremylt         ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE,
26389c6efa4Sjeremylt                                         lmode, vec, impl->evecs[i],
26489c6efa4Sjeremylt                                         request); CeedChk(ierr); CeedChk(ierr);
26516c359e6Sjeremylt         impl->inputstate[i] = state;
26616c359e6Sjeremylt       }
2674a2e7687Sjeremylt       // Get evec
2684a2e7687Sjeremylt       ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST,
2694a2e7687Sjeremylt                                     (const CeedScalar **) &impl->edata[i]);
2704a2e7687Sjeremylt       CeedChk(ierr);
2714a2e7687Sjeremylt     }
2724a2e7687Sjeremylt   }
2731d102b48SJeremy L Thompson   return 0;
2744a2e7687Sjeremylt }
2754a2e7687Sjeremylt 
2761d102b48SJeremy L Thompson // Input basis action
2771d102b48SJeremy L Thompson static inline int CeedOperatorInputBasis_Blocked(CeedInt e, CeedInt Q,
2781d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
2791d102b48SJeremy L Thompson     CeedInt numinputfields, CeedInt blksize, bool skipactive,
2801d102b48SJeremy L Thompson     CeedOperator_Blocked *impl) {
2811d102b48SJeremy L Thompson   CeedInt ierr;
2821d102b48SJeremy L Thompson   CeedInt dim, elemsize, size;
2831d102b48SJeremy L Thompson   CeedElemRestriction Erestrict;
2841d102b48SJeremy L Thompson   CeedEvalMode emode;
2851d102b48SJeremy L Thompson   CeedBasis basis;
2861d102b48SJeremy L Thompson 
2874a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
2881d102b48SJeremy L Thompson     // Skip active input
2891d102b48SJeremy L Thompson     if (skipactive) {
2901d102b48SJeremy L Thompson       CeedVector vec;
2911d102b48SJeremy L Thompson       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
2921d102b48SJeremy L Thompson       if (vec == CEED_VECTOR_ACTIVE)
2931d102b48SJeremy L Thompson         continue;
2941d102b48SJeremy L Thompson     }
2951d102b48SJeremy L Thompson 
2964d537eeaSYohann     // Get elemsize, emode, size
297d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict);
298d1bcdac9Sjeremylt     CeedChk(ierr);
299d1bcdac9Sjeremylt     ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
300d1bcdac9Sjeremylt     CeedChk(ierr);
301d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
302d1bcdac9Sjeremylt     CeedChk(ierr);
3034d537eeaSYohann     ierr = CeedQFunctionFieldGetSize(qfinputfields[i], &size); CeedChk(ierr);
3044a2e7687Sjeremylt     // Basis action
3054a2e7687Sjeremylt     switch(emode) {
3064a2e7687Sjeremylt     case CEED_EVAL_NONE:
307aedaa0e5Sjeremylt       ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
308aedaa0e5Sjeremylt                                 CEED_USE_POINTER,
3094d537eeaSYohann                                 &impl->edata[i][e*Q*size]); CeedChk(ierr);
3104a2e7687Sjeremylt       break;
3114a2e7687Sjeremylt     case CEED_EVAL_INTERP:
31289c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
31391703d3fSjeremylt       ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
314aedaa0e5Sjeremylt                                 CEED_USE_POINTER,
3154d537eeaSYohann                                 &impl->edata[i][e*elemsize*size]);
3164a2e7687Sjeremylt       CeedChk(ierr);
317d1bcdac9Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
31891703d3fSjeremylt                             CEED_EVAL_INTERP, impl->evecsin[i],
319aedaa0e5Sjeremylt                             impl->qvecsin[i]); CeedChk(ierr);
3204a2e7687Sjeremylt       break;
3214a2e7687Sjeremylt     case CEED_EVAL_GRAD:
32289c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
3234d537eeaSYohann       ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
32491703d3fSjeremylt       ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
325aedaa0e5Sjeremylt                                 CEED_USE_POINTER,
3264d537eeaSYohann                                 &impl->edata[i][e*elemsize*size/dim]);
3274a2e7687Sjeremylt       CeedChk(ierr);
328d1bcdac9Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
32991703d3fSjeremylt                             CEED_EVAL_GRAD, impl->evecsin[i],
330aedaa0e5Sjeremylt                             impl->qvecsin[i]); CeedChk(ierr);
3314a2e7687Sjeremylt       break;
3324a2e7687Sjeremylt     case CEED_EVAL_WEIGHT:
3334a2e7687Sjeremylt       break;  // No action
3344a2e7687Sjeremylt     case CEED_EVAL_DIV:
3351d102b48SJeremy L Thompson     case CEED_EVAL_CURL: {
3361d102b48SJeremy L Thompson       // LCOV_EXCL_START
3371d102b48SJeremy L Thompson       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis);
3381d102b48SJeremy L Thompson       CeedChk(ierr);
3391d102b48SJeremy L Thompson       Ceed ceed;
3401d102b48SJeremy L Thompson       ierr = CeedBasisGetCeed(basis, &ceed); CeedChk(ierr);
3411d102b48SJeremy L Thompson       return CeedError(ceed, 1, "Ceed evaluation mode not implemented");
3421d102b48SJeremy L Thompson       // LCOV_EXCL_STOP
3438c91a0c9SJeremy L Thompson       break; // Not implemented
3444a2e7687Sjeremylt     }
3454a2e7687Sjeremylt     }
34689c6efa4Sjeremylt   }
3471d102b48SJeremy L Thompson   return 0;
34889c6efa4Sjeremylt }
3494a2e7687Sjeremylt 
3501d102b48SJeremy L Thompson // Output basis action
3511d102b48SJeremy L Thompson static inline int CeedOperatorOutputBasis_Blocked(CeedInt e, CeedInt Q,
3521d102b48SJeremy L Thompson     CeedQFunctionField *qfoutputfields, CeedOperatorField *opoutputfields,
3531d102b48SJeremy L Thompson     CeedInt blksize, CeedInt numinputfields, CeedInt numoutputfields,
3541d102b48SJeremy L Thompson     CeedOperator op, CeedOperator_Blocked *impl) {
3551d102b48SJeremy L Thompson   CeedInt ierr;
3561d102b48SJeremy L Thompson   CeedInt dim, elemsize, size;
3571d102b48SJeremy L Thompson   CeedElemRestriction Erestrict;
3581d102b48SJeremy L Thompson   CeedEvalMode emode;
3591d102b48SJeremy L Thompson   CeedBasis basis;
3601d102b48SJeremy L Thompson 
3614a2e7687Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
3624d537eeaSYohann     // Get elemsize, emode, size
363d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict);
364d1bcdac9Sjeremylt     CeedChk(ierr);
36589c6efa4Sjeremylt     ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
36689c6efa4Sjeremylt     CeedChk(ierr);
367d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
368d1bcdac9Sjeremylt     CeedChk(ierr);
3694d537eeaSYohann     ierr = CeedQFunctionFieldGetSize(qfoutputfields[i], &size); CeedChk(ierr);
3704a2e7687Sjeremylt     // Basis action
3714a2e7687Sjeremylt     switch(emode) {
3724a2e7687Sjeremylt     case CEED_EVAL_NONE:
3734a2e7687Sjeremylt       break; // No action
3744a2e7687Sjeremylt     case CEED_EVAL_INTERP:
375d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
376d1bcdac9Sjeremylt       CeedChk(ierr);
37789c6efa4Sjeremylt       ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST,
37889c6efa4Sjeremylt                                 CEED_USE_POINTER,
3794d537eeaSYohann                                 &impl->edata[i + numinputfields][e*elemsize*size]);
38089c6efa4Sjeremylt       CeedChk(ierr);
381aedaa0e5Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
382aedaa0e5Sjeremylt                             CEED_EVAL_INTERP, impl->qvecsout[i],
38391703d3fSjeremylt                             impl->evecsout[i]); CeedChk(ierr);
3844a2e7687Sjeremylt       break;
3854a2e7687Sjeremylt     case CEED_EVAL_GRAD:
386d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
387d1bcdac9Sjeremylt       CeedChk(ierr);
3884d537eeaSYohann       ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
38989c6efa4Sjeremylt       ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST,
39089c6efa4Sjeremylt                                 CEED_USE_POINTER,
3914d537eeaSYohann                                 &impl->edata[i + numinputfields][e*elemsize*size/dim]);
39289c6efa4Sjeremylt       CeedChk(ierr);
393d1bcdac9Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
394aedaa0e5Sjeremylt                             CEED_EVAL_GRAD, impl->qvecsout[i],
39591703d3fSjeremylt                             impl->evecsout[i]); CeedChk(ierr);
3964a2e7687Sjeremylt       break;
3974ce2993fSjeremylt     case CEED_EVAL_WEIGHT: {
398c042f62fSJeremy L Thompson       // LCOV_EXCL_START
3994ce2993fSjeremylt       Ceed ceed;
4004ce2993fSjeremylt       ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
4011d102b48SJeremy L Thompson       return CeedError(ceed, 1, "CEED_EVAL_WEIGHT cannot be an output "
4021d102b48SJeremy L Thompson                        "evaluation mode");
403c042f62fSJeremy L Thompson       // LCOV_EXCL_STOP
4044a2e7687Sjeremylt       break; // Should not occur
4054ce2993fSjeremylt     }
4064a2e7687Sjeremylt     case CEED_EVAL_DIV:
4071d102b48SJeremy L Thompson     case CEED_EVAL_CURL: {
4081d102b48SJeremy L Thompson       // LCOV_EXCL_START
4091d102b48SJeremy L Thompson       Ceed ceed;
4101d102b48SJeremy L Thompson       ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
4111d102b48SJeremy L Thompson       return CeedError(ceed, 1, "Ceed evaluation mode not implemented");
4121d102b48SJeremy L Thompson       // LCOV_EXCL_STOP
4138c91a0c9SJeremy L Thompson       break; // Not implemented
4144a2e7687Sjeremylt     }
41589c6efa4Sjeremylt     }
41689c6efa4Sjeremylt   }
4171d102b48SJeremy L Thompson   return 0;
4181d102b48SJeremy L Thompson }
4191d102b48SJeremy L Thompson 
4201d102b48SJeremy L Thompson // Restore Inputs
4211d102b48SJeremy L Thompson static inline int CeedOperatorRestoreInputs_Blocked(CeedInt numinputfields,
4221d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
4231d102b48SJeremy L Thompson     bool skipactive, CeedOperator_Blocked *impl) {
4241d102b48SJeremy L Thompson   CeedInt ierr;
4251d102b48SJeremy L Thompson   CeedEvalMode emode;
4261d102b48SJeremy L Thompson 
4271d102b48SJeremy L Thompson   for (CeedInt i=0; i<numinputfields; i++) {
4281d102b48SJeremy L Thompson     // Skip active inputs
4291d102b48SJeremy L Thompson     if (skipactive) {
4301d102b48SJeremy L Thompson       CeedVector vec;
4311d102b48SJeremy L Thompson       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
4321d102b48SJeremy L Thompson       if (vec == CEED_VECTOR_ACTIVE)
4331d102b48SJeremy L Thompson         continue;
4341d102b48SJeremy L Thompson     }
4351d102b48SJeremy L Thompson     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
4361d102b48SJeremy L Thompson     CeedChk(ierr);
4371d102b48SJeremy L Thompson     if (emode == CEED_EVAL_WEIGHT) { // Skip
4381d102b48SJeremy L Thompson     } else {
4391d102b48SJeremy L Thompson       ierr = CeedVectorRestoreArrayRead(impl->evecs[i],
4401d102b48SJeremy L Thompson                                         (const CeedScalar **) &impl->edata[i]);
4411d102b48SJeremy L Thompson       CeedChk(ierr);
4421d102b48SJeremy L Thompson     }
4431d102b48SJeremy L Thompson   }
4441d102b48SJeremy L Thompson   return 0;
4451d102b48SJeremy L Thompson }
4461d102b48SJeremy L Thompson 
4471d102b48SJeremy L Thompson // Apply Ceed Operator
4481d102b48SJeremy L Thompson static int CeedOperatorApply_Blocked(CeedOperator op, CeedVector invec,
4491d102b48SJeremy L Thompson                                      CeedVector outvec,
4501d102b48SJeremy L Thompson                                      CeedRequest *request) {
4511d102b48SJeremy L Thompson   int ierr;
4521d102b48SJeremy L Thompson   CeedOperator_Blocked *impl;
4531d102b48SJeremy L Thompson   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
4541d102b48SJeremy L Thompson   const CeedInt blksize = 8;
4551d102b48SJeremy L Thompson   CeedInt Q, numinputfields, numoutputfields, numelements, size;
4561d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr);
4571d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
4581d102b48SJeremy L Thompson   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
4591d102b48SJeremy L Thompson   CeedQFunction qf;
4601d102b48SJeremy L Thompson   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
4611d102b48SJeremy L Thompson   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
4621d102b48SJeremy L Thompson   CeedChk(ierr);
4631d102b48SJeremy L Thompson   CeedTransposeMode lmode;
4641d102b48SJeremy L Thompson   CeedOperatorField *opinputfields, *opoutputfields;
4651d102b48SJeremy L Thompson   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
4661d102b48SJeremy L Thompson   CeedChk(ierr);
4671d102b48SJeremy L Thompson   CeedQFunctionField *qfinputfields, *qfoutputfields;
4681d102b48SJeremy L Thompson   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
4691d102b48SJeremy L Thompson   CeedChk(ierr);
4701d102b48SJeremy L Thompson   CeedEvalMode emode;
4711d102b48SJeremy L Thompson   CeedVector vec;
4721d102b48SJeremy L Thompson 
4731d102b48SJeremy L Thompson   // Setup
4741d102b48SJeremy L Thompson   ierr = CeedOperatorSetup_Blocked(op); CeedChk(ierr);
4751d102b48SJeremy L Thompson 
4761d102b48SJeremy L Thompson   // Input Evecs and Restriction
4771d102b48SJeremy L Thompson   ierr = CeedOperatorSetupInputs_Blocked(numinputfields, qfinputfields,
4787f823360Sjeremylt                                          opinputfields, invec, false, impl,
4797f823360Sjeremylt                                          request); CeedChk(ierr);
4801d102b48SJeremy L Thompson 
4811d102b48SJeremy L Thompson   // Output Evecs
4821d102b48SJeremy L Thompson   for (CeedInt i=0; i<numoutputfields; i++) {
4831d102b48SJeremy L Thompson     ierr = CeedVectorGetArray(impl->evecs[i+impl->numein], CEED_MEM_HOST,
4841d102b48SJeremy L Thompson                               &impl->edata[i + numinputfields]); CeedChk(ierr);
4851d102b48SJeremy L Thompson   }
4861d102b48SJeremy L Thompson 
4871d102b48SJeremy L Thompson   // Loop through elements
4881d102b48SJeremy L Thompson   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
4891d102b48SJeremy L Thompson     // Output pointers
4901d102b48SJeremy L Thompson     for (CeedInt i=0; i<numoutputfields; i++) {
4911d102b48SJeremy L Thompson       ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
4921d102b48SJeremy L Thompson       CeedChk(ierr);
4931d102b48SJeremy L Thompson       if (emode == CEED_EVAL_NONE) {
4941d102b48SJeremy L Thompson         ierr = CeedQFunctionFieldGetSize(qfoutputfields[i], &size);
4951d102b48SJeremy L Thompson         CeedChk(ierr);
4961d102b48SJeremy L Thompson         ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST,
4971d102b48SJeremy L Thompson                                   CEED_USE_POINTER,
4981d102b48SJeremy L Thompson                                   &impl->edata[i + numinputfields][e*Q*size]);
4991d102b48SJeremy L Thompson         CeedChk(ierr);
5001d102b48SJeremy L Thompson       }
5011d102b48SJeremy L Thompson     }
5021d102b48SJeremy L Thompson 
503*16911fdaSjeremylt     // Input basis apply
504*16911fdaSjeremylt     ierr = CeedOperatorInputBasis_Blocked(e, Q, qfinputfields, opinputfields,
505*16911fdaSjeremylt                                           numinputfields, blksize, false, impl);
506*16911fdaSjeremylt     CeedChk(ierr);
507*16911fdaSjeremylt 
5081d102b48SJeremy L Thompson     // Q function
509*16911fdaSjeremylt     if (!impl->identityqf) {
5101d102b48SJeremy L Thompson       ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
5111d102b48SJeremy L Thompson       CeedChk(ierr);
512*16911fdaSjeremylt     }
5131d102b48SJeremy L Thompson 
5141d102b48SJeremy L Thompson     // Output basis apply
5151d102b48SJeremy L Thompson     ierr = CeedOperatorOutputBasis_Blocked(e, Q, qfoutputfields, opoutputfields,
5167f823360Sjeremylt                                            blksize, numinputfields,
5177f823360Sjeremylt                                            numoutputfields, op, impl);
5181d102b48SJeremy L Thompson     CeedChk(ierr);
5191d102b48SJeremy L Thompson   }
52089c6efa4Sjeremylt 
52189c6efa4Sjeremylt   // Zero lvecs
52289c6efa4Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
52389c6efa4Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
52489c6efa4Sjeremylt     if (vec == CEED_VECTOR_ACTIVE) {
52589c6efa4Sjeremylt       if (!impl->add) {
52689c6efa4Sjeremylt         vec = outvec;
52789c6efa4Sjeremylt         ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr);
52889c6efa4Sjeremylt       }
52989c6efa4Sjeremylt     } else {
53089c6efa4Sjeremylt       ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr);
53189c6efa4Sjeremylt     }
53289c6efa4Sjeremylt   }
53389c6efa4Sjeremylt   impl->add = false;
53489c6efa4Sjeremylt 
53589c6efa4Sjeremylt   // Output restriction
53689c6efa4Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
53789c6efa4Sjeremylt     // Restore evec
53889c6efa4Sjeremylt     ierr = CeedVectorRestoreArray(impl->evecs[i+impl->numein],
53989c6efa4Sjeremylt                                   &impl->edata[i + numinputfields]); CeedChk(ierr);
540d1bcdac9Sjeremylt     // Get output vector
541d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
54289c6efa4Sjeremylt     // Active
543d1bcdac9Sjeremylt     if (vec == CEED_VECTOR_ACTIVE)
544d1bcdac9Sjeremylt       vec = outvec;
5454a2e7687Sjeremylt     // Restrict
54689c6efa4Sjeremylt     ierr = CeedOperatorFieldGetLMode(opoutputfields[i], &lmode); CeedChk(ierr);
54789c6efa4Sjeremylt     ierr = CeedElemRestrictionApply(impl->blkrestr[i+impl->numein], CEED_TRANSPOSE,
54889c6efa4Sjeremylt                                     lmode, impl->evecs[i+impl->numein], vec,
54989c6efa4Sjeremylt                                     request); CeedChk(ierr);
55089c6efa4Sjeremylt 
5514a2e7687Sjeremylt   }
5524a2e7687Sjeremylt 
5534a2e7687Sjeremylt   // Restore input arrays
5541d102b48SJeremy L Thompson   ierr = CeedOperatorRestoreInputs_Blocked(numinputfields, qfinputfields,
5557f823360Sjeremylt          opinputfields, false, impl); CeedChk(ierr);
5561d102b48SJeremy L Thompson 
5571d102b48SJeremy L Thompson   return 0;
5581d102b48SJeremy L Thompson }
5591d102b48SJeremy L Thompson 
5601d102b48SJeremy L Thompson // Assemble Linear QFunction
5611d102b48SJeremy L Thompson static int CeedOperatorAssembleLinearQFunction_Blocked(CeedOperator op,
5621d102b48SJeremy L Thompson     CeedVector *assembled, CeedElemRestriction *rstr, CeedRequest *request) {
5631d102b48SJeremy L Thompson   int ierr;
5641d102b48SJeremy L Thompson   CeedOperator_Blocked *impl;
5651d102b48SJeremy L Thompson   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
5661d102b48SJeremy L Thompson   const CeedInt blksize = 8;
5671d102b48SJeremy L Thompson   CeedInt Q, numinputfields, numoutputfields, numelements, size;
5681d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr);
5691d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
5701d102b48SJeremy L Thompson   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
5711d102b48SJeremy L Thompson   CeedQFunction qf;
5721d102b48SJeremy L Thompson   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
5731d102b48SJeremy L Thompson   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
5741d102b48SJeremy L Thompson   CeedChk(ierr);
5751d102b48SJeremy L Thompson   CeedOperatorField *opinputfields, *opoutputfields;
5761d102b48SJeremy L Thompson   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
5771d102b48SJeremy L Thompson   CeedChk(ierr);
5781d102b48SJeremy L Thompson   CeedQFunctionField *qfinputfields, *qfoutputfields;
5791d102b48SJeremy L Thompson   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
5801d102b48SJeremy L Thompson   CeedChk(ierr);
5811d102b48SJeremy L Thompson   CeedVector vec, lvec;
5821d102b48SJeremy L Thompson   CeedInt numactivein = 0, numactiveout = 0;
58342ea3801Sjeremylt   CeedVector *activein = NULL;
5841d102b48SJeremy L Thompson   CeedScalar *a, *tmp;
5851d102b48SJeremy L Thompson   Ceed ceed;
5861d102b48SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
5871d102b48SJeremy L Thompson 
5881d102b48SJeremy L Thompson   // Setup
5891d102b48SJeremy L Thompson   ierr = CeedOperatorSetup_Blocked(op); CeedChk(ierr);
5901d102b48SJeremy L Thompson 
591*16911fdaSjeremylt   // Check for identity
592*16911fdaSjeremylt   if (impl->identityqf)
593*16911fdaSjeremylt    // LCOV_EXCL_START
594*16911fdaSjeremylt     return CeedError(ceed, 1, "Assembling identity qfunction does not make sense");
595*16911fdaSjeremylt   // LCOV_EXCL_STOP
596*16911fdaSjeremylt 
5971d102b48SJeremy L Thompson   // Input Evecs and Restriction
5981d102b48SJeremy L Thompson   ierr = CeedOperatorSetupInputs_Blocked(numinputfields, qfinputfields,
5991d102b48SJeremy L Thompson                                          opinputfields, NULL, true, impl,
6001d102b48SJeremy L Thompson                                          request); CeedChk(ierr);
6011d102b48SJeremy L Thompson 
6021d102b48SJeremy L Thompson   // Count number of active input fields
6034a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
6041d102b48SJeremy L Thompson     // Get input vector
6051d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
6061d102b48SJeremy L Thompson     // Check if active input
6071d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
6081d102b48SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qfinputfields[i], &size); CeedChk(ierr);
6091d102b48SJeremy L Thompson       ierr = CeedVectorSetValue(impl->qvecsin[i], 0.0); CeedChk(ierr);
6101d102b48SJeremy L Thompson       ierr = CeedVectorGetArray(impl->qvecsin[i], CEED_MEM_HOST, &tmp);
611d1bcdac9Sjeremylt       CeedChk(ierr);
6121d102b48SJeremy L Thompson       ierr = CeedRealloc(numactivein + size, &activein); CeedChk(ierr);
6131d102b48SJeremy L Thompson       for (CeedInt field=0; field<size; field++) {
61442ea3801Sjeremylt         ierr = CeedVectorCreate(ceed, Q*blksize, &activein[numactivein+field]);
61542ea3801Sjeremylt         CeedChk(ierr);
61642ea3801Sjeremylt         ierr = CeedVectorSetArray(activein[numactivein+field], CEED_MEM_HOST,
61742ea3801Sjeremylt                                   CEED_USE_POINTER, &tmp[field*Q*blksize]);
618112e3f70Sjeremylt         CeedChk(ierr);
6191d102b48SJeremy L Thompson       }
6201d102b48SJeremy L Thompson       numactivein += size;
6211d102b48SJeremy L Thompson       ierr = CeedVectorRestoreArray(impl->qvecsin[i], &tmp); CeedChk(ierr);
6221d102b48SJeremy L Thompson     }
6231d102b48SJeremy L Thompson   }
6241d102b48SJeremy L Thompson 
6251d102b48SJeremy L Thompson   // Count number of active output fields
6261d102b48SJeremy L Thompson   for (CeedInt i=0; i<numoutputfields; i++) {
6271d102b48SJeremy L Thompson     // Get output vector
6281d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
6291d102b48SJeremy L Thompson     // Check if active output
6301d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
6311d102b48SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qfoutputfields[i], &size); CeedChk(ierr);
6321d102b48SJeremy L Thompson       numactiveout += size;
6331d102b48SJeremy L Thompson     }
6341d102b48SJeremy L Thompson   }
6351d102b48SJeremy L Thompson 
6361d102b48SJeremy L Thompson   // Check sizes
6371d102b48SJeremy L Thompson   if (!numactivein || !numactiveout)
6381d102b48SJeremy L Thompson     // LCOV_EXCL_START
6391d102b48SJeremy L Thompson     return CeedError(ceed, 1, "Cannot assemble QFunction without active inputs "
6401d102b48SJeremy L Thompson                      "and outputs");
6411d102b48SJeremy L Thompson   // LCOV_EXCL_STOP
6421d102b48SJeremy L Thompson 
6431d102b48SJeremy L Thompson   // Setup lvec
6441d102b48SJeremy L Thompson   ierr = CeedVectorCreate(ceed, nblks*blksize*Q*numactivein*numactiveout,
6451d102b48SJeremy L Thompson                           &lvec); CeedChk(ierr);
6461d102b48SJeremy L Thompson   ierr = CeedVectorGetArray(lvec, CEED_MEM_HOST, &a); CeedChk(ierr);
6471d102b48SJeremy L Thompson 
6481d102b48SJeremy L Thompson   // Create output restriction
6491d102b48SJeremy L Thompson   ierr = CeedElemRestrictionCreateIdentity(ceed, numelements, Q,
6507f823360Sjeremylt          numelements*Q, numactivein*numactiveout, rstr); CeedChk(ierr);
6511d102b48SJeremy L Thompson   // Create assembled vector
6521d102b48SJeremy L Thompson   ierr = CeedVectorCreate(ceed, numelements*Q*numactivein*numactiveout,
6531d102b48SJeremy L Thompson                           assembled); CeedChk(ierr);
6541d102b48SJeremy L Thompson 
6551d102b48SJeremy L Thompson   // Loop through elements
6561d102b48SJeremy L Thompson   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
6571d102b48SJeremy L Thompson     // Input basis apply
6581d102b48SJeremy L Thompson     ierr = CeedOperatorInputBasis_Blocked(e, Q, qfinputfields, opinputfields,
6591d102b48SJeremy L Thompson                                           numinputfields, blksize, true, impl);
6601d102b48SJeremy L Thompson     CeedChk(ierr);
6611d102b48SJeremy L Thompson 
6621d102b48SJeremy L Thompson     // Assemble QFunction
6631d102b48SJeremy L Thompson     for (CeedInt in=0; in<numactivein; in++) {
6641d102b48SJeremy L Thompson       // Set Inputs
66542ea3801Sjeremylt       ierr = CeedVectorSetValue(activein[in], 1.0); CeedChk(ierr);
66642ea3801Sjeremylt       if (numactivein > 1) {
66742ea3801Sjeremylt         ierr = CeedVectorSetValue(activein[(in+numactivein-1)%numactivein],
66842ea3801Sjeremylt                                   0.0); CeedChk(ierr);
66942ea3801Sjeremylt       }
6701d102b48SJeremy L Thompson       // Set Outputs
6711d102b48SJeremy L Thompson       for (CeedInt out=0; out<numoutputfields; out++) {
6721d102b48SJeremy L Thompson         // Get output vector
6731d102b48SJeremy L Thompson         ierr = CeedOperatorFieldGetVector(opoutputfields[out], &vec);
6741d102b48SJeremy L Thompson         CeedChk(ierr);
6751d102b48SJeremy L Thompson         // Check if active output
6761d102b48SJeremy L Thompson         if (vec == CEED_VECTOR_ACTIVE) {
6771d102b48SJeremy L Thompson           CeedVectorSetArray(impl->qvecsout[out], CEED_MEM_HOST,
6781d102b48SJeremy L Thompson                              CEED_USE_POINTER, a); CeedChk(ierr);
6791d102b48SJeremy L Thompson           ierr = CeedQFunctionFieldGetSize(qfoutputfields[out], &size);
6801d102b48SJeremy L Thompson           CeedChk(ierr);
6811d102b48SJeremy L Thompson           a += size*Q*blksize; // Advance the pointer by the size of the output
6821d102b48SJeremy L Thompson         }
6831d102b48SJeremy L Thompson       }
6841d102b48SJeremy L Thompson       // Apply QFunction
6851d102b48SJeremy L Thompson       ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
686d1bcdac9Sjeremylt       CeedChk(ierr);
6874a2e7687Sjeremylt     }
6884a2e7687Sjeremylt   }
6894a2e7687Sjeremylt 
6901d102b48SJeremy L Thompson   // Un-set output Qvecs to prevent accidental overwrite of Assembled
6911d102b48SJeremy L Thompson   for (CeedInt out=0; out<numoutputfields; out++) {
6921d102b48SJeremy L Thompson     // Get output vector
6931d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opoutputfields[out], &vec);
6941d102b48SJeremy L Thompson     CeedChk(ierr);
6951d102b48SJeremy L Thompson     // Check if active output
6961d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
6971d102b48SJeremy L Thompson       CeedVectorSetArray(impl->qvecsout[out], CEED_MEM_HOST, CEED_COPY_VALUES,
6981d102b48SJeremy L Thompson                          NULL); CeedChk(ierr);
6991d102b48SJeremy L Thompson     }
7001d102b48SJeremy L Thompson   }
7011d102b48SJeremy L Thompson 
7021d102b48SJeremy L Thompson   // Restore input arrays
7031d102b48SJeremy L Thompson   ierr = CeedOperatorRestoreInputs_Blocked(numinputfields, qfinputfields,
7047f823360Sjeremylt          opinputfields, true, impl); CeedChk(ierr);
7051d102b48SJeremy L Thompson 
7061d102b48SJeremy L Thompson   // Output blocked restriction
7071d102b48SJeremy L Thompson   ierr = CeedVectorRestoreArray(lvec, &a); CeedChk(ierr);
7081d102b48SJeremy L Thompson   ierr = CeedVectorSetValue(*assembled, 0.0); CeedChk(ierr);
7091d102b48SJeremy L Thompson   CeedElemRestriction blkrstr;
7101d102b48SJeremy L Thompson   ierr = CeedElemRestrictionCreateBlocked(ceed, numelements, Q, blksize,
7111d102b48SJeremy L Thompson                                           numelements*Q,
7121d102b48SJeremy L Thompson                                           numactivein*numactiveout,
7131d102b48SJeremy L Thompson                                           CEED_MEM_HOST, CEED_COPY_VALUES,
7141d102b48SJeremy L Thompson                                           NULL, &blkrstr); CeedChk(ierr);
7151d102b48SJeremy L Thompson   ierr = CeedElemRestrictionApply(blkrstr, CEED_TRANSPOSE, CEED_NOTRANSPOSE,
7161d102b48SJeremy L Thompson                                   lvec, *assembled, request); CeedChk(ierr);
7171d102b48SJeremy L Thompson 
7181d102b48SJeremy L Thompson   // Cleanup
71942ea3801Sjeremylt   for (CeedInt i=0; i<numactivein; i++) {
72042ea3801Sjeremylt     ierr = CeedVectorDestroy(&activein[i]); CeedChk(ierr);
72142ea3801Sjeremylt   }
7221d102b48SJeremy L Thompson   ierr = CeedFree(&activein); CeedChk(ierr);
7231d102b48SJeremy L Thompson   ierr = CeedVectorDestroy(&lvec); CeedChk(ierr);
7241d102b48SJeremy L Thompson   ierr = CeedElemRestrictionDestroy(&blkrstr); CeedChk(ierr);
7251d102b48SJeremy L Thompson 
7264a2e7687Sjeremylt   return 0;
7274a2e7687Sjeremylt }
7284a2e7687Sjeremylt 
7294a2e7687Sjeremylt int CeedOperatorCreate_Blocked(CeedOperator op) {
7304a2e7687Sjeremylt   int ierr;
731fe2413ffSjeremylt   Ceed ceed;
732fe2413ffSjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
7334ce2993fSjeremylt   CeedOperator_Blocked *impl;
7344a2e7687Sjeremylt 
7354a2e7687Sjeremylt   ierr = CeedCalloc(1, &impl); CeedChk(ierr);
736de686571SJeremy L Thompson   ierr = CeedOperatorSetData(op, (void *)&impl); CeedChk(ierr);
737fe2413ffSjeremylt 
7381d102b48SJeremy L Thompson   ierr = CeedSetBackendFunction(ceed, "Operator", op, "AssembleLinearQFunction",
7391d102b48SJeremy L Thompson                                 CeedOperatorAssembleLinearQFunction_Blocked);
7401d102b48SJeremy L Thompson   CeedChk(ierr);
741fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Apply",
742fe2413ffSjeremylt                                 CeedOperatorApply_Blocked); CeedChk(ierr);
743fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy",
744fe2413ffSjeremylt                                 CeedOperatorDestroy_Blocked); CeedChk(ierr);
7454a2e7687Sjeremylt   return 0;
7464a2e7687Sjeremylt }
747