xref: /libCEED/rust/libceed-sys/c-src/backends/blocked/ceed-blocked-operator.c (revision de686571da193802df261a10fbc2ea743b9830da)
14a2e7687Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC.
24a2e7687Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707.
34a2e7687Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details.
44a2e7687Sjeremylt //
54a2e7687Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software
64a2e7687Sjeremylt // libraries and APIs for efficient high-order finite element and spectral
74a2e7687Sjeremylt // element discretizations for exascale applications. For more information and
84a2e7687Sjeremylt // source code availability see http://github.com/ceed.
94a2e7687Sjeremylt //
104a2e7687Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC,
114a2e7687Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office
124a2e7687Sjeremylt // of Science and the National Nuclear Security Administration) responsible for
134a2e7687Sjeremylt // the planning and preparation of a capable exascale ecosystem, including
144a2e7687Sjeremylt // software, applications, hardware, advanced system engineering and early
154a2e7687Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative.
164a2e7687Sjeremylt 
174a2e7687Sjeremylt #include <string.h>
184a2e7687Sjeremylt #include "ceed-blocked.h"
194a2e7687Sjeremylt #include "../ref/ceed-ref.h"
204a2e7687Sjeremylt 
214a2e7687Sjeremylt static int CeedOperatorDestroy_Blocked(CeedOperator op) {
224a2e7687Sjeremylt   int ierr;
234ce2993fSjeremylt   CeedOperator_Blocked *impl;
244ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
254a2e7687Sjeremylt 
264a2e7687Sjeremylt   for (CeedInt i=0; i<impl->numein+impl->numeout; i++) {
274a2e7687Sjeremylt     ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChk(ierr);
284a2e7687Sjeremylt     ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChk(ierr);
294a2e7687Sjeremylt   }
304a2e7687Sjeremylt   ierr = CeedFree(&impl->blkrestr); CeedChk(ierr);
314a2e7687Sjeremylt   ierr = CeedFree(&impl->evecs); CeedChk(ierr);
324a2e7687Sjeremylt   ierr = CeedFree(&impl->edata); CeedChk(ierr);
3316c359e6Sjeremylt   ierr = CeedFree(&impl->inputstate); CeedChk(ierr);
344a2e7687Sjeremylt 
35aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numein; i++) {
3691703d3fSjeremylt     ierr = CeedVectorDestroy(&impl->evecsin[i]); CeedChk(ierr);
37aedaa0e5Sjeremylt     ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChk(ierr);
384a2e7687Sjeremylt   }
3991703d3fSjeremylt   ierr = CeedFree(&impl->evecsin); CeedChk(ierr);
40aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsin); CeedChk(ierr);
414a2e7687Sjeremylt 
42aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numeout; i++) {
4391703d3fSjeremylt     ierr = CeedVectorDestroy(&impl->evecsout[i]); CeedChk(ierr);
44aedaa0e5Sjeremylt     ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr);
45aedaa0e5Sjeremylt   }
4691703d3fSjeremylt   ierr = CeedFree(&impl->evecsout); CeedChk(ierr);
47aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsout); CeedChk(ierr);
484a2e7687Sjeremylt 
49fe2413ffSjeremylt   ierr = CeedFree(&impl); CeedChk(ierr);
504a2e7687Sjeremylt   return 0;
514a2e7687Sjeremylt }
524a2e7687Sjeremylt 
534a2e7687Sjeremylt /*
544a2e7687Sjeremylt   Setup infields or outfields
554a2e7687Sjeremylt  */
56fe2413ffSjeremylt static int CeedOperatorSetupFields_Blocked(CeedQFunction qf, CeedOperator op,
57fe2413ffSjeremylt     bool inOrOut,
584a2e7687Sjeremylt     CeedElemRestriction *blkrestr,
5991703d3fSjeremylt     CeedVector *fullevecs, CeedVector *evecs,
60aedaa0e5Sjeremylt     CeedVector *qvecs, CeedInt starte,
614a2e7687Sjeremylt     CeedInt numfields, CeedInt Q) {
624d1cd9fcSJeremy L Thompson   CeedInt dim, ierr, ncomp, P;
63aedaa0e5Sjeremylt   Ceed ceed;
64aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
65d1bcdac9Sjeremylt   CeedBasis basis;
66d1bcdac9Sjeremylt   CeedElemRestriction r;
67aedaa0e5Sjeremylt   CeedOperatorField *opfields;
68aedaa0e5Sjeremylt   CeedQFunctionField *qffields;
69fe2413ffSjeremylt   if (inOrOut) {
70aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, NULL, &opfields);
71fe2413ffSjeremylt     CeedChk(ierr);
72aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, NULL, &qffields);
73fe2413ffSjeremylt     CeedChk(ierr);
74fe2413ffSjeremylt   } else {
75aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, &opfields, NULL);
76fe2413ffSjeremylt     CeedChk(ierr);
77aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, &qffields, NULL);
78fe2413ffSjeremylt     CeedChk(ierr);
79fe2413ffSjeremylt   }
804a2e7687Sjeremylt   const CeedInt blksize = 8;
814a2e7687Sjeremylt 
824a2e7687Sjeremylt   // Loop over fields
834a2e7687Sjeremylt   for (CeedInt i=0; i<numfields; i++) {
84d1bcdac9Sjeremylt     CeedEvalMode emode;
85aedaa0e5Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChk(ierr);
864a2e7687Sjeremylt 
874a2e7687Sjeremylt     if (emode != CEED_EVAL_WEIGHT) {
88aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r);
89d1bcdac9Sjeremylt       CeedChk(ierr);
90fe2413ffSjeremylt       CeedElemRestriction_Ref *data;
91*de686571SJeremy L Thompson       ierr = CeedElemRestrictionGetData(r, (void *)&data); CeedChk(ierr);
924ce2993fSjeremylt       Ceed ceed;
934ce2993fSjeremylt       ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChk(ierr);
944d1cd9fcSJeremy L Thompson       CeedInt nelem, elemsize, ndof;
954ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChk(ierr);
964ce2993fSjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChk(ierr);
974ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumDoF(r, &ndof); CeedChk(ierr);
984ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChk(ierr);
994ce2993fSjeremylt       ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize,
1004ce2993fSjeremylt                                               blksize, ndof, ncomp,
1014a2e7687Sjeremylt                                               CEED_MEM_HOST, CEED_COPY_VALUES,
102aedaa0e5Sjeremylt                                               data->indices, &blkrestr[i+starte]);
1034a2e7687Sjeremylt       CeedChk(ierr);
104aedaa0e5Sjeremylt       ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL,
10591703d3fSjeremylt                                              &fullevecs[i+starte]);
1064a2e7687Sjeremylt       CeedChk(ierr);
1074a2e7687Sjeremylt     }
1084a2e7687Sjeremylt 
1094a2e7687Sjeremylt     switch(emode) {
1104a2e7687Sjeremylt     case CEED_EVAL_NONE:
111aedaa0e5Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp);
112d1bcdac9Sjeremylt       CeedChk(ierr);
11391703d3fSjeremylt       ierr = CeedVectorCreate(ceed, Q*ncomp*blksize, &qvecs[i]); CeedChk(ierr);
114aedaa0e5Sjeremylt       break;
115aedaa0e5Sjeremylt     case CEED_EVAL_INTERP:
116aedaa0e5Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp);
117aedaa0e5Sjeremylt       CeedChk(ierr);
1184d1cd9fcSJeremy L Thompson       ierr = CeedElemRestrictionGetElementSize(r, &P);
1194d1cd9fcSJeremy L Thompson       CeedChk(ierr);
1204d1cd9fcSJeremy L Thompson       ierr = CeedVectorCreate(ceed, P*ncomp*blksize, &evecs[i]); CeedChk(ierr);
121aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*ncomp*blksize, &qvecs[i]); CeedChk(ierr);
1224a2e7687Sjeremylt       break;
1234a2e7687Sjeremylt     case CEED_EVAL_GRAD:
124aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
125aedaa0e5Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp);
126*de686571SJeremy L Thompson       CeedChk(ierr);
127d1bcdac9Sjeremylt       ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
1284d1cd9fcSJeremy L Thompson       ierr = CeedElemRestrictionGetElementSize(r, &P);
1294d1cd9fcSJeremy L Thompson       CeedChk(ierr);
1304d1cd9fcSJeremy L Thompson       ierr = CeedVectorCreate(ceed, P*ncomp*blksize, &evecs[i]); CeedChk(ierr);
131aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*ncomp*dim*blksize, &qvecs[i]); CeedChk(ierr);
1324a2e7687Sjeremylt       break;
1334a2e7687Sjeremylt     case CEED_EVAL_WEIGHT: // Only on input fields
134aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
135aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChk(ierr);
136d1bcdac9Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
137aedaa0e5Sjeremylt                             CEED_EVAL_WEIGHT, NULL, qvecs[i]); CeedChk(ierr);
138aedaa0e5Sjeremylt 
1394a2e7687Sjeremylt       break;
1404a2e7687Sjeremylt     case CEED_EVAL_DIV:
1414a2e7687Sjeremylt       break; // Not implimented
1424a2e7687Sjeremylt     case CEED_EVAL_CURL:
1434a2e7687Sjeremylt       break; // Not implimented
1444a2e7687Sjeremylt     }
1454a2e7687Sjeremylt   }
1464a2e7687Sjeremylt   return 0;
1474a2e7687Sjeremylt }
1484a2e7687Sjeremylt 
1494a2e7687Sjeremylt /*
1504a2e7687Sjeremylt   CeedOperator needs to connect all the named fields (be they active or passive)
1514a2e7687Sjeremylt   to the named inputs and outputs of its CeedQFunction.
1524a2e7687Sjeremylt  */
1534a2e7687Sjeremylt static int CeedOperatorSetup_Blocked(CeedOperator op) {
1544a2e7687Sjeremylt   int ierr;
1554ce2993fSjeremylt   bool setupdone;
1564ce2993fSjeremylt   ierr = CeedOperatorGetSetupStatus(op, &setupdone); CeedChk(ierr);
1574ce2993fSjeremylt   if (setupdone) return 0;
158aedaa0e5Sjeremylt   Ceed ceed;
159aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
1604ce2993fSjeremylt   CeedOperator_Blocked *impl;
1614ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
1624ce2993fSjeremylt   CeedQFunction qf;
1634ce2993fSjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
1644ce2993fSjeremylt   CeedInt Q, numinputfields, numoutputfields;
1654ce2993fSjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
1664a2e7687Sjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
1674a2e7687Sjeremylt   CeedChk(ierr);
168d1bcdac9Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
169d1bcdac9Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
170d1bcdac9Sjeremylt   CeedChk(ierr);
171d1bcdac9Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
172d1bcdac9Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
173d1bcdac9Sjeremylt   CeedChk(ierr);
1744a2e7687Sjeremylt 
1754a2e7687Sjeremylt   // Allocate
176aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr);
1774a2e7687Sjeremylt   CeedChk(ierr);
178aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs);
1794a2e7687Sjeremylt   CeedChk(ierr);
180aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata);
1814a2e7687Sjeremylt   CeedChk(ierr);
1824a2e7687Sjeremylt 
18316c359e6Sjeremylt   ierr = CeedCalloc(16, &impl->inputstate); CeedChk(ierr);
18491703d3fSjeremylt   ierr = CeedCalloc(16, &impl->evecsin); CeedChk(ierr);
18591703d3fSjeremylt   ierr = CeedCalloc(16, &impl->evecsout); CeedChk(ierr);
186aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsin); CeedChk(ierr);
187aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsout); CeedChk(ierr);
1884a2e7687Sjeremylt 
189aedaa0e5Sjeremylt   impl->numein = numinputfields; impl->numeout = numoutputfields;
190aedaa0e5Sjeremylt 
1914a2e7687Sjeremylt   // Set up infield and outfield pointer arrays
1924a2e7687Sjeremylt   // Infields
193aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 0, impl->blkrestr,
19491703d3fSjeremylt                                          impl->evecs, impl->evecsin,
19591703d3fSjeremylt                                          impl->qvecsin, 0,
196aedaa0e5Sjeremylt                                          numinputfields, Q);
1974a2e7687Sjeremylt   CeedChk(ierr);
1984a2e7687Sjeremylt   // Outfields
199aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 1, impl->blkrestr,
20091703d3fSjeremylt                                          impl->evecs, impl->evecsout,
20191703d3fSjeremylt                                          impl->qvecsout, numinputfields,
20291703d3fSjeremylt                                          numoutputfields, Q);
2034a2e7687Sjeremylt   CeedChk(ierr);
204aedaa0e5Sjeremylt 
2054ce2993fSjeremylt   ierr = CeedOperatorSetSetupDone(op); CeedChk(ierr);
2064a2e7687Sjeremylt 
2074a2e7687Sjeremylt   return 0;
2084a2e7687Sjeremylt }
2094a2e7687Sjeremylt 
2104a2e7687Sjeremylt static int CeedOperatorApply_Blocked(CeedOperator op, CeedVector invec,
2114a2e7687Sjeremylt                                      CeedVector outvec, CeedRequest *request) {
2124a2e7687Sjeremylt   int ierr;
2134ce2993fSjeremylt   CeedOperator_Blocked *impl;
2144ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
2154ce2993fSjeremylt   const CeedInt blksize = 8;
216d1bcdac9Sjeremylt   CeedInt Q, elemsize, numinputfields, numoutputfields, numelements, ncomp;
2174ce2993fSjeremylt   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr);
2184ce2993fSjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
2194ce2993fSjeremylt   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
2204ce2993fSjeremylt   CeedQFunction qf;
2214ce2993fSjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
2224ce2993fSjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
2234ce2993fSjeremylt   CeedChk(ierr);
2244dccadb6Sjeremylt   CeedTransposeMode lmode;
225d1bcdac9Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
226d1bcdac9Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
227d1bcdac9Sjeremylt   CeedChk(ierr);
228d1bcdac9Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
229d1bcdac9Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
230d1bcdac9Sjeremylt   CeedChk(ierr);
231d1bcdac9Sjeremylt   CeedEvalMode emode;
232d1bcdac9Sjeremylt   CeedVector vec;
233d1bcdac9Sjeremylt   CeedBasis basis;
234d1bcdac9Sjeremylt   CeedElemRestriction Erestrict;
23516c359e6Sjeremylt   uint64_t state;
2364a2e7687Sjeremylt 
2374a2e7687Sjeremylt   // Setup
2384a2e7687Sjeremylt   ierr = CeedOperatorSetup_Blocked(op); CeedChk(ierr);
2394a2e7687Sjeremylt 
2404a2e7687Sjeremylt   // Input Evecs and Restriction
2414a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
242d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
243d1bcdac9Sjeremylt     CeedChk(ierr);
2444a2e7687Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
2454a2e7687Sjeremylt     } else {
246d1bcdac9Sjeremylt       // Get input vector
247d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
248d1bcdac9Sjeremylt       if (vec == CEED_VECTOR_ACTIVE)
249d1bcdac9Sjeremylt         vec = invec;
2504a2e7687Sjeremylt       // Restrict
25116c359e6Sjeremylt       ierr = CeedVectorGetState(vec, &state); CeedChk(ierr);
2528d713cf6Sjeremylt       if (state != impl->inputstate[i] || vec == invec) {
2534dccadb6Sjeremylt         ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode); CeedChk(ierr);
2544a2e7687Sjeremylt         ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE,
255d1bcdac9Sjeremylt                                         lmode, vec, impl->evecs[i],
2564a2e7687Sjeremylt                                         request); CeedChk(ierr); CeedChk(ierr);
25716c359e6Sjeremylt         impl->inputstate[i] = state;
25816c359e6Sjeremylt       }
2594a2e7687Sjeremylt       // Get evec
2604a2e7687Sjeremylt       ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST,
2614a2e7687Sjeremylt                                     (const CeedScalar **) &impl->edata[i]);
2624a2e7687Sjeremylt       CeedChk(ierr);
2634a2e7687Sjeremylt     }
2644a2e7687Sjeremylt   }
2654a2e7687Sjeremylt 
2664a2e7687Sjeremylt   // Output Evecs
2674a2e7687Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
2684a2e7687Sjeremylt     ierr = CeedVectorGetArray(impl->evecs[i+impl->numein], CEED_MEM_HOST,
2694a2e7687Sjeremylt                               &impl->edata[i + numinputfields]); CeedChk(ierr);
2704a2e7687Sjeremylt   }
2714a2e7687Sjeremylt 
2724a2e7687Sjeremylt   // Loop through elements
2734a2e7687Sjeremylt   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
2744a2e7687Sjeremylt     // Input basis apply if needed
2754a2e7687Sjeremylt     for (CeedInt i=0; i<numinputfields; i++) {
2764a2e7687Sjeremylt       // Get elemsize, emode, ncomp
277d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict);
278d1bcdac9Sjeremylt       CeedChk(ierr);
279d1bcdac9Sjeremylt       ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
280d1bcdac9Sjeremylt       CeedChk(ierr);
281d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
282d1bcdac9Sjeremylt       CeedChk(ierr);
283d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qfinputfields[i], &ncomp);
284d1bcdac9Sjeremylt       CeedChk(ierr);
2854a2e7687Sjeremylt       // Basis action
2864a2e7687Sjeremylt       switch(emode) {
2874a2e7687Sjeremylt       case CEED_EVAL_NONE:
288aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
289aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
290aedaa0e5Sjeremylt                                   &impl->edata[i][e*Q*ncomp]); CeedChk(ierr);
2914a2e7687Sjeremylt         break;
2924a2e7687Sjeremylt       case CEED_EVAL_INTERP:
293aedaa0e5Sjeremylt         ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
29491703d3fSjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
295aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
296aedaa0e5Sjeremylt                                   &impl->edata[i][e*elemsize*ncomp]);
2974a2e7687Sjeremylt         CeedChk(ierr);
298d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
29991703d3fSjeremylt                               CEED_EVAL_INTERP, impl->evecsin[i],
300aedaa0e5Sjeremylt                               impl->qvecsin[i]); CeedChk(ierr);
3014a2e7687Sjeremylt         break;
3024a2e7687Sjeremylt       case CEED_EVAL_GRAD:
303aedaa0e5Sjeremylt         ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
30491703d3fSjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
305aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
306aedaa0e5Sjeremylt                                   &impl->edata[i][e*elemsize*ncomp]);
3074a2e7687Sjeremylt         CeedChk(ierr);
308d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
30991703d3fSjeremylt                               CEED_EVAL_GRAD, impl->evecsin[i],
310aedaa0e5Sjeremylt                               impl->qvecsin[i]); CeedChk(ierr);
3114a2e7687Sjeremylt         break;
3124a2e7687Sjeremylt       case CEED_EVAL_WEIGHT:
3134a2e7687Sjeremylt         break;  // No action
3144a2e7687Sjeremylt       case CEED_EVAL_DIV:
3154a2e7687Sjeremylt         break; // Not implimented
3164a2e7687Sjeremylt       case CEED_EVAL_CURL:
3174a2e7687Sjeremylt         break; // Not implimented
3184a2e7687Sjeremylt       }
3194a2e7687Sjeremylt     }
3204a2e7687Sjeremylt 
3214a2e7687Sjeremylt     // Output pointers
3224a2e7687Sjeremylt     for (CeedInt i=0; i<numoutputfields; i++) {
323d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
324d1bcdac9Sjeremylt       CeedChk(ierr);
3254a2e7687Sjeremylt       if (emode == CEED_EVAL_NONE) {
326d1bcdac9Sjeremylt         ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp);
327d1bcdac9Sjeremylt         CeedChk(ierr);
328aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST,
329aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
330aedaa0e5Sjeremylt                                   &impl->edata[i + numinputfields][e*Q*ncomp]);
331aedaa0e5Sjeremylt         CeedChk(ierr);
3324a2e7687Sjeremylt       }
3334a2e7687Sjeremylt     }
3344a2e7687Sjeremylt     // Q function
335aedaa0e5Sjeremylt     ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
336aedaa0e5Sjeremylt     CeedChk(ierr);
3374a2e7687Sjeremylt 
3384a2e7687Sjeremylt     // Output basis apply if needed
3394a2e7687Sjeremylt     for (CeedInt i=0; i<numoutputfields; i++) {
3404a2e7687Sjeremylt       // Get elemsize, emode, ncomp
341d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict);
342d1bcdac9Sjeremylt       CeedChk(ierr);
343d1bcdac9Sjeremylt       ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
344d1bcdac9Sjeremylt       CeedChk(ierr);
345d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
346d1bcdac9Sjeremylt       CeedChk(ierr);
347d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp);
348d1bcdac9Sjeremylt       CeedChk(ierr);
3494a2e7687Sjeremylt       // Basis action
3504a2e7687Sjeremylt       switch(emode) {
3514a2e7687Sjeremylt       case CEED_EVAL_NONE:
3524a2e7687Sjeremylt         break; // No action
3534a2e7687Sjeremylt       case CEED_EVAL_INTERP:
354d1bcdac9Sjeremylt         ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
355d1bcdac9Sjeremylt         CeedChk(ierr);
35691703d3fSjeremylt         ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST,
357aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
3584a2e7687Sjeremylt                                   &impl->edata[i + numinputfields][e*elemsize*ncomp]);
359*de686571SJeremy L Thompson         CeedChk(ierr);
360aedaa0e5Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
361aedaa0e5Sjeremylt                               CEED_EVAL_INTERP, impl->qvecsout[i],
36291703d3fSjeremylt                               impl->evecsout[i]); CeedChk(ierr);
3634a2e7687Sjeremylt         break;
3644a2e7687Sjeremylt       case CEED_EVAL_GRAD:
365d1bcdac9Sjeremylt         ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
366d1bcdac9Sjeremylt         CeedChk(ierr);
36791703d3fSjeremylt         ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST,
368aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
369aedaa0e5Sjeremylt                                   &impl->edata[i + numinputfields][e*elemsize*ncomp]);
370*de686571SJeremy L Thompson         CeedChk(ierr);
371d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
372aedaa0e5Sjeremylt                               CEED_EVAL_GRAD, impl->qvecsout[i],
37391703d3fSjeremylt                               impl->evecsout[i]); CeedChk(ierr);
3744a2e7687Sjeremylt         break;
3754ce2993fSjeremylt       case CEED_EVAL_WEIGHT: {
3764ce2993fSjeremylt         Ceed ceed;
3774ce2993fSjeremylt         ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
3784ce2993fSjeremylt         return CeedError(ceed, 1,
3794a2e7687Sjeremylt                          "CEED_EVAL_WEIGHT cannot be an output evaluation mode");
3804a2e7687Sjeremylt         break; // Should not occur
3814ce2993fSjeremylt       }
3824a2e7687Sjeremylt       case CEED_EVAL_DIV:
3834a2e7687Sjeremylt         break; // Not implimented
3844a2e7687Sjeremylt       case CEED_EVAL_CURL:
3854a2e7687Sjeremylt         break; // Not implimented
3864a2e7687Sjeremylt       }
3874a2e7687Sjeremylt     }
3884a2e7687Sjeremylt   }
3894a2e7687Sjeremylt 
3904a2e7687Sjeremylt   // Zero lvecs
391d1bcdac9Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
392d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
39352d6035fSJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
39452d6035fSJeremy L Thompson       if (!impl->add) {
395d1bcdac9Sjeremylt         vec = outvec;
396d1bcdac9Sjeremylt         ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr);
3974a2e7687Sjeremylt       }
39852d6035fSJeremy L Thompson     } else {
39952d6035fSJeremy L Thompson       ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr);
40052d6035fSJeremy L Thompson     }
40152d6035fSJeremy L Thompson   }
40252d6035fSJeremy L Thompson   impl->add = false;
4034a2e7687Sjeremylt 
4044a2e7687Sjeremylt   // Output restriction
4054a2e7687Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
4064a2e7687Sjeremylt     // Restore evec
4074a2e7687Sjeremylt     ierr = CeedVectorRestoreArray(impl->evecs[i+impl->numein],
4084a2e7687Sjeremylt                                   &impl->edata[i + numinputfields]); CeedChk(ierr);
409d1bcdac9Sjeremylt     // Get output vector
410d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
4114a2e7687Sjeremylt     // Active
412d1bcdac9Sjeremylt     if (vec == CEED_VECTOR_ACTIVE)
413d1bcdac9Sjeremylt       vec = outvec;
4144a2e7687Sjeremylt     // Restrict
4154dccadb6Sjeremylt     ierr = CeedOperatorFieldGetLMode(opoutputfields[i], &lmode); CeedChk(ierr);
4164a2e7687Sjeremylt     ierr = CeedElemRestrictionApply(impl->blkrestr[i+impl->numein], CEED_TRANSPOSE,
417d1bcdac9Sjeremylt                                     lmode, impl->evecs[i+impl->numein], vec,
4184a2e7687Sjeremylt                                     request); CeedChk(ierr);
419d1bcdac9Sjeremylt 
4204a2e7687Sjeremylt   }
4214a2e7687Sjeremylt 
4224a2e7687Sjeremylt   // Restore input arrays
4234a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
424d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
425d1bcdac9Sjeremylt     CeedChk(ierr);
4264a2e7687Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
4274a2e7687Sjeremylt     } else {
4284a2e7687Sjeremylt       ierr = CeedVectorRestoreArrayRead(impl->evecs[i],
429d1bcdac9Sjeremylt                                         (const CeedScalar **) &impl->edata[i]);
430d1bcdac9Sjeremylt       CeedChk(ierr);
4314a2e7687Sjeremylt     }
4324a2e7687Sjeremylt   }
4334a2e7687Sjeremylt 
4344a2e7687Sjeremylt   return 0;
4354a2e7687Sjeremylt }
4364a2e7687Sjeremylt 
4374a2e7687Sjeremylt int CeedOperatorCreate_Blocked(CeedOperator op) {
4384a2e7687Sjeremylt   int ierr;
439fe2413ffSjeremylt   Ceed ceed;
440fe2413ffSjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
4414ce2993fSjeremylt   CeedOperator_Blocked *impl;
4424a2e7687Sjeremylt 
4434a2e7687Sjeremylt   ierr = CeedCalloc(1, &impl); CeedChk(ierr);
444*de686571SJeremy L Thompson   ierr = CeedOperatorSetData(op, (void *)&impl); CeedChk(ierr);
445fe2413ffSjeremylt 
446fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Apply",
447fe2413ffSjeremylt                                 CeedOperatorApply_Blocked); CeedChk(ierr);
448fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy",
449fe2413ffSjeremylt                                 CeedOperatorDestroy_Blocked); CeedChk(ierr);
4504a2e7687Sjeremylt   return 0;
4514a2e7687Sjeremylt }
452