xref: /libCEED/rust/libceed-sys/c-src/backends/blocked/ceed-blocked-operator.c (revision aedaa0e5fd3a3e03ad33ad8a6308ac527f4f900e)
14a2e7687Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC.
24a2e7687Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707.
34a2e7687Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details.
44a2e7687Sjeremylt //
54a2e7687Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software
64a2e7687Sjeremylt // libraries and APIs for efficient high-order finite element and spectral
74a2e7687Sjeremylt // element discretizations for exascale applications. For more information and
84a2e7687Sjeremylt // source code availability see http://github.com/ceed.
94a2e7687Sjeremylt //
104a2e7687Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC,
114a2e7687Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office
124a2e7687Sjeremylt // of Science and the National Nuclear Security Administration) responsible for
134a2e7687Sjeremylt // the planning and preparation of a capable exascale ecosystem, including
144a2e7687Sjeremylt // software, applications, hardware, advanced system engineering and early
154a2e7687Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative.
164a2e7687Sjeremylt 
174a2e7687Sjeremylt #include <string.h>
184a2e7687Sjeremylt #include "ceed-blocked.h"
194a2e7687Sjeremylt #include "../ref/ceed-ref.h"
204a2e7687Sjeremylt 
214a2e7687Sjeremylt static int CeedOperatorDestroy_Blocked(CeedOperator op) {
224a2e7687Sjeremylt   int ierr;
234ce2993fSjeremylt   CeedOperator_Blocked *impl;
244ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void*)&impl); CeedChk(ierr);
254a2e7687Sjeremylt 
264a2e7687Sjeremylt   for (CeedInt i=0; i<impl->numein+impl->numeout; i++) {
274a2e7687Sjeremylt     ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChk(ierr);
284a2e7687Sjeremylt     ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChk(ierr);
294a2e7687Sjeremylt   }
304a2e7687Sjeremylt   ierr = CeedFree(&impl->blkrestr); CeedChk(ierr);
314a2e7687Sjeremylt   ierr = CeedFree(&impl->evecs); CeedChk(ierr);
324a2e7687Sjeremylt   ierr = CeedFree(&impl->edata); CeedChk(ierr);
334a2e7687Sjeremylt 
34*aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numein; i++) {
35*aedaa0e5Sjeremylt     ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChk(ierr);
364a2e7687Sjeremylt   }
37*aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsin); CeedChk(ierr);
384a2e7687Sjeremylt 
39*aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numeout; i++) {
40*aedaa0e5Sjeremylt     ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr);
41*aedaa0e5Sjeremylt   }
42*aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsout); CeedChk(ierr);
434a2e7687Sjeremylt 
44fe2413ffSjeremylt   ierr = CeedFree(&impl); CeedChk(ierr);
454a2e7687Sjeremylt   return 0;
464a2e7687Sjeremylt }
474a2e7687Sjeremylt 
484a2e7687Sjeremylt /*
494a2e7687Sjeremylt   Setup infields or outfields
504a2e7687Sjeremylt  */
51fe2413ffSjeremylt static int CeedOperatorSetupFields_Blocked(CeedQFunction qf, CeedOperator op,
52fe2413ffSjeremylt     bool inOrOut,
534a2e7687Sjeremylt     CeedElemRestriction *blkrestr,
54*aedaa0e5Sjeremylt     CeedVector *evecs,
55*aedaa0e5Sjeremylt     CeedVector *qvecs, CeedInt starte,
564a2e7687Sjeremylt     CeedInt numfields, CeedInt Q) {
57*aedaa0e5Sjeremylt   CeedInt dim, ierr, ncomp;
58*aedaa0e5Sjeremylt   Ceed ceed;
59*aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
60d1bcdac9Sjeremylt   CeedBasis basis;
61d1bcdac9Sjeremylt   CeedElemRestriction r;
62*aedaa0e5Sjeremylt   CeedOperatorField *opfields;
63*aedaa0e5Sjeremylt   CeedQFunctionField *qffields;
64fe2413ffSjeremylt   if (inOrOut) {
65*aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, NULL, &opfields);
66fe2413ffSjeremylt     CeedChk(ierr);
67*aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, NULL, &qffields);
68fe2413ffSjeremylt     CeedChk(ierr);
69fe2413ffSjeremylt   } else {
70*aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, &opfields, NULL);
71fe2413ffSjeremylt     CeedChk(ierr);
72*aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, &qffields, NULL);
73fe2413ffSjeremylt     CeedChk(ierr);
74fe2413ffSjeremylt   }
754a2e7687Sjeremylt   const CeedInt blksize = 8;
764a2e7687Sjeremylt 
774a2e7687Sjeremylt   // Loop over fields
784a2e7687Sjeremylt   for (CeedInt i=0; i<numfields; i++) {
79d1bcdac9Sjeremylt     CeedEvalMode emode;
80*aedaa0e5Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChk(ierr);
814a2e7687Sjeremylt 
824a2e7687Sjeremylt     if (emode != CEED_EVAL_WEIGHT) {
83*aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r);
84d1bcdac9Sjeremylt       CeedChk(ierr);
85fe2413ffSjeremylt       CeedElemRestriction_Ref *data;
86fe2413ffSjeremylt       ierr = CeedElemRestrictionGetData(r, (void *)&data);
874ce2993fSjeremylt       Ceed ceed;
884ce2993fSjeremylt       ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChk(ierr);
894ce2993fSjeremylt       CeedInt nelem, elemsize, ndof, ncomp;
904ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChk(ierr);
914ce2993fSjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChk(ierr);
924ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumDoF(r, &ndof); CeedChk(ierr);
934ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChk(ierr);
944ce2993fSjeremylt       ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize,
954ce2993fSjeremylt                                               blksize, ndof, ncomp,
964a2e7687Sjeremylt                                               CEED_MEM_HOST, CEED_COPY_VALUES,
97*aedaa0e5Sjeremylt                                               data->indices, &blkrestr[i+starte]);
984a2e7687Sjeremylt       CeedChk(ierr);
99*aedaa0e5Sjeremylt       ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL,
100*aedaa0e5Sjeremylt                                              &evecs[i+starte]);
1014a2e7687Sjeremylt       CeedChk(ierr);
1024a2e7687Sjeremylt     }
1034a2e7687Sjeremylt 
1044a2e7687Sjeremylt     switch(emode) {
1054a2e7687Sjeremylt     case CEED_EVAL_NONE:
106*aedaa0e5Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp);
107d1bcdac9Sjeremylt       CeedChk(ierr);
108*aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*ncomp, &qvecs[i]); CeedChk(ierr);
109*aedaa0e5Sjeremylt       break;
110*aedaa0e5Sjeremylt     case CEED_EVAL_INTERP:
111*aedaa0e5Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp);
112*aedaa0e5Sjeremylt       CeedChk(ierr);
113*aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*ncomp*blksize, &qvecs[i]); CeedChk(ierr);
1144a2e7687Sjeremylt       break;
1154a2e7687Sjeremylt     case CEED_EVAL_GRAD:
116*aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
117*aedaa0e5Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp);
118d1bcdac9Sjeremylt       ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
119*aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*ncomp*dim*blksize, &qvecs[i]); CeedChk(ierr);
1204a2e7687Sjeremylt       break;
1214a2e7687Sjeremylt     case CEED_EVAL_WEIGHT: // Only on input fields
122*aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
123*aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChk(ierr);
124d1bcdac9Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
125*aedaa0e5Sjeremylt                             CEED_EVAL_WEIGHT, NULL, qvecs[i]); CeedChk(ierr);
126*aedaa0e5Sjeremylt 
1274a2e7687Sjeremylt       break;
1284a2e7687Sjeremylt     case CEED_EVAL_DIV:
1294a2e7687Sjeremylt       break; // Not implimented
1304a2e7687Sjeremylt     case CEED_EVAL_CURL:
1314a2e7687Sjeremylt       break; // Not implimented
1324a2e7687Sjeremylt     }
1334a2e7687Sjeremylt   }
1344a2e7687Sjeremylt   return 0;
1354a2e7687Sjeremylt }
1364a2e7687Sjeremylt 
1374a2e7687Sjeremylt /*
1384a2e7687Sjeremylt   CeedOperator needs to connect all the named fields (be they active or passive)
1394a2e7687Sjeremylt   to the named inputs and outputs of its CeedQFunction.
1404a2e7687Sjeremylt  */
1414a2e7687Sjeremylt static int CeedOperatorSetup_Blocked(CeedOperator op) {
1424a2e7687Sjeremylt   int ierr;
1434ce2993fSjeremylt   bool setupdone;
1444ce2993fSjeremylt   ierr = CeedOperatorGetSetupStatus(op, &setupdone); CeedChk(ierr);
1454ce2993fSjeremylt   if (setupdone) return 0;
146*aedaa0e5Sjeremylt   Ceed ceed;
147*aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
1484ce2993fSjeremylt   CeedOperator_Blocked *impl;
1494ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void*)&impl); CeedChk(ierr);
1504ce2993fSjeremylt   CeedQFunction qf;
1514ce2993fSjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
1524ce2993fSjeremylt   CeedInt Q, numinputfields, numoutputfields;
1534ce2993fSjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
1544a2e7687Sjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
1554a2e7687Sjeremylt   CeedChk(ierr);
156d1bcdac9Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
157d1bcdac9Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
158d1bcdac9Sjeremylt   CeedChk(ierr);
159d1bcdac9Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
160d1bcdac9Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
161d1bcdac9Sjeremylt   CeedChk(ierr);
1624a2e7687Sjeremylt 
1634a2e7687Sjeremylt   // Allocate
164*aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr);
1654a2e7687Sjeremylt   CeedChk(ierr);
166*aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs);
1674a2e7687Sjeremylt   CeedChk(ierr);
168*aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata);
1694a2e7687Sjeremylt   CeedChk(ierr);
1704a2e7687Sjeremylt 
171*aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsin); CeedChk(ierr);
172*aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsout); CeedChk(ierr);
1734a2e7687Sjeremylt 
174*aedaa0e5Sjeremylt   impl->numein = numinputfields; impl->numeout = numoutputfields;
175*aedaa0e5Sjeremylt 
1764a2e7687Sjeremylt   // Set up infield and outfield pointer arrays
1774a2e7687Sjeremylt   // Infields
178*aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 0, impl->blkrestr,
179*aedaa0e5Sjeremylt                                          impl->evecs, impl->qvecsin, 0,
180*aedaa0e5Sjeremylt                                          numinputfields, Q);
1814a2e7687Sjeremylt   CeedChk(ierr);
1824a2e7687Sjeremylt   // Outfields
183*aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 1, impl->blkrestr,
184*aedaa0e5Sjeremylt                                          impl->evecs, impl->qvecsout,
185*aedaa0e5Sjeremylt                                          numinputfields, numoutputfields, Q);
1864a2e7687Sjeremylt   CeedChk(ierr);
187*aedaa0e5Sjeremylt 
188*aedaa0e5Sjeremylt   // Temporary Vector
189*aedaa0e5Sjeremylt   ierr = CeedVectorCreate(ceed, 0, &impl->tempvec); CeedChk(ierr);
1904a2e7687Sjeremylt 
1914ce2993fSjeremylt   ierr = CeedOperatorSetSetupDone(op); CeedChk(ierr);
1924a2e7687Sjeremylt 
1934a2e7687Sjeremylt   return 0;
1944a2e7687Sjeremylt }
1954a2e7687Sjeremylt 
1964a2e7687Sjeremylt static int CeedOperatorApply_Blocked(CeedOperator op, CeedVector invec,
1974a2e7687Sjeremylt                                      CeedVector outvec, CeedRequest *request) {
1984a2e7687Sjeremylt   int ierr;
1994ce2993fSjeremylt   CeedOperator_Blocked *impl;
2004ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void*)&impl); CeedChk(ierr);
2014ce2993fSjeremylt   const CeedInt blksize = 8;
202d1bcdac9Sjeremylt   CeedInt Q, elemsize, numinputfields, numoutputfields, numelements, ncomp;
2034ce2993fSjeremylt   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr);
2044ce2993fSjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
2054ce2993fSjeremylt   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
2064ce2993fSjeremylt   CeedQFunction qf;
2074ce2993fSjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
2084ce2993fSjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
2094ce2993fSjeremylt   CeedChk(ierr);
2104dccadb6Sjeremylt   CeedTransposeMode lmode;
211d1bcdac9Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
212d1bcdac9Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
213d1bcdac9Sjeremylt   CeedChk(ierr);
214d1bcdac9Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
215d1bcdac9Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
216d1bcdac9Sjeremylt   CeedChk(ierr);
217d1bcdac9Sjeremylt   CeedEvalMode emode;
218d1bcdac9Sjeremylt   CeedVector vec;
219d1bcdac9Sjeremylt   CeedBasis basis;
220d1bcdac9Sjeremylt   CeedElemRestriction Erestrict;
2214a2e7687Sjeremylt 
2224a2e7687Sjeremylt   // Setup
2234a2e7687Sjeremylt   ierr = CeedOperatorSetup_Blocked(op); CeedChk(ierr);
2244a2e7687Sjeremylt 
2254a2e7687Sjeremylt   // Input Evecs and Restriction
2264a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
227d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
228d1bcdac9Sjeremylt     CeedChk(ierr);
2294a2e7687Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
2304a2e7687Sjeremylt     } else {
231d1bcdac9Sjeremylt       // Get input vector
232d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
233d1bcdac9Sjeremylt       if (vec == CEED_VECTOR_ACTIVE)
234d1bcdac9Sjeremylt         vec = invec;
2354a2e7687Sjeremylt       // Restrict
2364dccadb6Sjeremylt       ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode); CeedChk(ierr);
2374a2e7687Sjeremylt       ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE,
238d1bcdac9Sjeremylt                                       lmode, vec, impl->evecs[i],
2394a2e7687Sjeremylt                                       request); CeedChk(ierr); CeedChk(ierr);
2404a2e7687Sjeremylt       // Get evec
2414a2e7687Sjeremylt       ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST,
2424a2e7687Sjeremylt                                     (const CeedScalar **) &impl->edata[i]);
2434a2e7687Sjeremylt       CeedChk(ierr);
2444a2e7687Sjeremylt     }
2454a2e7687Sjeremylt   }
2464a2e7687Sjeremylt 
2474a2e7687Sjeremylt   // Output Evecs
2484a2e7687Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
2494a2e7687Sjeremylt     ierr = CeedVectorGetArray(impl->evecs[i+impl->numein], CEED_MEM_HOST,
2504a2e7687Sjeremylt                               &impl->edata[i + numinputfields]); CeedChk(ierr);
2514a2e7687Sjeremylt   }
2524a2e7687Sjeremylt 
2534a2e7687Sjeremylt   // Loop through elements
2544a2e7687Sjeremylt   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
2554a2e7687Sjeremylt     // Input basis apply if needed
2564a2e7687Sjeremylt     for (CeedInt i=0; i<numinputfields; i++) {
2574a2e7687Sjeremylt       // Get elemsize, emode, ncomp
258d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict);
259d1bcdac9Sjeremylt       CeedChk(ierr);
260d1bcdac9Sjeremylt       ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
261d1bcdac9Sjeremylt       CeedChk(ierr);
262d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
263d1bcdac9Sjeremylt       CeedChk(ierr);
264d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qfinputfields[i], &ncomp);
265d1bcdac9Sjeremylt       CeedChk(ierr);
2664a2e7687Sjeremylt       // Basis action
2674a2e7687Sjeremylt       switch(emode) {
2684a2e7687Sjeremylt       case CEED_EVAL_NONE:
269*aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
270*aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
271*aedaa0e5Sjeremylt                                   &impl->edata[i][e*Q*ncomp]); CeedChk(ierr);
2724a2e7687Sjeremylt         break;
2734a2e7687Sjeremylt       case CEED_EVAL_INTERP:
274*aedaa0e5Sjeremylt         ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
275*aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST,
276*aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
277*aedaa0e5Sjeremylt                                   &impl->edata[i][e*elemsize*ncomp]);
2784a2e7687Sjeremylt         CeedChk(ierr);
279d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
280*aedaa0e5Sjeremylt                               CEED_EVAL_INTERP, impl->tempvec,
281*aedaa0e5Sjeremylt                               impl->qvecsin[i]); CeedChk(ierr);
2824a2e7687Sjeremylt         break;
2834a2e7687Sjeremylt       case CEED_EVAL_GRAD:
284*aedaa0e5Sjeremylt         ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
285*aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST,
286*aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
287*aedaa0e5Sjeremylt                                   &impl->edata[i][e*elemsize*ncomp]);
2884a2e7687Sjeremylt         CeedChk(ierr);
289d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
290*aedaa0e5Sjeremylt                               CEED_EVAL_GRAD, impl->tempvec,
291*aedaa0e5Sjeremylt                               impl->qvecsin[i]); CeedChk(ierr);
2924a2e7687Sjeremylt         break;
2934a2e7687Sjeremylt       case CEED_EVAL_WEIGHT:
2944a2e7687Sjeremylt         break;  // No action
2954a2e7687Sjeremylt       case CEED_EVAL_DIV:
2964a2e7687Sjeremylt         break; // Not implimented
2974a2e7687Sjeremylt       case CEED_EVAL_CURL:
2984a2e7687Sjeremylt         break; // Not implimented
2994a2e7687Sjeremylt       }
3004a2e7687Sjeremylt     }
3014a2e7687Sjeremylt 
3024a2e7687Sjeremylt     // Output pointers
3034a2e7687Sjeremylt     for (CeedInt i=0; i<numoutputfields; i++) {
304d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
305d1bcdac9Sjeremylt       CeedChk(ierr);
3064a2e7687Sjeremylt       if (emode == CEED_EVAL_NONE) {
307d1bcdac9Sjeremylt         ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp);
308d1bcdac9Sjeremylt         CeedChk(ierr);
309*aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST,
310*aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
311*aedaa0e5Sjeremylt                                   &impl->edata[i + numinputfields][e*Q*ncomp]);
312*aedaa0e5Sjeremylt         CeedChk(ierr);
3134a2e7687Sjeremylt       }
3144a2e7687Sjeremylt     }
3154a2e7687Sjeremylt     // Q function
316*aedaa0e5Sjeremylt     ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
317*aedaa0e5Sjeremylt     CeedChk(ierr);
3184a2e7687Sjeremylt 
3194a2e7687Sjeremylt     // Output basis apply if needed
3204a2e7687Sjeremylt     for (CeedInt i=0; i<numoutputfields; i++) {
3214a2e7687Sjeremylt       // Get elemsize, emode, ncomp
322d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict);
323d1bcdac9Sjeremylt       CeedChk(ierr);
324d1bcdac9Sjeremylt       ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
325d1bcdac9Sjeremylt       CeedChk(ierr);
326d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
327d1bcdac9Sjeremylt       CeedChk(ierr);
328d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp);
329d1bcdac9Sjeremylt       CeedChk(ierr);
3304a2e7687Sjeremylt       // Basis action
3314a2e7687Sjeremylt       switch(emode) {
3324a2e7687Sjeremylt       case CEED_EVAL_NONE:
3334a2e7687Sjeremylt         break; // No action
3344a2e7687Sjeremylt       case CEED_EVAL_INTERP:
335d1bcdac9Sjeremylt         ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
336d1bcdac9Sjeremylt         CeedChk(ierr);
337*aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST,
338*aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
3394a2e7687Sjeremylt                                   &impl->edata[i + numinputfields][e*elemsize*ncomp]);
340*aedaa0e5Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
341*aedaa0e5Sjeremylt                               CEED_EVAL_INTERP, impl->qvecsout[i],
342*aedaa0e5Sjeremylt                               impl->tempvec); CeedChk(ierr);
3434a2e7687Sjeremylt         break;
3444a2e7687Sjeremylt       case CEED_EVAL_GRAD:
345d1bcdac9Sjeremylt         ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
346d1bcdac9Sjeremylt         CeedChk(ierr);
347*aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->tempvec, CEED_MEM_HOST,
348*aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
349*aedaa0e5Sjeremylt                                   &impl->edata[i + numinputfields][e*elemsize*ncomp]);
350d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
351*aedaa0e5Sjeremylt                               CEED_EVAL_GRAD, impl->qvecsout[i],
352*aedaa0e5Sjeremylt                               impl->tempvec); CeedChk(ierr);
3534a2e7687Sjeremylt         break;
3544ce2993fSjeremylt       case CEED_EVAL_WEIGHT: {
3554ce2993fSjeremylt         Ceed ceed;
3564ce2993fSjeremylt         ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
3574ce2993fSjeremylt         return CeedError(ceed, 1,
3584a2e7687Sjeremylt                          "CEED_EVAL_WEIGHT cannot be an output evaluation mode");
3594a2e7687Sjeremylt         break; // Should not occur
3604ce2993fSjeremylt       }
3614a2e7687Sjeremylt       case CEED_EVAL_DIV:
3624a2e7687Sjeremylt         break; // Not implimented
3634a2e7687Sjeremylt       case CEED_EVAL_CURL:
3644a2e7687Sjeremylt         break; // Not implimented
3654a2e7687Sjeremylt       }
3664a2e7687Sjeremylt     }
3674a2e7687Sjeremylt   }
3684a2e7687Sjeremylt 
3694a2e7687Sjeremylt   // Zero lvecs
370d1bcdac9Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
371d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
372d1bcdac9Sjeremylt     if (vec == CEED_VECTOR_ACTIVE)
373d1bcdac9Sjeremylt       vec = outvec;
374d1bcdac9Sjeremylt     ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr);
3754a2e7687Sjeremylt   }
3764a2e7687Sjeremylt 
3774a2e7687Sjeremylt   // Output restriction
3784a2e7687Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
3794a2e7687Sjeremylt     // Restore evec
3804a2e7687Sjeremylt     ierr = CeedVectorRestoreArray(impl->evecs[i+impl->numein],
3814a2e7687Sjeremylt                                   &impl->edata[i + numinputfields]); CeedChk(ierr);
382d1bcdac9Sjeremylt     // Get output vector
383d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
3844a2e7687Sjeremylt     // Active
385d1bcdac9Sjeremylt     if (vec == CEED_VECTOR_ACTIVE)
386d1bcdac9Sjeremylt       vec = outvec;
3874a2e7687Sjeremylt     // Restrict
3884dccadb6Sjeremylt     ierr = CeedOperatorFieldGetLMode(opoutputfields[i], &lmode); CeedChk(ierr);
3894a2e7687Sjeremylt     ierr = CeedElemRestrictionApply(impl->blkrestr[i+impl->numein], CEED_TRANSPOSE,
390d1bcdac9Sjeremylt                                     lmode, impl->evecs[i+impl->numein], vec,
3914a2e7687Sjeremylt                                     request); CeedChk(ierr);
392d1bcdac9Sjeremylt 
3934a2e7687Sjeremylt   }
3944a2e7687Sjeremylt 
3954a2e7687Sjeremylt   // Restore input arrays
3964a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
397d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
398d1bcdac9Sjeremylt     CeedChk(ierr);
3994a2e7687Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
4004a2e7687Sjeremylt     } else {
4014a2e7687Sjeremylt       ierr = CeedVectorRestoreArrayRead(impl->evecs[i],
402d1bcdac9Sjeremylt                                         (const CeedScalar **) &impl->edata[i]);
403d1bcdac9Sjeremylt       CeedChk(ierr);
4044a2e7687Sjeremylt     }
4054a2e7687Sjeremylt   }
4064a2e7687Sjeremylt 
4074a2e7687Sjeremylt   return 0;
4084a2e7687Sjeremylt }
4094a2e7687Sjeremylt 
4104a2e7687Sjeremylt int CeedOperatorCreate_Blocked(CeedOperator op) {
4114a2e7687Sjeremylt   int ierr;
412fe2413ffSjeremylt   Ceed ceed;
413fe2413ffSjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
4144ce2993fSjeremylt   CeedOperator_Blocked *impl;
4154a2e7687Sjeremylt 
4164a2e7687Sjeremylt   ierr = CeedCalloc(1, &impl); CeedChk(ierr);
417fe2413ffSjeremylt   ierr = CeedOperatorSetData(op, (void *)&impl);
418fe2413ffSjeremylt 
419fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Apply",
420fe2413ffSjeremylt                                 CeedOperatorApply_Blocked); CeedChk(ierr);
421fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy",
422fe2413ffSjeremylt                                 CeedOperatorDestroy_Blocked); CeedChk(ierr);
4234a2e7687Sjeremylt   return 0;
4244a2e7687Sjeremylt }
425